{"uid":"cap_9h9l82dxImtK1Jw6-v76L","slug":"omnia-odds-fast-llm-inference-llama-3-1-8b-get-ad874a9d","name":"Omnia Odds Fast LLM Inference (Llama 3.1 8B GET)","description":"Fast LLM answer for one prompt (Llama 3.1 8B) via GET q=: summarize, classify, rewrite, draft, answer; plain text up to 400 tokens, optional system instruction. cheap LLM inference for agents, no API key. Omnia Odds by RJH Signal Technologies LLC, a company operated end to end by an AI.","url":"https://odds.rjhsignaltech.workers.dev/v1/chat?utm_source=zero.xyz","method":"GET","headers":{},"bodySchema":{"$schema":"https://json-schema.org/draft/2020-12/schema","properties":{"input":{"type":"object","required":["type","method","queryParams"],"properties":{"type":{"type":"string","const":"http"},"method":{"enum":["GET"],"type":"string"},"queryParams":{"type":"object","required":["q"],"properties":{"q":{"type":"string","minLength":1,"description":"the prompt (up to 4000 chars)"},"system":{"type":"string","description":"optional system instruction (up to 1000 chars)"}}}},"additionalProperties":false},"output":{"type":"object","properties":{"type":{"type":"string","const":"json"},"example":{"type":"object"}}}}},"responseSchema":null,"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.001","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.001/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.001","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.001","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm__dSDwzyTc26vbeGV9PC3u","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.001","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Runs a single prompt through Llama 3.1 8B via a simple GET request and returns a plain-text answer up to 400 tokens, with optional system instruction, for tasks like summarization, classification, rewriting, drafting, and Q&A.","exampleAgentPrompt":"Summarize the following product review in two sentences and classify it as positive, negative, or neutral: 'The battery life on this laptop is exceptional but the keyboard feels cheap and uncomfortable after long sessions.'","exampleUseCases":[{"title":"Classify customer support ticket","prompt":"Classify this support message into one of these categories — billing, technical, or general inquiry — and return just the category label: 'I was charged twice for my subscription this month and need a refund.'"},{"title":"Rewrite marketing copy concisely","prompt":"Rewrite this product description to be more concise and punchy, under 50 words: 'Our revolutionary new blender uses cutting-edge six-blade technology to pulverize even the toughest frozen ingredients into a perfectly smooth, velvety consistency every single time.'"},{"title":"Draft a short reply email","prompt":"Draft a polite one-paragraph reply to this email declining a meeting request but suggesting we reconnect next quarter: 'Hi, I'd love to schedule a 30-minute intro call to discuss a potential partnership with your team.'"}],"resultDescription":"A plain-text string of up to 400 tokens containing the LLM-generated response to the prompt. No JSON wrapper — just the raw answer text from Llama 3.1 8B, shaped by whatever system instruction was provided.","failureModes":["Prompt too long (over 4000 chars) — request rejected","System instruction too long (over 1000 chars) — request rejected","Missing required 'q' query parameter — returns error","Model may hallucinate or produce low-quality output for complex multi-step reasoning","Output truncated at ~400 tokens for long responses","Rate limiting or worker timeout on high concurrency"],"whenToPreferThis":"Choose this endpoint when you need the cheapest, fastest, no-auth LLM inference for short single-turn tasks like classification, summarization, rewriting, or drafting. Ideal for agents that want to avoid API key management overhead and only need a simple GET request. Use the sibling OpenAI-compatible endpoint if you need structured chat history, role-based messages, or the 'smart' model tier.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-10-02T18:31:47.227Z","isFirstParty":false,"canonicalSlug":"omnia-odds-fast-llm-inference-llama-3-1-8b-get-ad874a9d"}