{"uid":"cap_WlhuOnFZzl1YPOdUc0HJC","slug":"zeroreader-llama-4-scout-17b-chat-completion-a7756e88","name":"ZeroReader Llama 4 Scout 17B Chat Completion","description":"Llama 4 Scout 17B — Latest Llama 4 architecture. 17B with 16 experts.","url":"https://api.zeroreader.com/v1/ai/llama-4-scout?utm_source=zero.xyz","method":"POST","headers":{},"bodySchema":{"type":"object","properties":{"stream":{"type":"boolean","default":false},"messages":{"type":"array","items":{"type":"object","required":["role","content"],"properties":{"role":{"enum":["system","user","assistant"],"type":"string"},"content":{"type":"string"}}}},"max_tokens":{"type":"integer","default":1024,"maximum":4096},"temperature":{"type":"number","default":0.7,"maximum":2,"minimum":0}}},"responseSchema":{"id":"chatcmpl-example","object":"chat.completion","choices":[{"index":0,"message":{"role":"assistant","content":"I'm doing well!"},"finish_reason":"stop"}]},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.005","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.005/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.005","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.005","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_DR37XOA8tir9PFVvgS2DH","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.005","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Runs chat completions using Meta's Llama 4 Scout 17B mixture-of-experts model via ZeroReader's API, paid per call with USDC.","exampleAgentPrompt":"Send these messages to Llama 4 Scout 17B and get a response: system says 'You are a helpful assistant', user asks 'Explain mixture-of-experts models in simple terms' — use a temperature of 0.7 and limit the output to 512 tokens.","exampleUseCases":null,"resultDescription":"A chat completion object (OpenAI-compatible format) containing the assistant's generated reply, the finish reason (e.g. 'stop'), and metadata like the completion ID and choice index.","failureModes":["Payment failure if insufficient USDC balance — 402 Payment Required","Token limit exceeded if max_tokens > 4096 — validation error","Malformed messages array (missing role or content) — 400 Bad Request","Temperature out of range (>2 or <0) — 400 Bad Request","Model overload or timeout — 503 Service Unavailable"],"whenToPreferThis":"Choose this endpoint when you need fast, cost-effective chat completions from the latest Llama 4 architecture with mixture-of-experts efficiency, and you want to pay per-call in USDC without a subscription. Prefer it over DeepSeek R1 for general conversation/generation tasks (not heavy math/logic), and over larger models when latency and cost matter more than maximum capability.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-10-02T01:16:30.433Z","isFirstParty":false,"canonicalSlug":"zeroreader-llama-4-scout-17b-chat-completion-a7756e88"}