{"uid":"cap_RcTG0H_gMBMmVddHtHUR2","slug":"openai-gpt-6-astra-chat-completions-via-mm-family-x402-b53058d4","name":"OpenAI GPT-6 Astra Chat Completions (via mm.family x402)","description":"gpt-6-astra: OpenAI GPT 6 Astra (<=272K input tokens) chat completions, paid per call in USDC. Per 1M tokens in/out: Flex (default) 5.0/25.0, Standard 10.0/50.0, Fast 20.0/100.0; pick with service_tier. Plus $0.0005 (Base) or $0.0005 (Solana) per call; the quote charges input plus 10% of max_tokens. Levels: low, medium, high, xhigh (model gpt-6-astra:<level>). Empty answers are free. Rates: https://openai.mm.family/x402/pricing","url":"https://openai.mm.family/x402/v1/models/gpt-6-astra/chat/completions?utm_source=zero.xyz","method":"POST","headers":{},"bodySchema":{"type":"object","properties":{"model":{"enum":["gpt-6-astra","gpt-6-astra:low","gpt-6-astra:medium","gpt-6-astra:high","gpt-6-astra:xhigh"],"type":"string"},"tools":{"type":"array"},"stream":{"type":"boolean"},"messages":{"type":"array","items":{"type":"object","required":["role","content"],"properties":{"role":{"type":"string"},"content":{"type":"string"}}}},"max_tokens":{"type":"integer"},"service_tier":{"enum":["flex","default","standard","auto","fast","priority"],"type":"string"},"response_format":{"type":"object"},"reasoning_effort":{"type":"string"}}},"responseSchema":{"type":"json","example":{"id":"chatcmpl-x","model":"gpt-6-astra","usage":{"total_tokens":9,"prompt_tokens":8,"completion_tokens":1},"object":"chat.completion","choices":[{"index":0,"message":{"role":"assistant","content":"Hello","refusal":null},"finish_reason":"stop"}],"service_tier":"flex"}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.00181","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.00181/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.00181","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.00181","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_-SNVM0MitPB66j2yeTIMI","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.00181","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Provides GPT-6 Astra chat completions with up to 272K input tokens, billed per call in USDC via x402 payment protocol with selectable service tiers","exampleAgentPrompt":"Using GPT-6 Astra on the mm.family x402 gateway, send this conversation to the model with flex pricing and up to 1000 output tokens: system 'You are a helpful assistant', user 'Explain quantum entanglement in simple terms'.","exampleUseCases":[{"title":"Complex legal document summarization","prompt":"Use GPT-6 Astra with standard tier to summarize this 200-page contract into key clauses and risks — set max_tokens to 2000 so I get a thorough answer."},{"title":"Autonomous agent reasoning task","prompt":"I need the xhigh reasoning level of GPT-6 Astra to solve this multi-step math proof — run it on flex tier and give me up to 4000 tokens to work through it."},{"title":"Customer support chatbot response","prompt":"Route this customer message through GPT-6 Astra on fast tier with max 500 tokens so I get a quick, high-quality support reply: the customer says their order hasn't arrived after two weeks."}],"resultDescription":"Returns a JSON chat completion object containing the assistant's reply message, finish reason (e.g. 'stop'), token usage breakdown (prompt, completion, total tokens), model identifier, and service tier used. The response mirrors the OpenAI chat completions API format.","failureModes":["Insufficient USDC balance — payment rejected before inference runs","Exceeded 272K token context window — use gpt-6.1-sol-long sibling instead","Invalid model variant string — must be one of the gpt-6-astra enum values","max_tokens too large for selected service tier capacity","Stream mode requested but client not handling chunked responses","Empty response returned free of charge when model produces no output"],"whenToPreferThis":"Choose this endpoint when you need access to GPT-6 Astra's frontier-level reasoning with up to 272K input context, want per-call USDC micropayment billing via x402 (no subscription), and want to tune cost/speed tradeoffs using flex, standard, or fast service tiers. Prefer it over cheaper sibling models (gpt-5-nano, gpt-5.4-mini) when task complexity demands the most capable model. Use the gpt-6.1-sol-long sibling if your input exceeds 272K tokens.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-10-02T12:40:57.359Z","isFirstParty":false,"canonicalSlug":"openai-gpt-6-astra-chat-completions-via-mm-family-x402-b53058d4"}