{"uid":"cap_vv-Orm9xFPxRxuGo5nVlS","slug":"openai-gpt-6-luna-chat-completions-x402-272k-tokens-ebb0fffa","name":"OpenAI GPT-6 Luna Chat Completions (x402, ≤272K tokens)","description":"gpt-6-luna: OpenAI GPT 6 Luna (<=272K input tokens) chat completions, paid per call in USDC. Per 1M tokens in/out: Flex (default) 0.05/0.25, Standard 0.1/0.5, Fast 0.2/1.0; pick with service_tier. Plus $0.0005 (Base) or $0.0005 (Solana) per call; the quote charges input plus 10% of max_tokens. Levels: none, low, medium, high, xhigh (model gpt-6-luna:<level>). Empty answers are free. Rates: https://openai.mm.family/x402/pricing","url":"https://openai.mm.family/x402/v1/models/gpt-6-luna/chat/completions?utm_source=zero.xyz","method":"POST","headers":{},"bodySchema":{"type":"object","properties":{"model":{"enum":["gpt-6-luna","gpt-6-luna:none","gpt-6-luna:low","gpt-6-luna:medium","gpt-6-luna:high","gpt-6-luna:xhigh"],"type":"string"},"tools":{"type":"array"},"stream":{"type":"boolean"},"messages":{"type":"array","items":{"type":"object","required":["role","content"],"properties":{"role":{"type":"string"},"content":{"type":"string"}}}},"max_tokens":{"type":"integer"},"service_tier":{"enum":["flex","default","standard","auto","fast","priority"],"type":"string"},"response_format":{"type":"object"},"reasoning_effort":{"type":"string"}}},"responseSchema":{"type":"json","example":{"id":"chatcmpl-x","model":"gpt-6-luna","usage":{"total_tokens":9,"prompt_tokens":8,"completion_tokens":1},"object":"chat.completion","choices":[{"index":0,"message":{"role":"assistant","content":"Hello","refusal":null},"finish_reason":"stop"}],"service_tier":"flex"}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.001","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.001/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.001","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.001","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_tYo8fCA7yDtVEmXRql2o4","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.001","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Runs GPT-6 Luna chat completions (up to 272K input tokens) via a pay-per-call x402 USDC-billed API with selectable speed/quality tiers","exampleAgentPrompt":"Use GPT-6 Luna on the standard service tier with up to 1000 output tokens to answer this question: what are the key differences between transformer and state-space models?","exampleUseCases":[{"title":"Cheap reasoning for agentic pipelines","prompt":"Send this multi-step reasoning task to GPT-6 Luna using the flex tier and up to 2000 tokens — I want the cheapest option that still gives me solid reasoning quality."},{"title":"Fast customer support response drafting","prompt":"Draft a polite, professional reply to this customer complaint using GPT-6 Luna on the fast tier, keeping it under 500 tokens so I get a quick turnaround."},{"title":"Long-document summarization with large context","prompt":"Summarize this 200,000-token research document using GPT-6 Luna on the standard tier — I need the full context window so nothing gets cut off, with up to 1500 tokens in the summary."}],"resultDescription":"A JSON chat completion object containing the assistant's reply message, finish reason (e.g. 'stop'), prompt and completion token counts, the model variant used, the service_tier applied, and a unique completion ID. Empty responses when the model produces no output are billed at zero.","failureModes":["Insufficient USDC balance or x402 payment failure — request rejected before inference","max_tokens set too high relative to balance — quote may exceed available funds","Input token count exceeds 272K limit — use the gpt-6-luna-long sibling endpoint instead","Invalid service_tier or model enum value — 422 validation error","Streaming interrupted mid-response — partial content returned","Rate limits or upstream OpenAI capacity issues — 429 or 503 responses"],"whenToPreferThis":"Choose this endpoint when you need GPT-6 Luna's capabilities within a 272K-token input window and want pay-per-call USDC billing via the x402 protocol with no subscription. Prefer the flex tier for cost-sensitive batch workloads, standard for balanced throughput, and fast/priority for latency-sensitive tasks. Use the gpt-6-luna-long sibling if your input exceeds 272K tokens. Prefer this over gpt-5.6-luna when you need GPT-6-class reasoning quality.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-10-02T12:41:24.468Z","isFirstParty":false,"canonicalSlug":"openai-gpt-6-luna-chat-completions-x402-272k-tokens-ebb0fffa"}