{"uid":"cap_e0uF6n5T56Y_Tj8IMOokJ","slug":"openai-compatible-chat-completions-via-x402-unykorn-local-gpu-7dc1b3ac","name":"OpenAI-Compatible Chat Completions via x402 (UnyKorn Local GPU)","description":"OpenAI-compatible chat completions over x402 (local GPU models, no API key) — Genesis402 / UnyKorn Operator Network","url":"https://twin.unykorn.org/v1/chat/completions?utm_source=zero.xyz","method":"POST","headers":{},"bodySchema":{"type":"object","properties":{"model":{"type":"string"},"prompt":{"type":"string"},"messages":{"type":"array","items":{"type":"object","required":["role","content"],"properties":{"role":{"enum":["system","user","assistant"],"type":"string"},"content":{"type":"string"}}}},"max_tokens":{"type":"integer","maximum":1024,"minimum":1},"temperature":{"type":"number","maximum":2,"minimum":0}}},"responseSchema":{"type":"json","example":{"ok":true,"model":"qwen2.5:7b","usage":{"total_tokens":44,"prompt_tokens":14,"completion_tokens":30},"object":"chat.completion","choices":[{"index":0,"message":{"role":"assistant","content":"<text>"},"finish_reason":"stop"}],"receipt":{"tx_hash":"0x<64hex>","amount_usd":0.005,"receipt_id":"g402-<16hex>"}}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.005","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.005/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.005","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.005","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_wfKQ6fIBem18NINt3vxa_","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.005","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Runs OpenAI-compatible chat completions on local GPU models via the x402 micropayment protocol, requiring no API key.","exampleAgentPrompt":"Using the local GPU model 'mistral-7b', send a user message asking 'What are the key differences between TCP and UDP?' with a max of 512 tokens and temperature 0.7, and return the assistant's reply.","exampleUseCases":[{"title":"Ad-hoc LLM inference without API key","prompt":"I need to ask an AI model to summarize a paragraph for me, but I don't have an OpenAI API key — can you call a compatible endpoint using a per-call USDC payment with model 'llama3-8b', max 256 tokens, and temperature 0.5?"},{"title":"Autonomous agent sub-task delegation","prompt":"My workflow agent needs to call an LLM to classify a customer support ticket into one of these categories: billing, technical, or general — use the local GPU model available, set temperature to 0.2 and max tokens to 128."},{"title":"Cost-controlled creative writing assistant","prompt":"Generate a short 3-sentence product description for a wireless ergonomic keyboard using a local GPU chat model, keep max tokens at 200 and temperature at 1.0 for some creativity."}],"resultDescription":"Returns an OpenAI-format chat completion object containing the assistant's generated message content, model used, finish reason, and token usage statistics (prompt tokens, completion tokens, total tokens).","failureModes":["Payment not fulfilled — x402 payment handshake fails, returning 402 Payment Required","Model name not found or unsupported — returns an error indicating unknown model","max_tokens exceeds allowed maximum of 1024 — request rejected with validation error","Malformed messages array (missing role or content) — returns 400 Bad Request","GPU resource unavailable or overloaded — may return 503 or timeout","Temperature value out of range [0, 2] — validation error"],"whenToPreferThis":"Choose this endpoint when you need OpenAI-compatible chat completions without a subscription or API key, want per-call USDC micropayment pricing, or are building autonomous agents that pay only for what they use. It is ideal for cost-sensitive, high-volume, or decentralized AI workflows where vendor lock-in or monthly billing is undesirable. Prefer it over OpenAI or Anthropic direct APIs when x402 payment support is already in your stack and you want local GPU inference without centralized rate limits.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-10-02T01:16:57.423Z","isFirstParty":false,"canonicalSlug":"openai-compatible-chat-completions-via-x402-unykorn-local-gpu-7dc1b3ac"}