{"uid":"cap_n54Rp0rv19udWQ7JYhgi-","slug":"gpuops-ai-inference-proxy-64978e39","name":"GPUOps AI Inference Proxy","description":"OpenAI-compatible AI inference API with 63 models. x402 pay-per-call with USDC on Base.","url":"https://ai.gpuops.io/v1/chat/completions","method":"POST","headers":{},"bodySchema":{"type":"object","properties":{"model":{"type":"string"},"messages":{"type":"array"},"max_tokens":{"type":"integer"},"temperature":{"type":"number"}}},"responseSchema":null,"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.01","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.01/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.01","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.01","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm__omlYRcmLjnQoBbHoYMc4","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.01","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"OpenAI-compatible chat completions API supporting 63 models, billed per-call via USDC micropayments on Base","exampleAgentPrompt":"Using GPUOps, send this message to the mistral-7b-instruct model and get a reply: 'Summarize the key benefits of renewable energy in 3 bullet points' — allow up to 512 tokens and use a temperature of 0.7.","exampleUseCases":[{"title":"On-demand LLM for autonomous agent","prompt":"I need my agent to call an AI model on demand without a monthly subscription — use GPUOps to send 'Draft a polite follow-up email for a sales lead who went quiet after the demo' to llama-3-8b-instruct with max 300 tokens and temperature 0.5."},{"title":"Multi-model comparison pipeline","prompt":"Compare how two models answer the same question — use GPUOps to send 'What are the risks of deflation?' to both mixtral-8x7b-instruct and gpt-4o-mini with max 400 tokens each, temperature 0.3."},{"title":"Cost-controlled content generation","prompt":"Generate a product description for my new standing desk using GPUOps — send the prompt 'Write a 100-word punchy product description for an ergonomic bamboo standing desk' to mistral-7b-instruct, limit it to 200 tokens, and keep temperature at 0.8."}],"resultDescription":"Returns an OpenAI-compatible chat completion response object containing the assistant's generated message content, model used, finish reason, and token usage statistics (prompt tokens, completion tokens, total tokens).","failureModes":["Insufficient USDC balance causes payment failure with 402 response","Invalid or unsupported model name returns 400 or 404 error","Exceeding max_tokens limit may truncate output or return error","Malformed messages array structure causes 400 bad request","Network timeout if model inference takes too long for complex prompts","Rate limiting if too many concurrent requests are sent"],"whenToPreferThis":"Choose this endpoint when you need pay-per-call LLM inference without a subscription commitment, especially in agentic or automated pipelines that already handle USDC/Base payments. Ideal when you want access to a broad selection of 63 models through a single OpenAI-compatible interface, or when your application needs to pay for inference programmatically using x402 crypto micropayments rather than managing API keys and billing accounts.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-14T06:33:45.418Z","isFirstParty":false}