{"uid":"cap_VXUejPU41V4yf2scm7Uyo","slug":"gpuops-ai-inference-proxy-text-to-speech-9d4ecaed","name":"GPUOps AI Inference Proxy – Text-to-Speech","description":"OpenAI-compatible AI inference API with 63 models. x402 pay-per-call with USDC on Base.","url":"https://ai.gpuops.io/v1/audio/speech","method":"POST","headers":{},"bodySchema":{"type":"object","properties":{"input":{"type":"string"},"model":{"type":"string"}}},"responseSchema":null,"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.01","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.01/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.01","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.01","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_QQLbvrQBCny8frm9-K8QK","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.01","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Converts input text to synthesized speech audio using a specified AI model, billed per call via USDC on Base.","exampleAgentPrompt":"Can you turn this script into spoken audio — 'Welcome to our platform. We're glad you're here.' — using the tts-1 model?","exampleUseCases":[{"title":"Podcast intro voice-over","prompt":"Generate a spoken audio clip from this intro text: 'Hello and welcome to The Daily Digest, your source for morning news.' Use the tts-1-hd model for high quality."},{"title":"Accessibility audio for article","prompt":"Convert this article summary into speech so visually impaired users can listen to it: 'Scientists have discovered a new species of deep-sea fish near the Mariana Trench.' Use tts-1 model."},{"title":"Voice assistant response audio","prompt":"Synthesize this chatbot reply as audio for playback in my voice app: 'Your order has been confirmed and will arrive by Thursday.' Use the tts-1 model."}],"resultDescription":"An audio file (speech) synthesized from the provided input text, generated by the selected model. The response is the raw audio stream or file suitable for playback or storage.","failureModes":["Invalid or unsupported model identifier returns an error","Empty or missing input text causes a 400 bad request","Insufficient USDC balance or x402 payment failure blocks the call","Network timeout if the model inference takes too long","Unsupported audio format or model capability mismatch"],"whenToPreferThis":"Choose this endpoint when you need pay-per-call, crypto-native (USDC on Base) text-to-speech inference with access to 63 OpenAI-compatible models, without committing to a subscription. Ideal for agents that need autonomous on-chain micropayments for speech generation.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-15T18:36:59.502Z","isFirstParty":false}