{"uid":"cap_UmmXC5CsVgPCFXnzap3nW","slug":"agentutility-text-to-speech-api-e34e5d08","name":"AgentUtility Text-to-Speech API","description":"Converts text to speech with 30+ voices and 5 audio formats. Morpheus primary for Kokoro, Venice fallback and alternate TTS models (xAI / ElevenLabs / Orpheus / MiniMax / Gemini), with fal.ai storage for hosted audio URLs. Use it as a TTS API or voice generator.","url":"https://x402.agentutility.ai/text-to-speech","method":"POST","headers":{},"bodySchema":null,"responseSchema":{"type":"json","example":{"format":"mp3","source":"morpheus","audio_url":"https://...mp3","file_size_bytes":24000}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.05","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"registry","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.05/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.05","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.05","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_TPdMxODrF0_VwT8dRzP42","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.05","costPer":"request","priority":0,"asset":null,"unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Converts input text to spoken audio using 30+ voices across multiple TTS providers (Kokoro, Venice, ElevenLabs, xAI, Orpheus, MiniMax, Gemini), returning a hosted audio URL via fal.ai storage.","exampleAgentPrompt":"Read this blog post intro aloud using a warm female voice and give me a hosted audio URL I can embed: 'Welcome to The Daily Digest — your source for the stories that matter most today.'","exampleUseCases":[{"title":"Podcast intro voiceover generation","prompt":"Generate a spoken intro for my podcast using a deep, confident male voice: 'You're listening to The Builder's Edge — conversations with founders who shipped before they were ready.' I need an mp3 I can drop straight into my editor."},{"title":"Accessible article narration","prompt":"Turn this article summary into audio so visually impaired readers can listen to it: 'Scientists have discovered a new deep-sea species off the coast of New Zealand, challenging existing models of ocean biodiversity.' Use a clear, neutral voice and return a hosted audio link."},{"title":"Agent reply voice output","prompt":"My AI customer support bot needs to respond to a user verbally — synthesize this message using a friendly female voice: 'Thanks for reaching out! Your order has shipped and should arrive by Thursday.' Give me the audio URL so I can play it back."}],"resultDescription":"Returns a hosted audio URL (via fal.ai storage) pointing to the synthesized speech file in the requested audio format. The agent can use this URL to stream, embed, or download the generated audio.","failureModes":["Primary Kokoro model unavailable — fallback to Venice or alternate provider may introduce latency or voice quality differences","Unsupported voice or format combination — returns error if requested voice/model/format pairing is invalid","Text too long — very large inputs may be rejected or truncated depending on model limits","Payment failure — x402 USDC payment not confirmed, request rejected with 402 status","fal.ai storage unavailable — audio generated but URL hosting may fail, causing incomplete response"],"whenToPreferThis":"Choose this endpoint when you need a pay-per-call TTS API that supports 30+ voices and multiple audio formats without a subscription, especially when you want a hosted audio URL returned directly (via fal.ai) rather than raw audio bytes. Ideal for AI agents that need voice output in agentic pipelines settled via USDC on Base. Prefer it over fixed-subscription TTS services when usage is sporadic or when you need access to multiple underlying TTS providers (Kokoro, ElevenLabs, Gemini, etc.) through a single endpoint.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-14T13:13:23.620Z","isFirstParty":false}