{"uid":"cap_JLM4JsS8FisxtV7bDiChS","slug":"dt0ur-online-text-to-speech-speecht5-20065c29","name":"dt0ur.online Text-to-Speech (SpeechT5)","description":"Synthesizes natural human speech from text prompts using SpeechT5 neural voice generation models with multi-speaker support.","url":"https://dt0ur.online/api/inference/tts?utm_source=zero.xyz","method":"POST","headers":{},"bodySchema":null,"responseSchema":null,"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.15","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"down","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.15/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.15","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.15","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_V7RAUc-JAUMsJ_-SW645-","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.15","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Synthesizes natural-sounding human speech audio from a text string using SpeechT5 neural TTS models with selectable speaker voice profiles.","exampleAgentPrompt":"Can you synthesize this paragraph into spoken audio using the default voice: 'Welcome to our platform. We're glad you're here and hope you enjoy the experience.'","exampleUseCases":[{"title":"Podcast intro narration","prompt":"Turn this script into natural-sounding speech using the default voice: 'Hello and welcome to The Daily Digest — your five-minute briefing on everything that matters today.'"},{"title":"Accessibility audio for written article","prompt":"Can you read aloud this article summary for me? Use the default speaker voice: 'Researchers have discovered a new method for storing solar energy that could reduce costs by up to 40 percent.'"},{"title":"Multilingual e-learning voiceover","prompt":"I need a spoken audio version of this lesson text for my e-learning course — synthesize it with the default voice: 'In this module, you will learn the fundamentals of neural networks and how they power modern AI systems.'"}],"resultDescription":"Returns a synthesized audio file or audio stream representing the input text spoken aloud in the selected voice profile, generated by SpeechT5 neural TTS models. The audio is human-like in cadence and intonation.","failureModes":["Empty or missing text field returns a validation error","Unsupported or unknown voice profile identifier may fall back to default or return an error","Very long text inputs may exceed processing limits or increase latency","Payment not included or insufficient USDC causes 402 Payment Required response","Service unavailability results in 5xx error with no audio output"],"whenToPreferThis":"Choose this endpoint when you need on-demand neural speech synthesis from text with multi-speaker support and pay-per-call pricing via x402. It is well-suited for agents that need to generate spoken audio dynamically without a long-term TTS subscription, especially when integrated into agentic pipelines that already use x402 micropayments.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-10-03T00:28:39.170Z","isFirstParty":false,"canonicalSlug":"dt0ur-online-text-to-speech-speecht5-20065c29"}