{"uid":"cap_2Q_8kQ-WD5ZWjVhD4B_W6","slug":"openai-tts-1-text-to-speech-x402-cdca78c3","name":"OpenAI TTS-1 Text-to-Speech (x402)","description":"tts-1: OpenAI TTS 1 text to speech, paid per call in USDC. OpenAI list $15 per 1M characters, plus $0.0005 (Base) or $0.0005 (Solana) per call; up to 4096 characters; voices: alloy, ash, coral, echo, fable, onyx, nova, sage, shimmer. Returns the audio bytes (mp3 by default). Rates: https://openai.mm.family/x402/pricing","url":"https://openai.mm.family/x402/v1/models/tts-1/audio/speech?utm_source=zero.xyz","method":"POST","headers":{},"bodySchema":{"type":"object","properties":{"input":{"type":"string"},"model":{"enum":["tts-1"],"type":"string"},"speed":{"type":"number"},"voice":{"enum":["alloy","ash","coral","echo","fable","onyx","nova","sage","shimmer"],"type":"string"},"instructions":{"type":"string"},"response_format":{"enum":["mp3","opus","aac","flac","wav","pcm"],"type":"string"}}},"responseSchema":{"type":"json","example":{"body":"binary audio bytes (mp3 by default)","content_type":"audio/mpeg"}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.001","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.001/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.001","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.001","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_KYx8CMKg0ijIcVi30w5-Z","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.001","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Converts text to speech audio using OpenAI's TTS-1 model, billed per call in USDC via x402 micropayment protocol","exampleAgentPrompt":"Read this product description aloud using the nova voice and give me an MP3: 'Introducing our new smart home hub — seamless control for every device in your home.'","exampleUseCases":[{"title":"Podcast intro narration","prompt":"Turn this script into spoken audio using the shimmer voice in MP3 format: 'Welcome to Tech Horizons, the podcast where we explore the future of artificial intelligence and the humans shaping it.'"},{"title":"Accessibility audio for article","prompt":"Generate an audio version of this article summary using the alloy voice at 0.9 speed so I can add it to my blog post: 'Climate scientists have recorded the hottest average global temperature in recorded history this past July.'"},{"title":"Voice assistant response recording","prompt":"Synthesize this customer support response in the onyx voice and return it as a WAV file: 'Thank you for reaching out. Your order has been shipped and will arrive within 3 to 5 business days.'"}],"resultDescription":"Binary audio bytes (MP3 by default) containing synthesized speech of the input text. The response content-type is audio/mpeg for MP3, with other formats available (opus, aac, flac, wav, pcm) depending on the response_format parameter. Maximum input is 4096 characters per call.","failureModes":["Input text exceeds 4096 character limit — request will be rejected","Invalid voice enum value — must be one of alloy, ash, coral, echo, fable, onyx, nova, sage, shimmer","Invalid response_format — must be one of mp3, opus, aac, flac, wav, pcm","Payment failure via x402 micropayment protocol — USDC payment not accepted or insufficient","Model field not set to 'tts-1' — only this model is accepted at this endpoint","Empty input string — no audio generated"],"whenToPreferThis":"Choose this endpoint when you need standard-quality, cost-efficient AI speech synthesis billed per call via USDC micropayments (x402 protocol) without API key management. Prefer tts-1 over tts-1-hd when audio quality is acceptable at lower cost ($15/1M chars vs $30/1M chars). Best for agents needing pay-per-use audio generation without subscription overhead, or when integrating into crypto-native workflows on Base or Solana.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-10-02T12:30:04.388Z","isFirstParty":false,"canonicalSlug":"openai-tts-1-text-to-speech-x402-cdca78c3"}