{"uid":"cap_uoqjvGutAhTl0tQ_5UWQh","slug":"forgemesh-tts-pro-long-text-to-speech-501-2000-chars-bc68c998","name":"ForgeMesh TTS Pro Long – Text-to-Speech (501–2000 chars)","description":"Paid text-to-speech via x402. WAV audio in 31 languages, OpenAI-compatible, no API keys.","url":"https://tts.forgemesh.io/v1/tts/pro-long","method":"POST","headers":{},"bodySchema":{"type":"object","properties":{"lang":{"enum":["en","ko","ja","ar","bg","cs","da","de","el","es","et","fi","fr","hi","hr","hu","id","it","lt","lv","nl","pl","pt","ro","ru","sk","sl","sv","tr","uk","vi"],"type":"string","description":"Language code; 31 languages supported"},"text":{"type":"string","maxLength":2000,"minLength":1,"description":"Text to synthesize into speech, max 2000 characters. Use non-long routes for 1-500 chars and -long routes for 501-2000 chars."},"speed":{"type":"number","maximum":2,"minimum":0.7,"description":"Pro/Custom expressive speed control, 0.7-2.0. Presets: slow 0.7, normal 1.0, fast 1.3, rapid 1.6"},"steps":{"type":"integer","maximum":100,"minimum":1,"description":"Pro/Custom quality control, 1-100. Presets: draft 4, standard 8, high 16, ultra 24"},"voice":{"enum":["M1","M2","M3","M4","M5","F1","F2","F3","F4","F5"],"type":"string","description":"Standard voice: M1-M5 or F1-F5"}}},"responseSchema":{"type":"json","example":{"description":"44.1kHz 16-bit mono WAV speech audio returned inline","content_type":"audio/wav"}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.006","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.006/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.006","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.006","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_3akMdgay9hkY87SqZAnpJ","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.006","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Converts long-form text (up to 2000 characters) into 44.1kHz WAV audio across 31 languages with configurable voice, speed, and quality steps, paid via x402 micropayment.","exampleAgentPrompt":"Read this 800-character product description aloud in English using a female voice (F2), at normal speed and high quality (16 steps): 'Our new wireless headphones deliver 40 hours of battery life, active noise cancellation, and crystal-clear sound — available now in three colors.'","exampleUseCases":[{"title":"Multilingual e-learning narration","prompt":"Convert this 1200-character Spanish lesson summary into spoken audio using a male voice (M3) at normal speed (1.0) and standard quality (8 steps) — I need a WAV file to embed in our course module."},{"title":"Podcast intro script voiceover","prompt":"Generate a voiceover for my 900-character podcast intro in English with voice F1, speed set to 1.3 for a lively feel, and quality at 16 steps so it sounds polished."},{"title":"Automated customer announcement in Japanese","prompt":"Turn this 700-character Japanese store closing announcement into speech using voice M2, slow speed (0.7) so it's easy to understand, and draft quality (4 steps) for a quick review before we finalize it."}],"resultDescription":"Returns a binary WAV audio file (audio/wav content type) encoded at 44.1kHz, 16-bit, mono. The file contains the synthesized speech of the submitted text in the chosen language, voice, speed, and quality level, ready for playback or embedding.","failureModes":["Text exceeds 2000 character limit — use a shorter input or split the text","Unsupported language code — must be one of the 31 supported ISO codes","Invalid voice ID — must be M1–M5 or F1–F5","Speed out of range — must be between 0.7 and 2.0","Steps out of range — must be between 1 and 100","Payment failure via x402 — insufficient USDC balance or wallet misconfiguration","Empty or missing text field — minLength of 1 character required","Service timeout for very high step counts (e.g. steps=100) on long text"],"whenToPreferThis":"Choose this endpoint when you need to synthesize longer text (501–2000 characters) into high-quality WAV audio without managing API keys, when you need fine-grained control over voice gender, speed, and quality steps, or when you need multilingual support across 31 languages in a single service. Prefer this over shorter-text TTS routes when your input exceeds 500 characters. It is ideal for agents that pay per-call via x402 micropayments and need OpenAI-compatible audio generation without a subscription.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-14T17:51:15.195Z","isFirstParty":false}