{"uid":"cap_qw4VAfL4X7ZHEtq2YUyvB","slug":"forgemesh-tts-base-long-text-to-speech-501-2000-chars-36d240ef","name":"ForgeMesh TTS Base-Long (Text-to-Speech, 501–2000 chars)","description":"Paid text-to-speech via x402. WAV audio in 31 languages, OpenAI-compatible, no API keys.","url":"https://tts.forgemesh.io/v1/tts/base-long","method":"POST","headers":{},"bodySchema":{"type":"object","properties":{"lang":{"enum":["en","ko","ja","ar","bg","cs","da","de","el","es","et","fi","fr","hi","hr","hu","id","it","lt","lv","nl","pl","pt","ro","ru","sk","sl","sv","tr","uk","vi"],"type":"string","description":"Language code; 31 languages supported"},"text":{"type":"string","maxLength":2000,"minLength":1,"description":"Text to synthesize into speech, max 2000 characters. Use non-long routes for 1-500 chars and -long routes for 501-2000 chars."},"voice":{"enum":["M1","M2","M3","M4","M5","F1","F2","F3","F4","F5"],"type":"string","description":"Standard voice: M1-M5 or F1-F5"}}},"responseSchema":{"type":"json","example":{"description":"44.1kHz 16-bit mono WAV speech audio returned inline","content_type":"audio/wav"}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.003","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.003/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.003","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.003","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_lzSH-4v6QZaSao5_S3SAy","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.003","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Converts long-form text (up to 2000 characters) into 44.1kHz 16-bit mono WAV audio in 31 languages using a selectable male or female voice, paid per-call via x402 with no API key required.","exampleAgentPrompt":"Read this 3-paragraph blog introduction aloud in English using a female voice (F2): 'Welcome to our quarterly review. This year has seen remarkable growth across all divisions. We look forward to sharing detailed results with you today.'","exampleUseCases":[{"title":"Multilingual podcast intro narration","prompt":"Convert this 800-character Spanish intro script into speech using a male voice (M3) in Spanish — I need a WAV file to open my podcast episode: 'Bienvenidos al episodio de hoy. Hoy hablaremos sobre los avances más recientes en inteligencia artificial y su impacto en la industria.'"},{"title":"Accessible article audio version","prompt":"Turn this article summary into spoken audio in French with a female voice (F1) — the text is about 600 characters and I need a WAV file to embed on the page for accessibility."},{"title":"AI assistant voice response playback","prompt":"My chatbot just generated a 700-character reply in Japanese — can you synthesize it into speech using voice M2 in Japanese so I can play it back to the user as audio?"}],"resultDescription":"Returns a 44.1kHz 16-bit mono WAV audio file inline (Content-Type: audio/wav) containing the synthesized speech of the submitted text. The audio is ready for direct playback or embedding without additional decoding.","failureModes":["Text under 1 character or over 2000 characters returns a validation error (use non-long routes for texts under 501 chars)","Unsupported language code returns an enum validation error","Invalid voice ID (outside M1-M5, F1-F5) returns a validation error","Payment failure via x402 (insufficient USDC balance) results in a 402 Payment Required response","Network timeout if synthesis takes longer than the server deadline"],"whenToPreferThis":"Choose this endpoint when you need to synthesize longer text passages (501–2000 characters) into high-quality WAV audio without managing API keys or subscription accounts. It is ideal for agents that need multilingual voice output across 31 languages with selectable voice gender and style, and can handle micropayments via x402. Prefer the non-long sibling routes for shorter texts (1–500 chars) to avoid unnecessary cost. Prefer this over ElevenLabs or Google TTS when you need a key-free, pay-per-call model compatible with x402 payment flows.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-14T17:51:00.191Z","isFirstParty":false}