{"uid":"cap_8iy7Np1SeyPqCxC5RwSXU","slug":"delx-audio-transcription-c69bd144","name":"Delx Audio Transcription","description":"Delx is an independent AI agent lab building and studying continuity, recovery, verifiable work and governed exchange for autonomous systems.","url":"https://api.delx.ai/api/v1/x402/transcribe-audio","method":"POST","headers":{},"bodySchema":{"type":"object","properties":{"audio_base64":{"type":"string","maxLength":2800000,"minLength":1,"description":"One base64-encoded WAV up to 2 MiB decoded and 60 seconds. audio/wav data-URI prefixes are accepted. URLs, file paths, and model/version fields are rejected."}}},"responseSchema":{"type":"json","example":{"text":"hello agents","bytes":12,"model":"vaibhavs10/incredibly-fast-whisper","sha256":"2d6c7c2b3f3c8e1e7b0d3a5b6e7f8a9b1234567890abcdef1234567890abcdef","provider":"replicate","text_url":"https://api.delx.ai/api/v1/generated-transcripts/9c2e1b84f0a14d2e9c8b7a6543210fed.txt","media_type":"text/plain","input_sha256":"34a8815e9c57e321194d50f2ea8b64394019fb8038c9aefd7a9a306a5030fc31","generation_id":"9c2e1b84f0a14d2e9c8b7a6543210fed","sale_price_usdc":0.02,"gross_margin_usd":0.0164,"upstream_cost_usd":0.0036,"provider_prediction_id":"tr-example-prediction"}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.02","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.02/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.02","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.02","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_QhLmcgDwBHg-ex05o58B5","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.02","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Transcribes a single WAV audio clip (up to 60 seconds, 2 MiB) to text for $0.02 USDC per call, returning the transcript and a SHA-256 receipt with no API key or audio retention.","exampleAgentPrompt":"Can you transcribe this voice note for me? I'll give you the WAV file as base64 — it's about 30 seconds long and I need the text of what was said.","exampleUseCases":[{"title":"Voice note to text for field agents","prompt":"I recorded a 45-second voice note on my phone after the meeting — here's the base64 WAV — can you transcribe it so I can paste it into my notes?"},{"title":"Support call snippet transcription","prompt":"I have a 55-second WAV clip of a customer support call encoded in base64. Can you get me a transcript of it, including a receipt I can attach to the case?"},{"title":"Meeting action items extraction","prompt":"Transcribe this 30-second WAV snippet from our standup — it's base64 encoded — so I can pull out the action items mentioned."}],"resultDescription":"Returns the full text transcript of the spoken audio, along with a SHA-256 hash receipt that serves as a verifiable record of the transcription. No audio is retained after processing.","failureModes":["Audio exceeds 2 MiB decoded size — request rejected","Audio clip longer than 60 seconds — request rejected","Invalid or malformed base64 encoding — decoding error returned","Non-WAV audio format submitted — format rejection error","Payment of $0.02 USDC not completed — 402 Payment Required response","URL or file path submitted instead of base64 data — rejected per schema"],"whenToPreferThis":"Choose this endpoint when you need a quick, pay-per-use transcription of a short WAV clip (under 60 seconds) without setting up an API key or subscription. It is ideal for one-off voice notes, support call snippets, or meeting excerpts where you need a verifiable SHA-256 receipt and audio privacy (no retained audio). Prefer alternatives for longer audio, non-WAV formats, bulk batch transcription, or when you need speaker diarization or timestamps.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-15T07:14:54.201Z","isFirstParty":false}