{"uid":"cap_f3SSpkdSA0GKHEzubDSvi","slug":"aayat-ai-image-ocr-696d7d38","name":"Aayat AI Image OCR","description":"Read the text in an image (OCR): screenshots, photos of signs, receipts, slides, scanned pages. Returns the text in reading order with line breaks, and whether any text was found. JPEG, PNG, GIF or WebP up to 5 MB from a public URL.","url":"https://aayatai.com/image/ocr?utm_source=zero.xyz","method":"GET","headers":{},"bodySchema":{"type":"object","$schema":"https://json-schema.org/draft/2020-12/schema","required":["input"],"properties":{"input":{"type":"object","required":["type","method"],"properties":{"type":{"type":"string","const":"http"},"method":{"enum":["GET"],"type":"string"},"queryParams":{"type":"object","required":["url"],"properties":{"url":{"type":"string","format":"uri","maxLength":2048,"description":"Public image URL (JPEG, PNG, GIF or WebP, up to 5 MB)."}}}},"additionalProperties":false},"output":{"type":"object","required":["type"],"properties":{"type":{"type":"string"},"example":{"type":"object","required":["text","hasText","model"],"properties":{"url":{"type":"string"},"text":{"type":"string"},"model":{"type":"string"},"trust":{"type":"object","description":"Third-party text, cleaned: read trust.notice; removed = what we stripped."},"hasText":{"type":"boolean"}}}}}}},"responseSchema":null,"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.01","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.01/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.01","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.01","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_7eqluQEzlC3EBEg4LsrPl","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.01","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Extracts text from images (screenshots, photos, receipts, slides, scanned pages) via OCR and returns the text in reading order with a flag indicating whether any text was found","exampleAgentPrompt":"Can you read all the text from this image for me? Here's the public URL: https://example.com/receipt.jpg — I need it in reading order with line breaks.","exampleUseCases":[{"title":"Receipt text extraction for expenses","prompt":"I have a photo of a restaurant receipt at https://example.com/receipt.png — can you pull out all the text so I can log it as an expense?"},{"title":"Reading a screenshot of a document","prompt":"There's a screenshot of a legal notice at https://mycdn.com/notice.jpg — extract all the text from it so I can copy it into a document."},{"title":"Digitizing a sign or label photo","prompt":"I took a photo of a street sign and uploaded it to https://photos.example.com/sign.webp — what does it say?"}],"resultDescription":"Returns the full extracted text in reading order with line breaks preserved, a boolean 'hasText' flag indicating whether any text was detected, the source image URL, and the model name used for OCR.","failureModes":["Image URL is not publicly accessible — returns an error or empty result","Image format not supported (must be JPEG, PNG, GIF, or WebP) — request rejected","Image exceeds 5 MB size limit — request rejected","Image contains no readable text — returns hasText: false with empty text string","Network timeout fetching the remote image URL — error response","Malformed or invalid URL provided — validation error"],"whenToPreferThis":"Choose this endpoint when you need fast, pay-per-use OCR on a single public image URL without setting up your own vision model infrastructure. It is especially useful for AI agents handling receipts, screenshots, scanned documents, signs, or slides where text needs to be extracted programmatically. At $0.01 per call it is cost-effective for sporadic or on-demand extraction tasks within agentic workflows that already use the x402 payment protocol.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-10-02T01:11:42.233Z","isFirstParty":false,"canonicalSlug":"aayat-ai-image-ocr-696d7d38"}