{"uid":"cap_YJ3Re96bxt4YDpQ4PCNqa","slug":"agenttools-pdf-to-text-extractor-e0ea3a48","name":"AgentTools PDF to Text Extractor","description":"Fetch a PDF (by URL, SSRF-guarded) or accept it as base64 and extract all readable text with page count. Deterministic parsing — no language model, so nothing is invented. Text-based PDFs only; scanned images need OCR (not included). Use this when an agent needs to extract readable text from any PDF — agents can't read PDFs.","url":"https://agenttools-hub.vercel.app/api/v1/dev/pdf-to-text?utm_source=zero.xyz","method":"GET","headers":{},"bodySchema":{"type":"object","$schema":"https://json-schema.org/draft/2020-12/schema","required":["input"],"properties":{"input":{"type":"object","required":["type","method"],"properties":{"type":{"type":"string","const":"http"},"method":{"enum":["GET"],"type":"string"},"queryParams":{"type":"object","required":[],"properties":{"url":{"type":"string","description":"Or provide base64 instead"},"base64":{"type":"string","description":"Base64 content"},"maxChars":{"type":"number","description":"Max text length"}}}},"additionalProperties":false},"output":{"type":"object","required":["type"],"properties":{"type":{"type":"string"},"example":{"type":"object"}}}}},"responseSchema":null,"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.003","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.003/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.003","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.003","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_az6Ya1Vr3IYeFJa3iGJXB","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.003","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Fetches a PDF by URL or accepts it as base64 and extracts all readable text with page count using deterministic parsing (no LLM).","exampleAgentPrompt":"Can you extract all the text from this PDF for me? Here's the link: https://example.com/report.pdf — just pull out all the readable text.","exampleUseCases":[{"title":"Research paper content extraction","prompt":"I found this academic paper as a PDF at https://arxiv.org/pdf/2301.00000.pdf — can you pull out all the text so I can read and summarize it?"},{"title":"Contract text retrieval","prompt":"I need to read through a contract that's only available as a PDF. Here's the URL: https://legal.example.com/contract.pdf — can you extract all the text from it?"},{"title":"Base64 PDF parsing","prompt":"I have a PDF encoded in base64 that I received from an API. Can you convert it to readable text? Here's the base64 content: JVBERi0xLjQK..."}],"resultDescription":"Returns all extracted plain text from the PDF along with the total page count. Output is deterministic — no language model is involved, so only text actually present in the PDF is returned. Works on text-based PDFs only; scanned image PDFs require OCR which this endpoint does not provide.","failureModes":["URL is unreachable or returns non-PDF content","PDF is a scanned image (no embedded text layer) — returns empty or minimal text","SSRF guard blocks internal/private IP URLs","base64 input is malformed or not a valid PDF","maxChars limit truncates output unexpectedly","PDF is password-protected or encrypted"],"whenToPreferThis":"Use this endpoint when an agent needs to read the contents of a text-based PDF — either by URL or as base64 — and requires deterministic, faithful extraction with no hallucination risk. Prefer this over LLM-based PDF tools when accuracy and verbatim reproduction are critical. Not suitable for scanned/image PDFs that require OCR.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-10-02T12:43:14.724Z","isFirstParty":false,"canonicalSlug":"agenttools-pdf-to-text-extractor-e0ea3a48"}