{"uid":"cap_ASGJRs3zCpAqOD4n0fxTf","slug":"x402-readable-content-extractor-0aacfeef","name":"x402 Readable Content Extractor","description":"Fourteen pay-per-call web tools for AI agents, billed in USDC on Base with no account or API key. START WITH /audit -- one call returns a complete page health check (SEO + structured data + accessibility, optionally plus links and sitemap) with a single score and one worst-first findings list. The parts are also sold individually. SITE AUDIT: /seo (on-page SEO and meta tags), /links (broken links), /sitemap (robots.txt and XML sitemap crawlability), /structured (JSON-LD and social preview tags), /a11y (WCAG accessibility) -- each scored 0-100 with a prioritised issue list. READING: /read (clean markdown), /snap (screenshot, PDF or rendered HTML), /tables (every HTML table as structured JSON), /verify (does a claim actually appear on this page). MONITORING: /watch (has this page changed), /compare (what differs between two pages). Plus /rank (best x402 service for a task) and /batch (several operations behind one payment).","url":"https://x402-zadenworks.zaden-gaming.workers.dev/read?utm_source=zero.xyz","method":"POST","headers":{},"bodySchema":{"type":"object","$schema":"https://json-schema.org/draft/2020-12/schema","required":["input"],"properties":{"input":{"type":"object","required":["type","method","bodyType","body"],"properties":{"body":{"type":"object","required":["url"],"properties":{"url":{"type":"string","description":"Public http/https URL to read"},"waitMs":{"type":"integer","maximum":5000,"minimum":0},"selector":{"type":"string","description":"Optional CSS selector to scope extraction to"},"maxLength":{"type":"integer","maximum":200000,"minimum":1000},"includeLinks":{"type":"boolean"},"includeImages":{"type":"boolean"}}},"type":{"type":"string","const":"http"},"method":{"enum":["POST"],"type":"string"},"bodyType":{"enum":["json","form-data","text"],"type":"string"}},"additionalProperties":false},"output":{"type":"object","required":["type"],"properties":{"type":{"type":"string"},"example":{"type":"object","properties":{"url":{"type":"string"},"title":{"type":["string","null"]},"status":{"type":["integer","null"]},"finalUrl":{"type":"string"},"markdown":{"type":"string"},"truncated":{"type":"boolean"},"wordCount":{"type":"integer"}}}}}}},"responseSchema":{"type":"json","example":{"url":"https://example.com","title":"Example Domain","status":200,"finalUrl":"https://example.com/","markdown":"# Example Domain\n\nThis domain is for use in illustrative examples...","truncated":false,"wordCount":28}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.01","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"registry","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.01/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.01","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.01","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_ucW8sam1FxRVBtuo1JSU7","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.01","costPer":"request","priority":0,"asset":null,"unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Extracts clean, readable markdown or plain text from any URL, including JavaScript-rendered pages, stripping navigation, ads, and scripts while preserving title, byline, excerpt, and body content.","exampleAgentPrompt":"Can you pull the clean readable text from this article — https://example.com/some-long-news-article — and give me just the title, byline, and main body content without all the ads and navigation?","exampleUseCases":[{"title":"Research article summarization pipeline","prompt":"Fetch the clean readable content from https://techcrunch.com/2024/05/01/some-article and give me the title, author, and full article body as plain text so I can summarize it."},{"title":"Newsletter digest from multiple sources","prompt":"Extract the readable article text from https://www.theverge.com/some-story — I want just the headline, byline, and body with all the ads and menus stripped out."},{"title":"AI training data from JS-heavy pages","prompt":"Scrape the main readable content from https://medium.com/@author/some-post — it's a JavaScript-rendered page, so I need the actual article text in clean markdown, not the raw HTML."}],"resultDescription":"Returns structured content extracted from the target URL including: article title, author byline, short excerpt/summary, and the full readable body text in clean markdown or plain text format, with navigation menus, advertisements, scripts, and other non-content elements removed.","failureModes":["URL is inaccessible or returns a non-200 status — extraction fails with an error","Page has no readable main content (e.g. pure app shell) — returns empty or minimal body","JavaScript rendering times out for very slow pages — may return partial or no content","Paywalled or login-gated content returns only the preview text visible to unauthenticated users","Malformed or non-HTTP URLs cause immediate input validation errors"],"whenToPreferThis":"Choose this endpoint when you need clean, human-readable article or blog content from a URL, especially when the page is JavaScript-rendered and raw HTML scraping would yield noisy or incomplete results. It is ideal for summarization pipelines, research assistants, or any workflow where you want the meaningful text body without boilerplate. Prefer this over raw HTML fetchers or screenshot tools when structured text output (title, byline, body) is the goal rather than visual rendering or structured data like tables.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-10-02T04:55:58.480Z","isFirstParty":false,"canonicalSlug":"x402-readable-content-extractor-0aacfeef"}