{"uid":"cap_-eOBMTokbhDLnvU9UkvZ9","slug":"content-hugen-tokyo-474fdd53","name":"content.hugen.tokyo Web Content Extractor","description":"Extract clean, readable text from any web page — removes navigation, ads, scripts, and boilerplate. Dual-engine extraction with JS-rendered page fallback for SPAs and dynamic content. Returns content with title, author, word count, language, and reading time. No headless browser or scraping library needed. Accepts USDC payments on Base and Solana","url":"https://content.hugen.tokyo/content/extract","method":"GET","headers":{},"bodySchema":{"properties":{"input":{"required":["method"]}}},"responseSchema":null,"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.01","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"settled","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.01/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.01","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.01","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_49FoOLnNrUOtecRKRA7mu","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.01","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Extracts clean, readable text from any web page, stripping ads, navigation, scripts, and boilerplate, with JS-rendered fallback for SPAs and dynamic sites.","exampleAgentPrompt":"Pull the clean readable text from this article — https://techcrunch.com/2024/05/01/ai-agents-are-here/ — and tell me the title, author, word count, and estimated reading time.","exampleUseCases":null,"resultDescription":"Returns the main article body as clean readable text, along with the page title, author name, word count, detected language, and estimated reading time. Ads, navigation menus, scripts, and boilerplate HTML are stripped. Handles both static HTML and JavaScript-rendered pages via dual-engine fallback.","failureModes":["URL is behind a login wall or paywall — extraction fails or returns partial content","Page uses aggressive bot-detection (Cloudflare, reCAPTCHA) — may return empty or blocked response","Invalid or malformed URL — returns 4xx error","Page has no meaningful text content (e.g. image-only or video page) — returns empty body","Network timeout on slow or unreachable URLs — returns timeout error","JS-rendered fallback may add latency if primary extraction fails"],"whenToPreferThis":"Choose this endpoint when you need clean article text from a URL without setting up a headless browser or scraping library, especially for JavaScript-heavy SPAs or dynamic content pages. Prefer it over raw HTTP fetch when you want boilerplate stripped and metadata like author, word count, and reading time included automatically.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-15T16:33:58.712Z","isFirstParty":false}