{"uid":"cap_xKuq1aK7X6y0YLN_cO4yl","slug":"diffbot-via-locus-x402-web-data-extraction-3b3c176c","name":"Diffbot via Locus x402 — Web Data Extraction","description":"Web data extraction — articles, products, discussions, images, videos, and auto-detect.","url":"https://diffbot.x402.paywithlocus.com/diffbot/job","method":"POST","headers":{},"bodySchema":{"type":"object","properties":{"url":{"type":"string"},"fields":{"type":"string"},"timeout":{"type":"number"},"discussion":{"type":"boolean"}}},"responseSchema":{"type":"json","example":{"data":{},"payment":{"scheme":"exact","settledUsdc":"0.001000","authorizedMaxUsdc":"0.001000"},"request":{"id":"00000000-0000-4000-8000-000000000000","statusUrl":"/requests/00000000-0000-4000-8000-000000000000"},"success":true}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.0042","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.0042/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.0042","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.0042","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_nYmr9szENnVRLMjtW61hV","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.0042","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Extracts structured data (articles, products, discussions, images, videos) from any URL using Diffbot's AI-powered web extraction, billed per call via x402/USDC micropayment.","exampleAgentPrompt":"Can you extract the full article text, author, and publication date from this URL — https://techcrunch.com/2024/01/15/some-article — including any comment threads on the page?","exampleUseCases":[{"title":"News article content extraction","prompt":"Pull the full article text, author name, publication date, and any embedded links from https://www.bbc.com/news/technology-12345678 — I need it as clean structured data."},{"title":"E-commerce product data scraping","prompt":"Extract the product name, price, description, and specs from this Amazon listing: https://www.amazon.com/dp/B09XYZ1234 — I want everything Diffbot can pull from the product page."},{"title":"Job posting details parser","prompt":"Grab the job title, required skills, salary range, and company description from this careers page: https://jobs.example.com/senior-engineer — include any extra fields like breadcrumbs if available."}],"resultDescription":"Returns a JSON object with a `data` field containing the extracted structured content (varies by page type: article text, product fields, discussion threads, etc.), plus payment confirmation details including settled USDC amount and a request ID with a status URL for tracking.","failureModes":["URL is unreachable or returns a non-200 HTTP status — extraction fails with error in response","Page is heavily JavaScript-rendered and Diffbot cannot parse content — empty or partial data returned","Timeout exceeded (default 30000ms) if page loads slowly — request may return incomplete data","URL points to a non-extractable file type (PDF, raw binary) — may return empty data object","Invalid or malformed URL input — API returns error response","Payment authorization failure via x402 — request not processed"],"whenToPreferThis":"Choose this endpoint when you need structured, AI-parsed web content from arbitrary URLs without maintaining your own scraping infrastructure. It is especially strong for extracting article text, product data, or discussion threads from complex, JavaScript-heavy pages. Prefer it over raw HTTP fetches or basic scrapers when you need clean structured JSON output and are comfortable with per-call USDC micropayment billing via x402.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-15T06:42:35.501Z","isFirstParty":false}