{"uid":"cap_Fc8Kx8mBvIPxPodT-w3ZX","slug":"diffbot-product-extraction-via-locus-x402-3a5da709","name":"Diffbot Product Extraction via Locus x402","description":"Web data extraction — articles, products, discussions, images, videos, and auto-detect.","url":"https://diffbot.x402.paywithlocus.com/diffbot/product","method":"POST","headers":{},"bodySchema":{"type":"object","properties":{"url":{"type":"string"},"fields":{"type":"string"},"timeout":{"type":"number"},"discussion":{"type":"boolean"}}},"responseSchema":{"type":"json","example":{"data":{},"payment":{"scheme":"exact","settledUsdc":"0.001000","authorizedMaxUsdc":"0.001000"},"request":{"id":"00000000-0000-4000-8000-000000000000","statusUrl":"/requests/00000000-0000-4000-8000-000000000000"},"success":true}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.0042","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.0042/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.0042","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.0042","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_-mb3xKtfmcJP4BfnLZqes","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.0042","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Extracts structured product data (price, title, description, images, specs) from any e-commerce product page URL using Diffbot's AI web extraction.","exampleAgentPrompt":"Can you pull the structured product data from this page — https://www.amazon.com/dp/B09XYZ — including the price, title, description, images, and any breadcrumb navigation?","exampleUseCases":[{"title":"Competitor price monitoring","prompt":"Grab the current price, availability, and product title from this competitor's product page: https://www.bestbuy.com/site/apple-macbook-pro/1234567.p — I want to track if they discount it."},{"title":"Product catalog enrichment","prompt":"I have a list of supplier product URLs I need structured — start with https://supplier-site.com/products/widget-pro and pull back the title, price, description, images, and any breadcrumb info so I can populate our internal catalog."},{"title":"Marketplace listing research","prompt":"Can you extract all the product details from this Etsy listing — https://www.etsy.com/listing/987654321/handmade-ceramic-mug — including price, description, and any available specs or meta fields?"}],"resultDescription":"Returns a JSON object with structured product data extracted from the target URL, including fields like product name, price, images, description, brand, availability, and optionally breadcrumbs, links, and meta — plus payment confirmation details (settled USDC amount) and a request status URL.","failureModes":["Target URL is paywalled or requires login — returns empty or partial data","URL is not a recognizable product page — extraction may return no structured fields","Timeout exceeded if the target page is slow — controlled via the timeout parameter","Anti-bot protections on target site may block Diffbot's crawlers","Malformed URL input causes request failure","Rate limiting or payment failure via x402 protocol returns a 402 error"],"whenToPreferThis":"Choose this endpoint when you need structured product data (price, title, images, specs) extracted from a specific e-commerce product page URL without building a custom scraper. It is ideal for price comparison, catalog enrichment, and competitive intelligence tasks. Prefer it over generic HTML-fetching approaches when you want clean, field-level JSON output rather than raw HTML. Use the article or auto-detect sibling endpoints if the target page is editorial content rather than a product listing.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-14T00:35:58.707Z","isFirstParty":false}