{"uid":"cap_bSpcxjnPXromLmpGjsLnX","slug":"agent402-tools-content-extraction-workflow-88f42985","name":"agent402.tools Content Extraction Workflow","description":"Bundled execution of the Content extraction workflow - Turn arbitrary URLs and PDFs into clean structured text - articles, page metadata, PDF pages, OCR'd images, browser-rendered SPAs. One x402 payment runs 6 underlying tools (extract, meta, pdf-to-markdown, pdf-extract-pages, render, image-ocr); partial-success per step.","url":"https://agent402.tools/api/skill/content-extraction","method":"POST","headers":{},"bodySchema":{"type":"object","properties":{"urls":{"type":"string","description":"Newline- or comma-separated list of URLs / PDF links to ingest"}}},"responseSchema":{"type":"json","example":{"args":{"urls":"https://agent402.tools/"},"pack":"content-extraction","steps":[{"ok":true,"slug":"extract","result":{}},{"ok":true,"slug":"meta","result":{}},{"ok":true,"slug":"pdf-to-markdown","result":{}},{"ok":true,"slug":"pdf-extract-pages","result":{}},{"ok":true,"slug":"render","result":{}},{"ok":true,"slug":"image-ocr","result":{}}],"summary":"6/6 steps succeeded"}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.044","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.044/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.044","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.044","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_Y8fuyi8Z5VvxnYt10eX2X","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.044","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Converts arbitrary URLs and PDFs into clean structured text, including articles, page metadata, PDF pages, OCR'd images, and browser-rendered SPAs in a single bundled call","exampleAgentPrompt":"Can you extract the full article text and metadata from these two URLs: https://example.com/article and https://example.com/report.pdf — I need clean structured text from both, including any PDF pages.","exampleUseCases":null,"resultDescription":"Returns clean structured text extracted from the provided URLs or PDFs, including article body, page metadata (title, description, OpenGraph), PDF page content, OCR'd image text, and rendered content from JavaScript-heavy SPAs — all from a single bundled x402 payment covering 6 underlying extraction tools.","failureModes":["URL is unreachable or returns non-200 status — extraction fails for that URL","PDF is password-protected or corrupted — PDF parsing returns error","SPA requires authentication to render — browser rendering may return gated content","Image OCR fails if image resolution is too low or format is unsupported","Malformed or empty URL input — returns validation error","Rate limits or network timeouts on the target site — partial or empty extraction"],"whenToPreferThis":"Use this endpoint when you need a comprehensive, single-call extraction pipeline covering web pages, PDFs, and SPAs without stitching together multiple tools. It is ideal when you need both article content and metadata from the same URLs, or when URLs may be a mix of HTML pages and PDF documents. Prefer this over individual scraping tools when you want one x402 payment to cover the full extraction workflow.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-15T06:54:59.463Z","isFirstParty":false}