{"uid":"cap_kBcqbB6zOXrsmgpfg-H2k","slug":"spraay-url-content-extractor-1727d98f","name":"Spraay URL Content Extractor","description":"Extract clean content from URLs for RAG pipelines. Up to 5 URLs per request.","url":"https://gateway.spraay.app/api/v1/search/extract","method":"POST","headers":{},"bodySchema":{"type":"object","properties":{"urls":{"type":"array"}}},"responseSchema":null,"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.02","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.02/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.02","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.02","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_2FJCM2ZyGijMzhyEUm5Vg","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.02","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Extracts clean, readable content from up to 5 URLs in a single request, optimized for RAG pipelines and AI ingestion.","exampleAgentPrompt":"Extract clean, readable content from these three URLs so I can feed them into my RAG pipeline: https://example.com/article1, https://example.com/article2, and https://example.com/article3.","exampleUseCases":[{"title":"Research ingestion for AI knowledge base","prompt":"Pull the clean text content from these five blog posts so I can embed them into my vector database: https://blog.example.com/post1, https://blog.example.com/post2, https://blog.example.com/post3, https://blog.example.com/post4, https://blog.example.com/post5."},{"title":"On-demand document context for LLM","prompt":"Fetch and clean the content from https://docs.example.com/api-reference and https://docs.example.com/quickstart so I can pass them as context to my language model for answering user questions."},{"title":"Competitive research content scraping","prompt":"Extract the readable article text from these two competitor blog URLs — https://competitor.com/feature-announcement and https://competitor.com/pricing-update — so I can summarize and analyze them."}],"resultDescription":"Returns clean, extracted text content from each submitted URL, stripped of HTML markup, navigation, ads, and boilerplate — ready for embedding, indexing, or LLM context injection. Up to 5 URLs processed per request.","failureModes":["URL is unreachable or returns a non-200 status — extraction fails for that URL","Paywalled or login-gated pages return partial or empty content","JavaScript-rendered single-page apps may yield incomplete content if JS execution is not supported","Exceeding 5 URLs per request returns a validation error","Malformed or invalid URLs rejected at input validation","Rate limiting or payment failures result in 402 responses"],"whenToPreferThis":"Choose this endpoint when you need to convert raw web URLs into clean, LLM-ready text — especially for RAG pipeline ingestion, vector embedding, or providing fresh web context to a language model. It is optimized for batch extraction of up to 5 URLs per call, making it efficient for multi-source research tasks. Prefer this over generic scrapers when you need boilerplate-stripped, readable content rather than raw HTML.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-14T18:44:54.540Z","isFirstParty":false}