{"uid":"cap_lf4fteEmFl2rQbfQWzRX3","slug":"qa-judge-e58d9cae","name":"QA Judge","description":"QA Judge: grade an agent output against your rubric (structure, citations, claim support, length constraints) and receive a 0-100 score with a signed attestation. $0.10 USDC per grading.","url":"https://attester.dev/judge","method":"POST","headers":{},"bodySchema":null,"responseSchema":null,"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.1","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.1/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.1","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.1","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_frdfDLjbX1J_0jdShdaDS","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.1","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Evaluates and judges the quality of question-answer pairs or AI outputs, returning a quality assessment score or verdict.","exampleAgentPrompt":"Can you judge the quality of this QA pair — the question is 'What is the boiling point of water?' and the answer given is 'Water boils at 100 degrees Celsius at sea level' — is this a good, accurate answer?","exampleUseCases":[{"title":"Filter low-quality chatbot responses","prompt":"Before we store these AI-generated customer support replies in our knowledge base, can you run a quality check on each one to make sure the answers are actually accurate and complete — flag any that don't pass so we can review them manually?"},{"title":"Benchmark LLM answer quality","prompt":"I'm comparing two different language models and I have 50 question-answer pairs from each — can you score all of them for quality so I can see which model is giving better, more accurate responses overall?"},{"title":"Validate AI tutor answer accuracy","prompt":"We have a set of student questions paired with answers generated by our AI tutoring system — can you judge whether each answer is correct and adequately addresses what the student asked, so we can catch any misleading explanations before students see them?"}],"resultDescription":"Returns a quality judgment or verdict for the submitted QA pair, likely including a score or pass/fail classification indicating whether the answer adequately and accurately addresses the question.","failureModes":["Missing or malformed question/answer input returns 400 error","Payment not included or insufficient results in 402 Payment Required","Service unavailable returns 5xx error","Ambiguous or very short inputs may return low-confidence judgments"],"whenToPreferThis":"Choose this endpoint when you need an automated, paid quality assessment of question-answer pairs or AI-generated responses — particularly useful for LLM evaluation pipelines, benchmarking, or filtering low-quality outputs before storing or presenting them.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-14T00:34:18.176Z","isFirstParty":false}