{"uid":"cap_cQZ6xue7gO93jXz_i9U-L","slug":"agentprobe-v1-0-0-bilingual-llm-agent-adversarial-test-pack-grader-260a103c","name":"AgentProbe v1.0.0 – Bilingual LLM-Agent Adversarial Test Pack Grader","description":"53 hand-authored, OWASP-mapped adversarial cases (EN+ZH) plus a zero-dependency MIT harness that scores any OpenAI-compatible LLM agent. US$6.99.","url":"https://agentprobe.pythonanywhere.com/v1/grade","method":"POST","headers":{},"bodySchema":null,"responseSchema":{"type":"json","example":{"pack":"zh-agent-adversarial","count":2}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.003","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.003/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.003","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.003","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm__EOjiWv60i0nYYAn61evj","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.003","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Scores an LLM agent against 53 hand-authored, OWASP-mapped adversarial test cases in English and Chinese, returning a graded result JSON.","exampleAgentPrompt":"Run the Chinese adversarial test pack from AgentProbe against my agent and tell me how many cases it passed — I want to see the OWASP-mapped grade results.","exampleUseCases":[{"title":"Red-team safety audit of production agent","prompt":"Run the full English adversarial test pack against my customer support agent and give me a grade — I need to know if it's vulnerable to any OWASP-mapped attack categories before we go live."},{"title":"Bilingual compliance check for multilingual bot","prompt":"My chatbot serves both English and Chinese users — can you run both the EN and ZH adversarial packs from AgentProbe and show me the scores so I know where it fails?"},{"title":"CI pipeline security gate for LLM agent","prompt":"Before I merge this PR, run the English adversarial test pack on my updated agent and tell me the case count and grade — fail the build if it doesn't pass."}],"resultDescription":"A JSON object containing the test pack name (e.g. 'zh-agent-adversarial') and the count of adversarial cases evaluated, along with grading results indicating how the agent performed against each OWASP-mapped adversarial scenario.","failureModes":["Invalid or unrecognized pack name returns an error or empty result","Malformed POST body causes a 400 or parsing error","Agent endpoint being tested is unreachable, causing timeout or connection failure","Payment not processed correctly via x402, resulting in 402 Payment Required","Rate limiting or server overload on pythonanywhere hosting returns 5xx"],"whenToPreferThis":"Choose this endpoint when you need to evaluate an OpenAI-compatible LLM agent against a curated, OWASP-mapped adversarial test suite in English and/or Chinese without writing your own test harness. It is especially useful for red-teaming, safety auditing, or CI/CD gating of agents where bilingual (EN+ZH) adversarial coverage matters and you want structured, reproducible scores per call.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-14T12:43:10.911Z","isFirstParty":false}