{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2025:AZBDYHXMEMVZYGVIGOA5BBPLRC","short_pith_number":"pith:AZBDYHXM","schema_version":"1.0","canonical_sha256":"06423c1eec232b9c1aa83381d085eb88a5a1ec8bedefe451c4110c4a829edd6c","source":{"kind":"arxiv","id":"2505.13195","version":1},"attestation_state":"computed","paper":{"title":"Adversarial Testing in LLMs: Insights into Decision-Making Vulnerabilities","license":"http://creativecommons.org/licenses/by-nc-sa/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.AI","authors_text":"Haomiaomiao Wang, Libao Deng, Lili Zhang, Long Cheng, Tomas Ward","submitted_at":"2025-05-19T14:50:44Z","abstract_excerpt":"As Large Language Models (LLMs) become increasingly integrated into real-world decision-making systems, understanding their behavioural vulnerabilities remains a critical challenge for AI safety and alignment. While existing evaluation metrics focus primarily on reasoning accuracy or factual correctness, they often overlook whether LLMs are robust to adversarial manipulation or capable of using adaptive strategy in dynamic environments. This paper introduces an adversarial evaluation framework designed to systematically stress-test the decision-making processes of LLMs under interactive and ad"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2505.13195","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by-nc-sa/4.0/","primary_cat":"cs.AI","submitted_at":"2025-05-19T14:50:44Z","cross_cats_sorted":[],"title_canon_sha256":"b842b0821eca5eb807904abf06b8b44520095addd9a527bfee41e87aa0d2d50f","abstract_canon_sha256":"82506cdc18d26e5b3376b5ad1234b60449a4290cc2f01feb789e410d381f813d"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T11:05:26.493739Z","signature_b64":"ONxNIeJ8Jg4opvvtiyRPf/ePuBWF4ARdlUZoIA39D+7qgHuuWz3W7FSXEyaXRXbYSbzo7dCKO1q9GRde/GmVAg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"06423c1eec232b9c1aa83381d085eb88a5a1ec8bedefe451c4110c4a829edd6c","last_reissued_at":"2026-07-05T11:05:26.493158Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T11:05:26.493158Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Adversarial Testing in LLMs: Insights into Decision-Making Vulnerabilities","license":"http://creativecommons.org/licenses/by-nc-sa/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.AI","authors_text":"Haomiaomiao Wang, Libao Deng, Lili Zhang, Long Cheng, Tomas Ward","submitted_at":"2025-05-19T14:50:44Z","abstract_excerpt":"As Large Language Models (LLMs) become increasingly integrated into real-world decision-making systems, understanding their behavioural vulnerabilities remains a critical challenge for AI safety and alignment. While existing evaluation metrics focus primarily on reasoning accuracy or factual correctness, they often overlook whether LLMs are robust to adversarial manipulation or capable of using adaptive strategy in dynamic environments. This paper introduces an adversarial evaluation framework designed to systematically stress-test the decision-making processes of LLMs under interactive and ad"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2505.13195","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2505.13195/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2505.13195","created_at":"2026-07-05T11:05:26.493241+00:00"},{"alias_kind":"arxiv_version","alias_value":"2505.13195v1","created_at":"2026-07-05T11:05:26.493241+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2505.13195","created_at":"2026-07-05T11:05:26.493241+00:00"},{"alias_kind":"pith_short_12","alias_value":"AZBDYHXMEMVZ","created_at":"2026-07-05T11:05:26.493241+00:00"},{"alias_kind":"pith_short_16","alias_value":"AZBDYHXMEMVZYGVI","created_at":"2026-07-05T11:05:26.493241+00:00"},{"alias_kind":"pith_short_8","alias_value":"AZBDYHXM","created_at":"2026-07-05T11:05:26.493241+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":1,"sample":[{"citing_arxiv_id":"2608.03166","citing_title":"Adversarial Stress Testing of Role-Playing Language Agents using Multi-Agent Evaluation","ref_index":7,"is_internal_anchor":true}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/AZBDYHXMEMVZYGVIGOA5BBPLRC","json":"https://pith.science/pith/AZBDYHXMEMVZYGVIGOA5BBPLRC.json","graph_json":"https://pith.science/api/pith-number/AZBDYHXMEMVZYGVIGOA5BBPLRC/graph.json","events_json":"https://pith.science/api/pith-number/AZBDYHXMEMVZYGVIGOA5BBPLRC/events.json","paper":"https://pith.science/paper/AZBDYHXM"},"agent_actions":{"view_html":"https://pith.science/pith/AZBDYHXMEMVZYGVIGOA5BBPLRC","download_json":"https://pith.science/pith/AZBDYHXMEMVZYGVIGOA5BBPLRC.json","view_paper":"https://pith.science/paper/AZBDYHXM","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2505.13195&json=true","fetch_graph":"https://pith.science/api/pith-number/AZBDYHXMEMVZYGVIGOA5BBPLRC/graph.json","fetch_events":"https://pith.science/api/pith-number/AZBDYHXMEMVZYGVIGOA5BBPLRC/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/AZBDYHXMEMVZYGVIGOA5BBPLRC/action/timestamp_anchor","attest_storage":"https://pith.science/pith/AZBDYHXMEMVZYGVIGOA5BBPLRC/action/storage_attestation","attest_author":"https://pith.science/pith/AZBDYHXMEMVZYGVIGOA5BBPLRC/action/author_attestation","sign_citation":"https://pith.science/pith/AZBDYHXMEMVZYGVIGOA5BBPLRC/action/citation_signature","submit_replication":"https://pith.science/pith/AZBDYHXMEMVZYGVIGOA5BBPLRC/action/replication_record"}},"created_at":"2026-07-05T11:05:26.493241+00:00","updated_at":"2026-07-05T11:05:26.493241+00:00"}