{"as_of":"2026-08-24T01:09:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:8cbbf13a7e1ca7d95350a79f2d26be97b89950034be0a705fd81d7005c201bcd","coverage":[{"denominator":20,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":20,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-30T14:38:43.708578Z","state":"measured"},{"denominator":22,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":22,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-23T06:30:58.430688+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-02T02:38:46.490717Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"pith","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"pith"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.24134","snapshot_observed_at":"2026-08-02T02:38:46.490717Z","title":"Proofagent harness: Open infrastructure for adversarial evaluation of ai agents, 2026","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.14275","last_updated":"2026-07-15T18:33:02Z","snapshot_observed_at":"2026-08-16T07:30:28.246281Z","submitted_at":"2026-07-15T18:33:02Z","title":"AI Agents Do Not Fail Alone:The Context Fails First","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-02T02:38:46.490717Z"},"links":{"cited_paper":"/paper/2605.24134","citing_paper":"/paper/2607.14275"},"observation_digest":"sha256:70e92244e31666ddc3969e1b3c786a3021a35bdcf35939676ce76457abd69e32","observation_id":"fac75aa7-8f2a-4ae8-baa5-61d393044d77","resolution":{"observed_at":"2026-08-02T02:38:46.490717Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"cited_work":{"arxiv_id":"2605.24134","doi":"10.48550/arxiv.2605.24134","metadata_source":"pith","pith_arxiv_id":"2605.24134","snapshot_observed_at":"2026-08-05T02:49:54.815029Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","venue":"cs.MA","work_id":"295564ef-e302-47ea-b85f-a782bc33708f","year":2026},"citing_paper":{"arxiv_id":"2607.27677","last_updated":"2026-07-30T04:44:47Z","snapshot_observed_at":"2026-08-17T14:04:40.356913Z","submitted_at":"2026-07-30T04:44:47Z","title":"Stop Shipping AI Agents on Faith: Capability Is Not Production Readiness","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-01T03:31:34.165368Z"},"links":{"cited_paper":"/paper/2605.24134","citing_paper":"/paper/2607.27677"},"observation_digest":"sha256:de1f5f396f65d285315f248db05012648725bab7820d43aa43d1656f95b51627","observation_id":"cb2b457c-a9b9-421f-830d-a87abf592a58","resolution":{"observed_at":"2026-08-01T03:34:07.269196Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2605.24134/citation-record","integrity":"/paper/2605.24134/integrity","json":"/paper/2605.24134/citation-record.json","paper":"/paper/2605.24134"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2501.00881","last_updated":"2025-01-01T16:00:18Z","snapshot_observed_at":"2026-08-21T20:06:22.075873Z","submitted_at":"2025-01-01T16:00:18Z","title":"Agentic Systems: A Guide to Transforming Industries with Vertical AI Agents","version":1},"cited_work":{"arxiv_id":"2501.00881","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.00881","snapshot_observed_at":"2026-06-30T14:44:45.263422Z","title":"Agentic systems: A guide to transforming indus- tries with vertical ai agents.arXiv preprint arXiv:2501.00881, 2025","venue":null,"work_id":"90526c7b-cb28-4db5-b301-559d57095bcc","year":2025},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"cited_paper":"/paper/2501.00881","citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:e4c464a186b1f239159320bb5e04872be04a4d89a71a7690589cc8027811f20d","observation_id":"bb3c0a5c-6621-4924-bf36-fc53ab83b744","resolution":{"observed_at":"2026-06-30T14:44:45.266937Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.08944","last_updated":"2025-01-15T16:46:32Z","snapshot_observed_at":"2026-08-10T20:11:43.649680Z","submitted_at":"2025-01-15T16:46:32Z","title":"Physical AI Agents: Integrating Cognitive Intelligence with Real-World Action","version":1},"cited_work":{"arxiv_id":"2501.08944","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.08944","snapshot_observed_at":"2026-06-30T14:44:45.276821Z","title":"Physical ai agents: Integrating cognitive intelligence with real-world action","venue":null,"work_id":"e84609a2-b532-4086-ae12-703edb218062","year":2025},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"cited_paper":"/paper/2501.08944","citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:fd84ab50ed5b75e3ceeacce6180fe454f692279d95080e7741c0bae39dc43643","observation_id":"50dabbfc-72e2-48ed-863d-394171b495fc","resolution":{"observed_at":"2026-06-30T14:44:45.278771Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2601.11653","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T14:44:45.283811Z","title":"Ai agents need memory control over more context","venue":null,"work_id":"4bec663a-1fd9-4335-9c5a-b5261fec0cec","year":2026},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:5db6acde0c6ade2ef43a86a664bc0416588fddf582cb88dabf2518a6e6afcf64","observation_id":"01783b3b-c78f-43ca-b62e-d575367fa748","resolution":{"observed_at":"2026-06-30T14:44:45.285472Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T22:25:39.974465Z","title":"Human oversight in the eu artificial intelligence act","venue":null,"work_id":"2271f605-844e-4acb-8860-9a514164c2ea","year":2023},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:dd97989736eee652e0c3a6407e0d81ddc08a1edb3bd3a8157b25e889524d38ef","observation_id":"a6ac0793-c2ec-4c26-b002-b85dc545e483","resolution":{"observed_at":"2026-07-08T22:25:39.976106Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T22:25:39.959621Z","title":"Regulation (eu) 2024/1689: Artificial intelligence act, article 14 human oversight,","venue":null,"work_id":"8c70ec30-53ef-4483-b89d-c48d6b3e7b67","year":2024},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:afd3c7dd61b7511c4961755697c6bb6b3777188d360297cac7c08d1ccb234f20","observation_id":"328a27f3-e45e-4e13-8972-aa4ee9a058a3","resolution":{"observed_at":"2026-07-08T22:25:39.961252Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T22:25:39.962018Z","title":null,"venue":null,"work_id":"a71dc2aa-91f6-4e06-98bc-eaaad4b15c73","year":null},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:315df2ae965c37b33ec70e5e076ff305c4626ba756a3a93bda1a2c48e140a24f","observation_id":"c7a16a38-f02a-40c9-9661-d78c4635c1b8","resolution":{"observed_at":"2026-07-08T22:25:39.963678Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2209.07858","last_updated":"2022-11-22T19:12:57Z","snapshot_observed_at":"2026-08-21T13:43:54.063053Z","submitted_at":"2022-08-23T23:37:14Z","title":"Red Teaming Language Models to Reduce Harms: Methods, Scaling Behaviors, and Lessons Learned","version":2},"cited_work":{"arxiv_id":"2209.07858","doi":"10.1136/bcr-2013-201554","metadata_source":"pith","pith_arxiv_id":"2209.07858","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Red Teaming Language Models to Reduce Harms: Methods, Scaling Behaviors, and Lessons Learned","venue":"cs.CL","work_id":"1aabd84d-3779-4ba9-ba2f-15ce264a9b1e","year":2022},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"cited_paper":"/paper/2209.07858","citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:f9b44eba89f63673ca17c31cd7e1a7e57a8a96daba2b9295556b51ff26e8c40b","observation_id":"f8543b32-dc10-4cb0-968f-c667e5f46b08","resolution":{"observed_at":"2026-06-30T14:44:45.287996Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T22:25:39.964598Z","title":"Ethics guidelines for trustworthy ai, 2019","venue":null,"work_id":"292a8c72-f4ba-439a-bcb6-3dc5f986af99","year":2019},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:8ec11029f54c675c376ae6f7ad547828dcd774d0fc21da8cbde66e1b05bd6396","observation_id":"8ea776ba-8b8c-44d6-977e-9728b13b432b","resolution":{"observed_at":"2026-07-08T22:25:39.966339Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T22:25:39.954187Z","title":"Jimenez, John Yang, Alexander Wettig, Shunyu Yao, Kexin Pei, Ofir Press, and Karthik Narasimhan","venue":null,"work_id":"fa1e111f-1568-4159-abfe-1d8ae6be1c8d","year":2024},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:91a77e0dec019a50a0789d0bd6996dc8789fa1a90efc22c16bb4ecf892a58177","observation_id":"7b62f927-09a9-4dcb-8592-3f4af9e2edc2","resolution":{"observed_at":"2026-07-08T22:25:39.956195Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.03688","last_updated":"2025-10-04T03:54:18Z","snapshot_observed_at":"2026-08-20T10:21:12.032735Z","submitted_at":"2023-08-07T16:08:11Z","title":"AgentBench: Evaluating LLMs as Agents","version":3},"cited_work":{"arxiv_id":"2308.03688","doi":"10.1109/fllm63129.2024.10852426","metadata_source":"pith","pith_arxiv_id":"2308.03688","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"AgentBench: Evaluating LLMs as Agents","venue":"cs.AI","work_id":"a37549b4-4c94-412d-acc4-4efeb08509be","year":2023},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"cited_paper":"/paper/2308.03688","citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:0d4085986b3b548d843bc7ada576e9e29c78ec1855d2993448776c0930bd170a","observation_id":"03253d2e-b22e-47e7-b2fa-dfaf55264a55","resolution":{"observed_at":"2026-06-30T14:44:45.281931Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T22:25:39.976752Z","title":"G-eval: Nlg evaluation using gpt-4 with better human alignment","venue":null,"work_id":"1d126d4c-f424-4a4f-b5f7-b2d8096ebac7","year":2023},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:029a3807edbd049214007af4654f543919de8bce2fdc4a3a06becb4bb9166d7b","observation_id":"e122c7e9-9e48-4c2d-b6ef-7e03a0502408","resolution":{"observed_at":"2026-07-08T22:25:39.978219Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T22:47:43.036334Z","title":"Red teaming language models with language models","venue":null,"work_id":"80ef9cc0-1d94-4685-acc0-247155457835","year":2022},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:18133a33788dfdb93c1c3d1de44975d1cf709d0f98fd6fdd6fc5e3d1e5c756f9","observation_id":"ac35bca0-09d9-403c-895f-dc402dbd2fe4","resolution":{"observed_at":"2026-07-08T22:25:39.968973Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.16789","last_updated":"2023-10-03T14:45:48Z","snapshot_observed_at":"2026-08-16T04:18:30.717622Z","submitted_at":"2023-07-31T15:56:53Z","title":"ToolLLM: Facilitating Large Language Models to Master 16000+ Real-world APIs","version":2},"cited_work":{"arxiv_id":"2307.16789","doi":"10.48550/arxiv.2307.16789","metadata_source":"pith","pith_arxiv_id":"2307.16789","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"ToolLLM: Facilitating Large Language Models to Master 16000+ Real-world APIs","venue":"cs.AI","work_id":"3c555b48-a4d9-42dd-9fdd-0f6018fbe9cb","year":2023},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"cited_paper":"/paper/2307.16789","citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:545519d000d07806158d7c53ab1a7dbaa5cea9ef099d000d831c6679392b4f40","observation_id":"149bae42-7880-426f-9031-84d3040bbd04","resolution":{"observed_at":"2026-06-30T14:44:45.275213Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-07-12T03:19:33.730697+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T03:19:33.730697+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T22:25:39.952071Z","title":"Jailbroken: How does llm safety training fail?Advances in Neural Information Processing Systems, 2023","venue":null,"work_id":"895f1643-6031-432a-a0ad-44dc4adbaad9","year":2023},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:922a1ee0079ad7776e527e4ede33951ad821a3feea836a6dcfc1fa3a2ed63a56","observation_id":"990183fa-e47a-42b5-9f70-e8243d13016a","resolution":{"observed_at":"2026-07-08T22:25:39.953407Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.08155","last_updated":"2023-10-03T20:47:10Z","snapshot_observed_at":"2026-08-08T22:28:26.004138Z","submitted_at":"2023-08-16T05:57:52Z","title":"AutoGen: Enabling Next-Gen LLM Applications via Multi-Agent Conversation","version":2},"cited_work":{"arxiv_id":"2308.08155","doi":"10.48550/arxiv.2308.08155","metadata_source":"pith","pith_arxiv_id":"2308.08155","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"AutoGen: Enabling Next-Gen LLM Applications via Multi-Agent Conversation","venue":"cs.AI","work_id":"92b7eb9c-c3d8-4518-a376-06fa15dd895b","year":2023},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"cited_paper":"/paper/2308.08155","citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:53fec01c38a3b097a4c3831ac820667e0abad153f884909aba134ea3bb379a0b","observation_id":"1c2610c4-fae3-4db9-af7b-6f477a3a11e4","resolution":{"observed_at":"2026-06-30T14:44:45.272643Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-07-14T18:20:20.607676+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-14T18:20:20.607676+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T22:25:39.959040Z","title":"Web- shop: Towards scalable real-world web interaction with grounded lan- guage agents","venue":null,"work_id":"04bcce87-46e5-44f5-88cf-5b9eeec1469a","year":2022},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:1f4b9dea323db0540a2fa232e855033b90db06247250b7fa504adeb94d0d1281","observation_id":"6ba8dd86-e0ff-44ee-ae31-9ba0b34c9988","resolution":{"observed_at":"2026-07-08T22:25:39.960419Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T22:25:39.969895Z","title":"React: Synergizing reasoning and acting in language models","venue":null,"work_id":"04de70d2-36bb-4ae8-ae0f-334a0f463959","year":2023},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:7b3e0070b5778a7f37eb8939aaae53dd9bb99736f15f92d532fe77ff87ea71d9","observation_id":"d9751ca0-70fd-4ef2-9451-2f6f5f1fea40","resolution":{"observed_at":"2026-07-08T22:25:39.971452Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T22:25:39.972206Z","title":"Xing, 38 Hao Zhang, Joseph E","venue":null,"work_id":"c79a7f32-50df-43b0-89f7-99df64bd47b8","year":2023},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:1066f2308360f7a1373c862c9dbb592f905a3bf94d91dc1c2384f1ad71696d38","observation_id":"67cf3cf0-a356-41bd-9643-755e0726ad0e","resolution":{"observed_at":"2026-07-08T22:25:39.973730Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.13854","last_updated":"2024-04-16T15:13:18Z","snapshot_observed_at":"2026-08-14T11:14:55.351653Z","submitted_at":"2023-07-25T22:59:32Z","title":"WebArena: A Realistic Web Environment for Building Autonomous Agents","version":4},"cited_work":{"arxiv_id":"2307.13854","doi":"10.48550/arxiv.2307.13854","metadata_source":"pith","pith_arxiv_id":"2307.13854","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"WebArena: A Realistic Web Environment for Building Autonomous Agents","venue":"cs.AI","work_id":"7058ffd2-a339-4102-89eb-248eeb074652","year":2023},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"cited_paper":"/paper/2307.13854","citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:2d4114b7896d57287bcd42894d174d151ee0e29b4477a7671fb8c35e34998171","observation_id":"28172c0f-e066-4465-9f98-3e897d6c0f6f","resolution":{"observed_at":"2026-06-30T14:44:45.290322Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-15T07:08:14.424698+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-15T07:08:14.424698+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.15043","last_updated":"2023-12-20T20:48:57Z","snapshot_observed_at":"2026-08-12T09:06:50.363435Z","submitted_at":"2023-07-27T17:49:12Z","title":"Universal and Transferable Adversarial Attacks on Aligned Language Models","version":2},"cited_work":{"arxiv_id":"2307.15043","doi":"10.48550/arxiv.2307.15043","metadata_source":"pith","pith_arxiv_id":"2307.15043","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Universal and Transferable Adversarial Attacks on Aligned Language Models","venue":"cs.CL","work_id":"3322fa86-1768-4677-8425-dd326b45e078","year":2023},"citing_paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-06-30T14:38:43.708578Z"},"links":{"cited_paper":"/paper/2307.15043","citing_paper":"/paper/2605.24134"},"observation_digest":"sha256:4ec0d5cb5530e575c398bb0c25ee64da75ef38d3a74577b3ec0b3622bcb006f7","observation_id":"4b09c56d-48c9-40a3-9663-6108f62a622b","resolution":{"observed_at":"2026-06-30T14:44:45.269769Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-07-15T23:50:40.271168+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-15T23:50:40.271168+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2605.24134","last_updated":"2026-05-22T18:52:34Z","latest_version":1,"primary_category":"cs.MA","snapshot_observed_at":"2026-08-13T09:50:28.297111Z","submitted_at":"2026-05-22T18:52:34Z","title":"ProofAgent Harness: Open Infrastructure for Adversarial Evaluation of AI Agents"},"reference_resolution":{"displayed":20,"state_counts":{"malformed_identifier":0,"metadata_mismatch":2,"parse_uncertain":0,"unresolved":1,"verified_exact":7,"verified_fuzzy":10},"total_outbound_references":20},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"thesis":"As of 24 August 2026, this Paper Citation Record lists 20 of 20 outbound references and 2 inbound Pith citation observations for arXiv:2605.24134."}