{"as_of":"2026-08-23T08:11:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:142103d9fdc0b0f89c5e4500913080d1a23c805fcd57cd380167621e1b0f7933","coverage":[{"denominator":82,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":82,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-29T16:51:36.524194Z","state":"measured"},{"denominator":82,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":82,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-23T06:30:58.430688+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2605.27209/citation-record","integrity":"/paper/2605.27209/integrity","json":"/paper/2605.27209/citation-record.json","paper":"/paper/2605.27209"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Introducing gpt-5.2","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:2c81de8401d72a4250b5a092642ee6fe239d87f1b57c5ec98c3affcade681a33","observation_id":"713e2198-bc20-4571-ba04-6ad493b5de81","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Gemini 3 pro model card.https://storage.googleapis.com/deepmind-media/Model- Cards/Gemini-3-Pro-Model-Card.pdf, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:f56834451965e4d34af5a275fec3b461cf2232369533e9daea8b46b8f8f4187d","observation_id":"00bfa7d6-4078-4bfc-aa3a-5df18d048fd0","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2509.18883","doi":"10.48550/arxiv.2509.18883","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Introducing LongCat-flash-thinking: A technical report","venue":"arXiv (Cornell University)","work_id":"d83b56e1-68ce-4e32-b9c0-fc080b79c8fd","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:7685f6b687ce743339c893abba0efd85fe10081ca2af9c5ae7aaa0ffd2d35821","observation_id":"03deb8b2-7042-4539-811a-5bc0f35fb954","resolution":{"observed_at":"2026-06-29T16:53:40.563419Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.20534","last_updated":"2026-02-03T04:57:00Z","snapshot_observed_at":"2026-08-16T14:37:33.548231Z","submitted_at":"2025-07-28T05:35:43Z","title":"Kimi K2: Open Agentic Intelligence","version":2},"cited_work":{"arxiv_id":"2507.20534","doi":"10.1145/3448609","metadata_source":"pith","pith_arxiv_id":"2507.20534","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Kimi K2: Open Agentic Intelligence","venue":"cs.LG","work_id":"7f18284c-12d3-4137-bea1-1da97e8cf3c1","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2507.20534","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:1e5c89bc89d14a2ad8a713bb319cb9440ea1e07acdb1c5fce6368964d58ab5ec","observation_id":"6945d13c-f20e-4658-a1e2-4c6c5031ff8c","resolution":{"observed_at":"2026-06-29T16:53:40.534381Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-05-25T01:23:16.170083+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-25T01:23:16.170083+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.06471","last_updated":"2025-08-08T17:21:06Z","snapshot_observed_at":"2026-08-11T03:34:09.767397Z","submitted_at":"2025-08-08T17:21:06Z","title":"GLM-4.5: Agentic, Reasoning, and Coding (ARC) Foundation Models","version":1},"cited_work":{"arxiv_id":"2508.06471","doi":"10.48550/arxiv.2508.06471","metadata_source":"pith","pith_arxiv_id":"2508.06471","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"GLM-4.5: Agentic, Reasoning, and Coding (ARC) Foundation Models","venue":"cs.CL","work_id":"5bb4e5d7-985e-431b-bb0d-75576cdc2950","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2508.06471","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:10a56eca45e5976b13281ac4e705a1e1c84875d221f4ffd27671271eb77e89d3","observation_id":"f1eaeb71-4cc9-4b20-b5ef-66c9315e4ed2","resolution":{"observed_at":"2026-06-29T16:53:40.577567Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-05-23T05:23:01.474969+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T05:23:01.474969+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2509.01322","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-03T09:37:49.308056Z","title":"Longcat-flash technical report","venue":null,"work_id":"112bbed5-cf36-4773-82b3-0229bfea4626","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:35281970f16bbc208725d492c31618580c3d55aeb998259a54b9aa2dbc934974","observation_id":"2de6fdfc-0348-4a0d-b219-e2713c7aa540","resolution":{"observed_at":"2026-06-29T16:53:40.606326Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12045","last_updated":"2024-06-17T19:33:08Z","snapshot_observed_at":"2026-08-17T20:31:29.818313Z","submitted_at":"2024-06-17T19:33:08Z","title":"$\\tau$-bench: A Benchmark for Tool-Agent-User Interaction in Real-World Domains","version":1},"cited_work":{"arxiv_id":"2406.12045","doi":"10.48550/arxiv.2406.12045","metadata_source":"pith","pith_arxiv_id":"2406.12045","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"$\\tau$-bench: A Benchmark for Tool-Agent-User Interaction in Real-World Domains","venue":"cs.AI","work_id":"6a8d8dc4-0cc0-4052-8109-abbcdcd4a962","year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2406.12045","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:486a10457f6a31ad9f7bb377609cc4edd8ca25d1f4feb5795ae5b8194754e0d6","observation_id":"69514999-1d97-4372-8d9e-47b4cadd20bb","resolution":{"observed_at":"2026-06-29T16:53:40.072292Z","resolver_source":"local_arxiv","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-07-14T18:20:21.86453+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-14T18:20:21.86453+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.07982","last_updated":"2025-06-09T17:52:18Z","snapshot_observed_at":"2026-08-14T06:34:01.114459Z","submitted_at":"2025-06-09T17:52:18Z","title":"$\\tau^2$-Bench: Evaluating Conversational Agents in a Dual-Control Environment","version":1},"cited_work":{"arxiv_id":"2506.07982","doi":"10.48550/arxiv.2506.07982","metadata_source":"pith","pith_arxiv_id":"2506.07982","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"$\\tau^2$-Bench: Evaluating Conversational Agents in a Dual-Control Environment","venue":"cs.AI","work_id":"3a498b1a-455f-4667-b572-c5216c99a89c","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2506.07982","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:1a070f1fba85a2ce7925bc3bb5ecd59325e4f1052a9fa74555e32bad22c2659a","observation_id":"5db102f9-01ef-4e8c-b681-d52a25a652a9","resolution":{"observed_at":"2026-06-29T16:53:40.608406Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-07-14T18:20:22.129969+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-14T18:20:22.129969+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2509.26490","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T10:39:44.801876Z","title":"Vitabench: Benchmarking llm agents with versatile interactive tasks in real-world applications","venue":null,"work_id":"a9d216f9-8bb8-457e-a320-d1457744f75b","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:3abb84c38807684386b32b9460b48fd21c84cdc141c0014e0b1faa8a5264125d","observation_id":"3673533f-c22d-4334-bd58-d7853dca0dda","resolution":{"observed_at":"2026-06-29T16:53:40.522389Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Mind2web: Towards a generalist agent for the web.Advances in Neural Information Processing Systems, 36:28091–28114, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:0c8ec88d6dbd280c13984e95cd84eee84c8fdf7590f08a17df43f942551f8593","observation_id":"33d84901-61ff-4cac-b5bc-655b3ef7ba5c","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.13854","last_updated":"2024-04-16T15:13:18Z","snapshot_observed_at":"2026-08-14T11:14:55.351653Z","submitted_at":"2023-07-25T22:59:32Z","title":"WebArena: A Realistic Web Environment for Building Autonomous Agents","version":4},"cited_work":{"arxiv_id":"2307.13854","doi":"10.48550/arxiv.2307.13854","metadata_source":"pith","pith_arxiv_id":"2307.13854","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"WebArena: A Realistic Web Environment for Building Autonomous Agents","venue":"cs.AI","work_id":"7058ffd2-a339-4102-89eb-248eeb074652","year":2023},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2307.13854","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:7d52ea0f97e7f86322528a8b7a2e8881e2d139db26149a01282e119d0f957671","observation_id":"dd81adb2-6535-4e0f-be3a-f86a48702946","resolution":{"observed_at":"2026-06-29T16:53:40.621593Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-15T07:08:14.424698+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-15T07:08:14.424698+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2504.01382","doi":"10.48550/arxiv.2504.01382","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"An illusion of progress? assessing the current state of web agents","venue":"ArXiv.org","work_id":"faf58fb2-da4e-413b-a68e-c2c7438f7c62","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:baa20857ab09fa91c5e402396f28730fe569c544cb1c6bec95c3a54fa00e2454","observation_id":"7a0d0305-1d87-4abc-8ddd-f92200527cce","resolution":{"observed_at":"2026-06-29T16:53:40.582635Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Agenttuning: Enabling generalized agent abilities for llms","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:bca6d1d86196ac6bea0ff062b9b088426497c1e2ed99ec6de8f648c3abfe8bca","observation_id":"6b7c2918-a3c4-4cc8-a229-296ae267deec","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.02337","last_updated":"2025-01-27T11:56:15Z","snapshot_observed_at":"2026-08-18T10:06:51.290180Z","submitted_at":"2024-11-04T17:59:58Z","title":"WebRL: Training LLM Web Agents via Self-Evolving Online Curriculum Reinforcement Learning","version":3},"cited_work":{"arxiv_id":"2411.02337","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2411.02337","snapshot_observed_at":"2026-07-04T17:20:00.041093Z","title":"Webrl: Training llm web agents via self-evolving online curriculum reinforcement learning","venue":null,"work_id":"ceab60f7-8cb0-4202-b6d8-973b47e933fc","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2411.02337","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:2d66dfc24fdc5fcd1db49d706647297b19207628260edd9c381cb7837efea90e","observation_id":"3d93c6be-e952-4002-9d0c-81f60346f689","resolution":{"observed_at":"2026-06-29T16:53:40.605741Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Communication accommodation theory.Theo- rizing about intercultural communication, pages 121–148, 2005","venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:16fcf1c876cfe2d51831ac5f0862ee55e995639c07887183f7648175ea2da7bf","observation_id":"29efeb0b-611d-45cb-93d8-a333c88c2673","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"What do users really ask large language models? an initial log analysis of google bard interactions in the wild","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:1876eeb2d7da54e1bfe0302fba5bed9bdc48c5f8440dcc1d362664681ec694f2","observation_id":"e272e8ed-3f32-4595-9711-38fc902b6e5d","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2509.25238","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-02T16:47:09.209768Z","title":"PALADIN: Self-correcting language model agents to cure tool-failure cases","venue":null,"work_id":"6f7b83ca-36ff-4352-86a8-4d1bf09becdf","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:d2c99bd0da2ae3061518c17fb504ee60a21deeb01935488956d3449e15ca405b","observation_id":"e1bafd98-451d-46b0-b152-80b648d93340","resolution":{"observed_at":"2026-06-29T16:53:40.611542Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.15296","last_updated":"2025-07-21T06:55:37Z","snapshot_observed_at":"2026-08-21T05:28:04.149419Z","submitted_at":"2025-07-21T06:55:37Z","title":"Butterfly Effects in Toolchains: A Comprehensive Analysis of Failed Parameter Filling in LLM Tool-Agent Systems","version":1},"cited_work":{"arxiv_id":"2507.15296","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.15296","snapshot_observed_at":"2026-07-02T12:16:56.286769Z","title":"Ruijie Xu, Zengzhi Wang, Run-Ze Fan, and Pengfei Liu","venue":null,"work_id":"4bb064d6-7897-4d8f-855e-e2d0de443cb2","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2507.15296","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:8e4dea80ddc67e508e8400a5776a6d0146a876fc1005de3bfe8807603bc07711","observation_id":"87714e8d-ba27-49b3-94c8-09466e150bbe","resolution":{"observed_at":"2026-06-29T16:53:40.615048Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.14573","last_updated":"2025-04-06T20:37:50Z","snapshot_observed_at":"2026-08-19T11:49:14.266180Z","submitted_at":"2024-05-23T13:48:54Z","title":"AndroidWorld: A Dynamic Benchmarking Environment for Autonomous Agents","version":5},"cited_work":{"arxiv_id":"2405.14573","doi":"10.48550/arxiv.2405.14573","metadata_source":"pith","pith_arxiv_id":"2405.14573","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"AndroidWorld: A Dynamic Benchmarking Environment for Autonomous Agents","venue":"cs.AI","work_id":"c5116d19-d3d3-40fd-9620-f7489812a9ba","year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2405.14573","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:30e08376701e01e0803434dc12e1f4abd4490b2352a6742e2afe321082b6ca3e","observation_id":"11c90cc8-67e5-426a-939f-5e5cd3777ae8","resolution":{"observed_at":"2026-06-29T16:53:40.585685Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Gui-xplore: Empowering generalizable gui agents with one exploration","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:b28afabccb0cc3953bfe3430a52dec1ea147b20eeb2c30f848b064a00892cc90","observation_id":"6fcc4d92-9867-45d3-82a3-eb6be7fff78d","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Out-of-distribution segmentation in autonomous driving: Problems and state of the art","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:ebf2a102793f629ab33514fd3df769113046394ce29e08a8aa3ce8519026c5f8","observation_id":"38590bdf-db02-4e16-9400-a3af4308d3b4","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Domain randomization for transferring deep neural networks from simulation to the real world","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:7b0a13fefd1e86a269971d77a0930ed1b4820aaced0f7ed77e7b0163964eeb7c","observation_id":"2de9696a-7a3f-4c7e-8e7c-736369a64c9c","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Cad2rl: Real single-image flight without a single real image","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:ad1864ab9e5ba6fb673a271d18bac94f810a0c4405aaca71a4d76bb00db02822","observation_id":"f6a77746-e317-46aa-99b3-8e25a5e20796","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Robust reinforcement learning as a stackelberg game","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:a024b5bbd4d5da56450aca36da0009bb8ef91e7f7db3e2b001f303fac683fd51","observation_id":"4b06a032-97bd-4703-874d-4e9d2d116afb","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2509.02547","last_updated":"2026-04-17T18:09:08Z","snapshot_observed_at":"2026-08-03T09:07:42.489237Z","submitted_at":"2025-09-02T17:46:26Z","title":"The Landscape of Agentic Reinforcement Learning for LLMs: A Survey","version":5},"cited_work":{"arxiv_id":"2509.02547","doi":"10.48550/arxiv.2509.02547","metadata_source":"pith","pith_arxiv_id":"2509.02547","snapshot_observed_at":"2026-08-05T02:49:54.815029Z","title":"The Landscape of Agentic Reinforcement Learning for LLMs: A Survey","venue":"cs.AI","work_id":"87909127-da20-4ccc-8ae3-4a4a20ef81b7","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2509.02547","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:2b5e7b7773115bdca52fe9840c34b3daa700b001b8497581356017df9679b3d8","observation_id":"5a28c917-b334-4e04-b1b9-1c306c65e98e","resolution":{"observed_at":"2026-06-29T16:53:40.551897Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-17T17:38:13.840562+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-17T17:38:13.840562+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2509.01055","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-03T05:47:41.476341Z","title":"Verltool: Towards holistic agentic reinforcement learning with tool use","venue":null,"work_id":"b23dbbef-f1e5-4b1a-8ae5-6dcca5e78dc8","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:f97ebc6f1b8a1d5d83321bf2e7d15b49f9a9f37738d6e938171ad1e2703f2596","observation_id":"23cb580c-76b7-4ba2-92c4-0e4b53232f5c","resolution":{"observed_at":"2026-06-29T16:53:40.614622Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06347","last_updated":"2017-08-28T09:20:06Z","snapshot_observed_at":"2026-08-20T07:04:06.309989Z","submitted_at":"2017-07-20T02:32:33Z","title":"Proximal Policy Optimization Algorithms","version":2},"cited_work":{"arxiv_id":"1707.06347","doi":"10.1016/j.artint.2010.12.005","metadata_source":"pith","pith_arxiv_id":"1707.06347","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Proximal Policy Optimization Algorithms","venue":"cs.LG","work_id":"240c67fe-d14d-4520-91c1-38a4e272ca19","year":2017},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/1707.06347","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:cb7518132c27f9b1872eebaf8a3d341ed3ef302c2483fdf45d31d44a5919c1d5","observation_id":"fea48980-c37e-4bb5-9839-aa850c82a677","resolution":{"observed_at":"2026-06-29T16:53:40.591708Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2601.16725","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T19:50:10.434955Z","title":"Longcat-flash-thinking-2601 technical report","venue":null,"work_id":"1375347d-bbc3-47bf-a423-d4e39ac50144","year":2026},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:4b6bda43ae7a4c7b4f34dea7092649001a8335b2fafa647b03c41f7e09a47ba7","observation_id":"1b044b96-4093-4e80-a2f9-c78355cfa3f3","resolution":{"observed_at":"2026-06-29T16:53:40.612140Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2512.02556","last_updated":"2025-12-02T09:25:14Z","snapshot_observed_at":"2026-08-17T01:58:07.850738Z","submitted_at":"2025-12-02T09:25:14Z","title":"DeepSeek-V3.2: Pushing the Frontier of Open Large Language Models","version":1},"cited_work":{"arxiv_id":"2512.02556","doi":"10.18653/v1/d18-1512","metadata_source":"pith","pith_arxiv_id":"2512.02556","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"DeepSeek-V3.2: Pushing the Frontier of Open Large Language Models","venue":"cs.CL","work_id":"07c85cc5-4086-4abc-823b-6d0f4ff784d0","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2512.02556","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:8ca3bba56d787b21a3a9f76eb5e61405386047ffa4932e5eb7454b72df3746df","observation_id":"84e2b736-8760-4f70-b9e8-9fdf3271a84a","resolution":{"observed_at":"2026-06-29T16:53:40.554631Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2602.06820","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T20:05:33.983652Z","title":"Scaleenv: Scaling environment synthesis from scratch for generalist interactive tool-use agent training","venue":null,"work_id":"317dbc65-000f-4ec0-a39b-b32add74a01b","year":2026},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:b829f95b2c05cff90df678e6e5b0b9d2b0c6e10cfeacb03e67ede9f831eb127f","observation_id":"fa0fe347-3360-40c4-8658-2a4d3beb4f45","resolution":{"observed_at":"2026-06-29T16:53:40.563163Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2602.11348","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T20:00:08.729711Z","title":"tasks\": [ {","venue":null,"work_id":"ad2d9422-171a-4205-afae-f10160299e71","year":2026},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:bb524d95c84413f19137ae99d2b927c3127467a7adc54849f25ab1c855614c69","observation_id":"4b9bb57c-02f1-47ed-a722-c6027573ca75","resolution":{"observed_at":"2026-06-29T16:53:40.618707Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"React: Synergizing reasoning and acting in language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:c3d31ec38e4e0d5bb6b4711ca5c9af50804e483cf4e1c3e8c558d79969c93494","observation_id":"f6842076-5548-44a1-9c09-0086c9f3979a","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Toolformer: Language models can teach themselves to use tools.Advances in Neural Information Processing Systems, 36: 68539–68551, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:4a9a7ad6b67b63d3ab2e8338df0c3bd5874019ad7b6d44bf9516adb6d5aefcbe","observation_id":"d6869748-f129-4af6-b71b-5404e08b8007","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Reflexion: Language agents with verbal reinforcement learning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:c1ca445d759fbc16410f88ad97628effa01494d2b44de56c64531204a11f9ae1","observation_id":"d67cfd49-c45f-4a08-9463-de5f17d5b41d","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.16291","last_updated":"2023-10-19T16:27:03Z","snapshot_observed_at":"2026-08-13T16:55:46.498156Z","submitted_at":"2023-05-25T17:46:38Z","title":"Voyager: An Open-Ended Embodied Agent with Large Language Models","version":2},"cited_work":{"arxiv_id":"2305.16291","doi":"10.18653/v1/2023.emnlp-main.118","metadata_source":"pith","pith_arxiv_id":"2305.16291","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Voyager: An Open-Ended Embodied Agent with Large Language Models","venue":"cs.AI","work_id":"ffe0d207-86cf-4742-a100-e988ac8b9676","year":2023},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2305.16291","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:37e79e63956aa32202cc10821aa3f2fc72bea5086143bbc1d54c72abae3db72e","observation_id":"c021fb3c-31ce-44b0-bfe2-e8030e42188f","resolution":{"observed_at":"2026-06-29T16:53:40.506779Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.08155","last_updated":"2023-10-03T20:47:10Z","snapshot_observed_at":"2026-08-08T22:28:26.004138Z","submitted_at":"2023-08-16T05:57:52Z","title":"AutoGen: Enabling Next-Gen LLM Applications via Multi-Agent Conversation","version":2},"cited_work":{"arxiv_id":"2308.08155","doi":"10.48550/arxiv.2308.08155","metadata_source":"pith","pith_arxiv_id":"2308.08155","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"AutoGen: Enabling Next-Gen LLM Applications via Multi-Agent Conversation","venue":"cs.AI","work_id":"92b7eb9c-c3d8-4518-a376-06fa15dd895b","year":2023},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2308.08155","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:f2f2853b2fab58ab0503245b2846ee9b4ecace7af8f66f7fadf0a6bb9ac30c37","observation_id":"33ec84d1-a8e2-4ca3-85e3-30ddf0037b2d","resolution":{"observed_at":"2026-06-29T16:53:40.602980Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-07-14T18:20:20.607676+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-14T18:20:20.607676+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Metagpt: Meta programming for a multi-agent collaborative framework","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:eaa6130f29be3149d61dbeb964cefdf1bb1a3c58add39e6fb88e72be6112ab42","observation_id":"93e01972-2f06-4ec6-b610-15db4c1b357b","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Generative agents: Interactive simulacra of human behavior","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:38e7a715a580e91b7dbb81c8dc87e06db4ded884988c869ffd7327dcd93ee5f9","observation_id":"c02340fd-ee98-4379-b657-8db9db2413fb","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-08-15T12:33:55.451951Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":"2501.12948","doi":"10.1016/j.artmed.2024.103001","metadata_source":"pith","pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","venue":"cs.CL","work_id":"e6b75ad5-2877-4168-97c8-710407094d20","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:a535aba13a71ac459ed0f900fba85f2acb718aa5d0228466aaa5cbdccf157171","observation_id":"5089337f-f44f-4a9c-8c1a-7272975e7133","resolution":{"observed_at":"2026-06-29T16:53:40.510676Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.15124","last_updated":"2025-04-14T22:39:09Z","snapshot_observed_at":"2026-08-17T14:11:00.232598Z","submitted_at":"2024-11-22T18:44:04Z","title":"Tulu 3: Pushing Frontiers in Open Language Model Post-Training","version":5},"cited_work":{"arxiv_id":"2411.15124","doi":"10.48550/arxiv.2411.15124","metadata_source":"pith","pith_arxiv_id":"2411.15124","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Tulu 3: Pushing Frontiers in Open Language Model Post-Training","venue":"cs.CL","work_id":"28c9dbea-056a-48c2-8000-85f809827e45","year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2411.15124","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:22dbc09888cce54e8c56b9c375c5b9600e9fafb1c0285d1144af010302cb23b6","observation_id":"26998045-7c36-4d1e-955a-8c801dc0eb26","resolution":{"observed_at":"2026-06-29T16:53:40.498618Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:00.522112+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:00.522112+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03300","last_updated":"2024-04-27T15:25:53Z","snapshot_observed_at":"2026-08-06T14:58:42.911363Z","submitted_at":"2024-02-05T18:55:32Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","version":3},"cited_work":{"arxiv_id":"2402.03300","doi":"10.1016/0004-3702(73)90011-8","metadata_source":"pith","pith_arxiv_id":"2402.03300","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","venue":"cs.CL","work_id":"c5006563-f3ec-438a-9e35-b7b484f34828","year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2402.03300","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:bb22e093b2c37c60efce48bbecee95ea0c4f8f5d8b40554a241e82681bfedbc9","observation_id":"882c0224-ff88-4c58-979c-70cd0fe33887","resolution":{"observed_at":"2026-06-29T16:53:40.501402Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.14476","last_updated":"2025-05-20T01:37:34Z","snapshot_observed_at":"2026-08-18T05:01:20.543826Z","submitted_at":"2025-03-18T17:49:06Z","title":"DAPO: An Open-Source LLM Reinforcement Learning System at Scale","version":2},"cited_work":{"arxiv_id":"2503.14476","doi":"10.48550/arxiv.2503.14476","metadata_source":"pith","pith_arxiv_id":"2503.14476","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"DAPO: An Open-Source LLM Reinforcement Learning System at Scale","venue":"cs.LG","work_id":"64019d00-0b11-4bbd-b173-b46c8fad0157","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2503.14476","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:929bd34970ad36593900b6fe51932de18b8ce78abe4f0228d7b8ac906a2b621e","observation_id":"94ecf3d6-8a1f-40a2-bdcd-05a06c64b33f","resolution":{"observed_at":"2026-06-29T16:53:40.560597Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-08T16:08:20.547492+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T16:08:20.547492+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.18071","last_updated":"2025-07-28T11:11:33Z","snapshot_observed_at":"2026-08-21T02:04:30.074424Z","submitted_at":"2025-07-24T03:50:32Z","title":"Group Sequence Policy Optimization","version":2},"cited_work":{"arxiv_id":"2507.18071","doi":"10.48550/arxiv.2507.18071","metadata_source":"pith","pith_arxiv_id":"2507.18071","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Group Sequence Policy Optimization","venue":"cs.LG","work_id":"3a98b53b-9f52-4d95-adf7-89353c0a9a65","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2507.18071","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:dc03aeef3d0837c6203120007cfdf59b35ec8cca082286438f765ebb0605a0ff","observation_id":"b9e958a2-a17c-4a93-ba90-08115ae5904e","resolution":{"observed_at":"2026-06-29T16:53:40.527831Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-15T03:08:21.960036+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-15T03:08:21.960036+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.20783","last_updated":"2025-10-06T09:30:03Z","snapshot_observed_at":"2026-08-13T12:34:54.476684Z","submitted_at":"2025-03-26T17:59:14Z","title":"Understanding R1-Zero-Like Training: A Critical Perspective","version":2},"cited_work":{"arxiv_id":"2503.20783","doi":"10.48550/arxiv.2503.20783","metadata_source":"pith","pith_arxiv_id":"2503.20783","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Understanding R1-Zero-Like Training: A Critical Perspective","venue":"cs.LG","work_id":"ec354f3b-9484-4a0c-94c8-92d4d0260835","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2503.20783","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:f4b7d5cb6bb8570b05155e9cabd7dec8f4e3313cc92d268794f6b5210afa989d","observation_id":"625e3557-897d-4a35-b041-62190c0b70a5","resolution":{"observed_at":"2026-06-29T16:53:40.482519Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-05-24T09:23:05.84445+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-24T09:23:05.84445+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.05118","last_updated":"2025-04-11T02:54:58Z","snapshot_observed_at":"2026-08-14T16:01:52.772456Z","submitted_at":"2025-04-07T14:21:11Z","title":"VAPO: Efficient and Reliable Reinforcement Learning for Advanced Reasoning Tasks","version":3},"cited_work":{"arxiv_id":"2504.05118","doi":"10.1109/access.2024.3384487","metadata_source":"pith","pith_arxiv_id":"2504.05118","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"VAPO: Efficient and Reliable Reinforcement Learning for Advanced Reasoning Tasks","venue":"cs.AI","work_id":"c2351652-65f7-47cd-ae80-dbcd72a6eb20","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2504.05118","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:d78ce22439c35aa60601f81e22083ac0307327c336fff1b677aee08d3a9f2e69","observation_id":"1a170f90-bce1-40cf-8f98-67734c5206f0","resolution":{"observed_at":"2026-06-29T16:53:40.530817Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2602.03048","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-03T04:27:36.740990Z","title":"Coba-rl: Capability-oriented budget allocation for reinforcement learning in llms","venue":null,"work_id":"dfc92b4c-afc7-4195-b7d6-d28f98ea47f6","year":2026},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:24a1ce40a4cb434fa308b665de9f1f8bad03506ed9b49cb48cc39169527413cb","observation_id":"f6b814b9-6cb1-4ab2-8e64-de37b78cc4e7","resolution":{"observed_at":"2026-06-29T16:53:40.516552Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.11536","last_updated":"2025-04-17T16:46:07Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-04-15T18:10:22Z","title":"ReTool: Reinforcement Learning for Strategic Tool Use in LLMs","version":2},"cited_work":{"arxiv_id":"2504.11536","doi":"10.48550/arxiv.2504.11536","metadata_source":"pith","pith_arxiv_id":"2504.11536","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"ReTool: Reinforcement Learning for Strategic Tool Use in LLMs","venue":"cs.CL","work_id":"6da510e4-e55c-4412-b7c8-839984508e99","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2504.11536","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:98d755d43d37d396147608b0ef7b25e0bad9e55654fcebbc37eb3eb1ae367606","observation_id":"c7e31f8d-1aba-4ead-924d-234f4cba419b","resolution":{"observed_at":"2026-06-29T16:53:40.603449Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.09516","last_updated":"2025-08-05T19:08:38Z","snapshot_observed_at":"2026-08-15T13:17:00.526689Z","submitted_at":"2025-03-12T16:26:39Z","title":"Search-R1: Training LLMs to Reason and Leverage Search Engines with Reinforcement Learning","version":5},"cited_work":{"arxiv_id":"2503.09516","doi":"10.48550/arxiv.2503.09516","metadata_source":"pith","pith_arxiv_id":"2503.09516","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Search-R1: Training LLMs to Reason and Leverage Search Engines with Reinforcement Learning","venue":"cs.CL","work_id":"0e0b7549-2bc4-4574-aa7f-588ffa16eaae","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2503.09516","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:f4e574285e3455b7a331c433565be64c0bfeaa7df42085e2327469c43058303d","observation_id":"db3a0a16-5476-4024-b7b9-e90f6f1f6f9d","resolution":{"observed_at":"2026-06-29T16:53:40.630043Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.05592","last_updated":"2025-03-18T08:32:24Z","snapshot_observed_at":"2026-08-15T04:19:26.319323Z","submitted_at":"2025-03-07T17:14:44Z","title":"R1-Searcher: Incentivizing the Search Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":"2503.05592","doi":"10.48550/arxiv.2503.05592","metadata_source":"pith","pith_arxiv_id":"2503.05592","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"R1-Searcher: Incentivizing the Search Capability in LLMs via Reinforcement Learning","venue":"cs.AI","work_id":"f5ed73d2-f2ff-4cf6-853d-3586333e44ef","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2503.05592","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:3b860ff0a83deb5831f6acdd5645d04fefd4b739f3ec5451bc5c1cfd5e676013","observation_id":"01d5be0e-72d3-456a-b302-19cb2d04afee","resolution":{"observed_at":"2026-06-29T16:53:40.600912Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2509.23040","doi":"10.48550/arxiv.2509.23040","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Look back to reason forward: Revis- itable memory for long-context llm agents.arXiv preprint arXiv:2509.23040, 2025a","venue":"arXiv (Cornell University)","work_id":"74006986-b24a-48c9-a5c0-d5cfbc5a76ab","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:1fb1e77b6fbd926f41c97d0469bea75985a073de0533a9d34f9910252c8d26df","observation_id":"6daba692-88d3-4268-8629-83132e4cd289","resolution":{"observed_at":"2026-06-29T16:53:40.624121Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Jimenez, John Yang, Alexander Wettig, Shunyu Yao, Kexin Pei, Ofir Press, and Karthik Narasimhan","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:abc9513338557f9d47eb7abcc268d88fb50104ee2a209c6738e762c4d9c8e2da","observation_id":"9ae4d2e8-a9b4-4fd8-b4e6-941a76c35e22","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.21139","last_updated":"2025-06-06T07:53:20Z","snapshot_observed_at":"2026-08-16T19:05:11.273764Z","submitted_at":"2024-12-30T18:15:39Z","title":"Training Software Engineering Agents and Verifiers with SWE-Gym","version":2},"cited_work":{"arxiv_id":"2412.21139","doi":"10.48550/arxiv.2412.21139","metadata_source":"pith","pith_arxiv_id":"2412.21139","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Training Software Engineering Agents and Verifiers with SWE-Gym","venue":"cs.SE","work_id":"17189f19-7774-4b97-ab44-9966bf5d6d48","year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2412.21139","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:dcc4036662c88b8dabb8e002adbe3aafab3a7ea7b90d5b39d9649883f6edb574","observation_id":"f490a6e6-4bc6-4117-bc7b-c7a6467dcaaf","resolution":{"observed_at":"2026-06-29T16:53:40.583201Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.18449","last_updated":"2025-12-01T00:16:59Z","snapshot_observed_at":"2026-08-13T07:37:18.494967Z","submitted_at":"2025-02-25T18:45:04Z","title":"SWE-RL: Advancing LLM Reasoning via Reinforcement Learning on Open Software Evolution","version":2},"cited_work":{"arxiv_id":"2502.18449","doi":"10.48550/arxiv.2502.18449","metadata_source":"pith","pith_arxiv_id":"2502.18449","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SWE-RL: Advancing LLM Reasoning via Reinforcement Learning on Open Software Evolution","venue":"cs.SE","work_id":"4b93fb93-87c9-40fe-84d0-d7ecb4e11bed","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2502.18449","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:c68a363f9130f96a5520e854e4bcab292e524b04c0b2e40d1d4231ebd0165a91","observation_id":"a2569508-16e4-42e9-a68e-825d2831a632","resolution":{"observed_at":"2026-06-29T16:53:40.588646Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Xu, Hao Zhu, Xuhui Zhou, Robert Lo, Abishek Sridhar, Xianyi Cheng, Tianyue Ou, Yonatan Bisk, Daniel Fried, Uri Alon, and Graham Neubig","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:9bbe596cb0412c24c17121da544011e3bd52b9f5262e0d7b8009f89d05deb50f","observation_id":"daae1af8-4483-40e3-969c-48de82746a9a","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"OSWorld: Benchmarking multimodal agents for open-ended tasks in real computer environments","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:7bdba39c8fbf0db829ae7fee666a202d0b662d0e2c49910066cdb38d370aeccf","observation_id":"f635e77c-3d91-4871-b767-ed2f97d6e75b","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"AppWorld: A controllable world of apps and people for benchmarking interactive coding agents","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:10619c3a246f24f0b72fe2a0857580c4f0304e9af7e1e9e0505cc22b5315babe","observation_id":"0ae4f2f6-21d0-405c-b86d-f918b9788295","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2503.20197","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:53:40.584018Z","title":"arXiv preprint arXiv:2503.20197 , year=","venue":null,"work_id":"9e52cea0-b363-421e-b191-a1605631f30d","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:7ae168c0d7ff7d842239e2ca58cd6e4e21c8fe00ce1c824e1845ad2a028d2401","observation_id":"b9ea1352-8d00-4145-a1c0-b2fb50b1b807","resolution":{"observed_at":"2026-06-29T16:53:40.585738Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.17723","last_updated":"2025-07-24T08:03:09Z","snapshot_observed_at":"2026-08-20T23:44:53.009005Z","submitted_at":"2025-04-24T16:36:19Z","title":"Statistical Runtime Verification for LLMs via Robustness Estimation","version":2},"cited_work":{"arxiv_id":"2504.17723","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.17723","snapshot_observed_at":"2026-06-29T16:53:40.593038Z","title":"Towards robust LLMs: An adversarial robustness measurement framework","venue":null,"work_id":"6e9922f8-c4fb-460d-863e-0105e6c6b21b","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2504.17723","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:f3a4d4bb58ab2a1a7aab10e9b3add84373c60c767aecc4f683e63fc800c19301","observation_id":"18ac6053-ae45-400d-9a6f-a88958f5dad7","resolution":{"observed_at":"2026-06-29T16:53:40.595190Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.02733","last_updated":"2025-04-03T16:17:56Z","snapshot_observed_at":"2026-08-17T23:05:32.639190Z","submitted_at":"2025-04-03T16:17:56Z","title":"Enhancing LLM Robustness to Perturbed Instructions: An Empirical Study","version":1},"cited_work":{"arxiv_id":"2504.02733","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.02733","snapshot_observed_at":"2026-07-03T20:58:57.691842Z","title":"org/abs/2504.02733","venue":null,"work_id":"dd4a394b-b948-4d86-a07d-ea1173c52c83","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2504.02733","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:56645c84daa21059fa3393fb882ba0268ccc9f35059393b5276d18a055173eb8","observation_id":"6a0a9f1e-6e8e-4f2e-9341-2deb21de8970","resolution":{"observed_at":"2026-06-29T16:53:40.576801Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Anghel, Emilia Pecheanu, Adina Cocu, Adrian Istrate, and Con- stantin A","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:5a2c1e6943957fa498dd584a3462d3d0043a9bcff971f4fbd4fcaffc1896dd64","observation_id":"cf381734-3029-4d8f-b706-ce5ed0630c69","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Robust LLM training infrastructure at ByteDance","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:01ffba13db40bea2b099d181789eb0eb8c64bf2675c4e78024ec049afb260724","observation_id":"d3d6b478-3e21-4027-a5fb-66fc052d7740","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:6024e3a51fdbfb6ce98f34b1cb58318e25a772f529f5caa3815b3cbae2abcc9b","observation_id":"5a19f7c8-a44e-4485-bafe-6f5e23b7ae0e","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Evaluating the performance and robustness of LLMs in materials science Q&A and property predictions.Digital Discovery, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:0ec806ce54b4a93f05d7d10befa3da77f76cedfcc1dcf8c88a8a7b132f8a3efb","observation_id":"ce4bffc3-619e-4e50-87fa-45a5f0011c33","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.04550","last_updated":"2025-03-06T15:36:06Z","snapshot_observed_at":"2026-08-19T22:15:03.482573Z","submitted_at":"2025-03-06T15:36:06Z","title":"Benchmarking Reasoning Robustness in Large Language Models","version":1},"cited_work":{"arxiv_id":"2503.04550","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2503.04550","snapshot_observed_at":"2026-06-29T16:53:40.544660Z","title":"Benchmarking reasoning robustness in large language models","venue":null,"work_id":"98a2ee94-365f-4718-8a68-3a4e88e7d151","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2503.04550","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:7e47bb0212f74076054148170b6dd3f7045f6486259609a429325c35d63c2b13","observation_id":"6c9f0343-f6fb-446a-a119-2f5a828c2255","resolution":{"observed_at":"2026-06-29T16:53:40.546212Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.14494","last_updated":"2025-05-30T04:20:31Z","snapshot_observed_at":"2026-08-16T12:56:41.889446Z","submitted_at":"2025-02-20T12:22:18Z","title":"StructFlowBench: A Structured Flow Benchmark for Multi-turn Instruction Following","version":2},"cited_work":{"arxiv_id":"2502.14494","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.14494","snapshot_observed_at":"2026-06-29T16:53:40.564193Z","title":"StructFlowBench: A structured flow benchmark for multi-turn instruction following.arXiv preprint arXiv:2502.14494, 2025","venue":null,"work_id":"db9cbf63-f92b-442c-8a50-594b2dd4d724","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2502.14494","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:c9076cda9b607067703b9ae0c9c0fdd80b44674bdd15f1baf61733c07c1ffae5","observation_id":"5b9004d8-d374-4eb5-ae40-09c7731055da","resolution":{"observed_at":"2026-06-29T16:53:40.565668Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Multichallenge: A realistic multi-turn conversation evaluation benchmark challenging to frontier llms","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:941e51ec43d4b1ef3075902d5906af3f0184532a9be0bb3955449ceef01e8fa8","observation_id":"4c348a4e-e675-418e-b6e0-acd3d0fc1b05","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.16944","last_updated":"2025-05-22T17:31:10Z","snapshot_observed_at":"2026-08-16T18:48:25.262848Z","submitted_at":"2025-05-22T17:31:10Z","title":"AGENTIF: Benchmarking Instruction Following of Large Language Models in Agentic Scenarios","version":1},"cited_work":{"arxiv_id":"2505.16944","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.16944","snapshot_observed_at":"2026-07-03T15:18:33.541374Z","title":"Agentif: Benchmarking instruction following of large language models in agentic scenarios","venue":null,"work_id":"e63898a9-4e78-48d1-83f9-8f73b6015f9b","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2505.16944","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:ab70813271eaa5525754b6f65fa3c0442f137defcec4fe131797f54bf6131a03","observation_id":"7d556ebd-4c18-473b-8be4-ed6554e9953e","resolution":{"observed_at":"2026-06-29T16:53:40.609311Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08329","last_updated":"2024-01-16T12:49:00Z","snapshot_observed_at":"2026-08-20T06:21:06.151853Z","submitted_at":"2024-01-16T12:49:00Z","title":"Understanding User Experience in Large Language Model Interactions","version":1},"cited_work":{"arxiv_id":"2401.08329","doi":"10.48550/arxiv.2401.08329","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.08329","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2401.08329 , year=","venue":"arXiv (Cornell University)","work_id":"6af57789-c067-43d0-9a79-dbd67d83a3f3","year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2401.08329","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:83c2a7929b5ca878500a1a7fd03f760220be2e56d80881efb11cfbe977b36776","observation_id":"3c30d5d9-f7d8-43b4-a02f-6f7c8abe691a","resolution":{"observed_at":"2026-06-29T16:53:40.627476Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.06097","last_updated":"2024-09-14T20:55:13Z","snapshot_observed_at":"2026-08-18T13:44:20.068314Z","submitted_at":"2024-09-09T22:29:35Z","title":"ClarQ-LLM: A Benchmark for Models Clarifying and Requesting Information in Task-Oriented Dialog","version":2},"cited_work":{"arxiv_id":"2409.06097","doi":"10.48550/arxiv.2409.06097","metadata_source":"arxiv_reference","pith_arxiv_id":"2409.06097","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"doi:10.48550/arXiv.2409.06097 , abstract =","venue":"arXiv (Cornell University)","work_id":"5792f841-3f4b-4185-8200-a4ea06414633","year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2409.06097","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:eee1a35eecb3ec711c934a75a93780327013b10984631ae01250e078678ac4c6","observation_id":"f3565ad3-3a90-4818-9ead-bb7d061061c5","resolution":{"observed_at":"2026-06-29T16:53:40.571485Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-07-14T18:20:23.148615+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-14T18:20:23.148615+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.12063","last_updated":"2024-06-01T07:35:26Z","snapshot_observed_at":"2026-08-19T09:50:10.885772Z","submitted_at":"2024-05-20T14:34:01Z","title":"CLAMBER: A Benchmark of Identifying and Clarifying Ambiguous Information Needs in Large Language Models","version":2},"cited_work":{"arxiv_id":"2405.12063","doi":"10.48550/arxiv.2405.12063","metadata_source":"arxiv_reference","pith_arxiv_id":"2405.12063","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Clamber: A benchmark of identifying and clarifying ambiguous information needs in large language models.arXiv preprint arXiv:2405.12063, 2024","venue":"arXiv (Cornell University)","work_id":"0b25b085-5cc5-419d-89c0-e29a547c1915","year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2405.12063","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:3cda8a355e37c021744f0169a7c9fe6b1035dbdb494f242e2fd5b67973f2ebcc","observation_id":"5bf061a3-ae5a-49a4-a5a3-8a72b6ff1300","resolution":{"observed_at":"2026-06-29T16:53:40.624562Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.13360","last_updated":"2026-04-24T18:23:38Z","snapshot_observed_at":"2026-08-14T13:31:11.360543Z","submitted_at":"2025-05-19T17:03:42Z","title":"What Prompts Don't Say: Understanding and Managing Underspecification in LLM Prompts","version":3},"cited_work":{"arxiv_id":"2505.13360","doi":"10.48550/arxiv.2505.13360","metadata_source":"pith","pith_arxiv_id":"2505.13360","snapshot_observed_at":"2026-08-05T02:49:54.815029Z","title":"What Prompts Don't Say: Understanding and Managing Underspecification in LLM Prompts","venue":"cs.CL","work_id":"a5bcb9f4-7fed-4035-8449-f6d670f94fd3","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2505.13360","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:11012cb64704d3860dfe20741b84b9fa7f3a9e5cba72b269e7799a4ca15ed76f","observation_id":"ffd0a8a4-9325-42c5-a6b3-241e681feb2b","resolution":{"observed_at":"2026-06-29T16:53:40.565926Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.04141","last_updated":"2025-05-29T08:04:32Z","snapshot_observed_at":"2026-08-15T00:09:02.045110Z","submitted_at":"2024-12-05T13:10:54Z","title":"Reducing Tool Hallucination via Reliability Alignment","version":3},"cited_work":{"arxiv_id":"2412.04141","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2412.04141","snapshot_observed_at":"2026-06-29T16:53:40.545243Z","title":"Reducing tool hallucination via reliability alignment.arXiv preprint arXiv:2412.04141","venue":null,"work_id":"7b953ade-f109-41e6-8c17-d81d91f61258","year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2412.04141","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:7cf8355a48f017b090dbc55c65910ba067f3f7c45d0b2e8575768f2c89721fde","observation_id":"8d103b0d-6cc0-4c65-b181-8adff8013200","resolution":{"observed_at":"2026-06-29T16:53:40.546847Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20015","last_updated":"2024-10-04T07:51:29Z","snapshot_observed_at":"2026-08-18T13:56:56.933519Z","submitted_at":"2024-06-28T16:03:30Z","title":"ToolBeHonest: A Multi-level Hallucination Diagnostic Benchmark for Tool-Augmented Large Language Models","version":2},"cited_work":{"arxiv_id":"2406.20015","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.20015","snapshot_observed_at":"2026-06-29T16:53:40.615942Z","title":"Toolbehonest: A multi-level hallucination diagnostic benchmark for tool-augmented large language models","venue":null,"work_id":"36ea12fe-9d01-4388-8fd2-170e860ca452","year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2406.20015","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:fc35eeb77315e09f57f28fb5f2f2df03a9b0c8d0b9dfaca10c72652e0f48a307","observation_id":"27a4c51f-cec4-441a-9980-a3141ab48d2a","resolution":{"observed_at":"2026-06-29T16:53:40.617632Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Toolscan: A benchmark for characterizing errors in tool-use llms","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:4984bf8a20a364df660b1a14be8520b126440b5813608a3886bcf53b9d9eb5b7","observation_id":"2777f43e-a485-4c3c-9fe9-c434921fa492","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.02691","last_updated":"2024-08-04T04:52:35Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-03-05T06:21:45Z","title":"InjecAgent: Benchmarking Indirect Prompt Injections in Tool-Integrated Large Language Model Agents","version":3},"cited_work":{"arxiv_id":"2403.02691","doi":"10.1145/3696410.3714756","metadata_source":"pith","pith_arxiv_id":"2403.02691","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"InjecAgent: Benchmarking Indirect Prompt Injections in Tool-Integrated Large Language Model Agents","venue":"cs.CL","work_id":"5cbfcda4-ec26-44e4-be60-e1525956d71d","year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2403.02691","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:6c352df10aa226387b2e9d0cb15e803e84c2f3941321d96024230a16db53e0b2","observation_id":"dd2e087d-bb31-4f2f-8954-711561c22ca2","resolution":{"observed_at":"2026-06-29T16:53:40.497612Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"From allies to adversaries: Manipulating LLM tool-calling through adversarial injection","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:fa5b2cc164669329cd27ca6d04710e65e668f07f7479dc7049fb91b5eb8fedef","observation_id":"ec6920e2-1edf-4ee6-b32f-4137fc337b09","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Zhu et al","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:370fbdff62169ea0f5c23e66d35804f6c284a88304a0f3cebba04e0700f2f5fd","observation_id":"de86aa06-aaca-4af3-b11d-762d5192a831","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.02502","last_updated":"2024-07-10T17:36:25Z","snapshot_observed_at":"2026-08-16T14:12:52.561838Z","submitted_at":"2024-03-04T21:50:29Z","title":"Trial and Error: Exploration-Based Trajectory Optimization for LLM Agents","version":2},"cited_work":{"arxiv_id":"2403.02502","doi":"10.48550/arxiv.2403.02502","metadata_source":"arxiv_reference","pith_arxiv_id":"2403.02502","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Trial and error: Exploration-based trajectory optimization for llm agents","venue":"arXiv (Cornell University)","work_id":"e9f4e923-14da-4180-9f84-15011e539433","year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2403.02502","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:c4a079f95dabc2e5c4c77926898eedcd8dea7a688724bb90557ae99c59ec8352","observation_id":"c280ebd1-32be-4256-a36c-38317269654f","resolution":{"observed_at":"2026-06-29T16:53:40.485548Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.00137","last_updated":"2025-02-28T19:27:29Z","snapshot_observed_at":"2026-08-16T12:54:09.521599Z","submitted_at":"2025-02-28T19:27:29Z","title":"SCORE: Systematic COnsistency and Robustness Evaluation for Large Language Models","version":1},"cited_work":{"arxiv_id":"2503.00137","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2503.00137","snapshot_observed_at":"2026-06-29T16:53:40.555918Z","title":"Score: Systematic consistency and robustness evaluation for large language models.arXiv preprint arXiv:2503.00137, 2025","venue":null,"work_id":"fcf5242a-e7ea-4057-bb19-77c9bc332a69","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2503.00137","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:ccb1d7b5b5e004a817ad61d2e06d4630c5d2317558b8be477c58efe1d82de3c3","observation_id":"571a2810-b364-4da3-985f-75420d482925","resolution":{"observed_at":"2026-06-29T16:53:40.557491Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:5368a82c76882cb513166d93cf3a50077f70f6741bfe2fe07a11cf45385aee2b","observation_id":"4dfcb310-b390-4f06-94f3-686ec368a30d","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2509.04013","last_updated":"2025-09-04T08:43:27Z","snapshot_observed_at":"2026-08-14T16:00:37.515648Z","submitted_at":"2025-09-04T08:43:27Z","title":"On Robustness and Reliability of Benchmark-Based Evaluation of LLMs","version":1},"cited_work":{"arxiv_id":"2509.04013","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2509.04013","snapshot_observed_at":"2026-06-30T23:45:07.998622Z","title":"On robustness and reliability of benchmark-based evaluation of llms","venue":null,"work_id":"b9028098-f624-4214-ba6e-82cd64cb2a14","year":2025},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"cited_paper":"/paper/2509.04013","citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:3dd4c2158a2d07561706d131f3bbb5ca9072e574a2f28d7b781caec8b56982db","observation_id":"b0a32282-5a01-4201-ade7-b59c79931a23","resolution":{"observed_at":"2026-06-29T16:53:40.568253Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T16:51:36.524194Z","title":"Yes, please return the Mechanical Keyboard and the Gaming Mouse","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments","version":1},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-06-29T16:51:36.524194Z"},"links":{"citing_paper":"/paper/2605.27209"},"observation_digest":"sha256:be628270edcc6c9d263e2508543b31461d7434932e49f22180ed73fc039042a2","observation_id":"35858390-1ae4-4e5a-b942-a47b1507e872","resolution":{"observed_at":"2026-06-29T16:51:36.524194Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2605.27209","last_updated":"2026-05-26T16:02:00Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-13T12:46:09.126090Z","submitted_at":"2026-05-26T16:02:00Z","title":"Learning to Act under Noise: Enhancing Agent Robustness via Noisy Environments"},"reference_resolution":{"displayed":82,"state_counts":{"malformed_identifier":2,"metadata_mismatch":5,"parse_uncertain":0,"unresolved":29,"verified_exact":46,"verified_fuzzy":0},"total_outbound_references":82},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"thesis":"As of 23 August 2026, this Paper Citation Record lists 82 of 82 outbound references and 0 inbound Pith citation observations for arXiv:2605.27209."}