{"as_of":"2026-08-05T03:00:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:3d6bd830a2bc0d20beffc574ad74eb8f789f8030eb0525230cf7d194110ca54f","coverage":[{"denominator":2,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-21T18:04:48.156727Z","state":"measured"},{"denominator":70,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":70,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-04T06:34:03.388597+00:00","state":"measured"},{"denominator":68,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":68,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-04T21:37:35.836340Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-04T17:20:00.018935Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2601.12538","last_updated":"2026-01-18T18:58:23Z","snapshot_observed_at":"2026-08-04T22:42:23.171653Z","submitted_at":"2026-01-18T18:58:23Z","title":"Agentic Reasoning for Large Language Models","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-17T15:14:25.558878Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2601.12538"},"observation_digest":"sha256:cb0f4dce26ec2d8ea03f2687520c0a002ed548bd62043666e602db7f6358d608","observation_id":"0837b2ba-1b6a-4925-a54d-0c4e790718f1","resolution":{"observed_at":"2026-05-17T15:14:25.765131Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-08-03T09:21:45.040455Z","title":"Evo-memory: Benchmarking llm agent test-time learning with self-evolving memory.arXiv preprint arXiv:2511.20857, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.14192","last_updated":"2026-07-03T11:26:58Z","snapshot_observed_at":"2026-08-03T09:21:30.130466Z","submitted_at":"2026-01-20T17:51:56Z","title":"Toward Efficient Agents: Memory, Tool learning, and Planning","version":2},"reference_index":145,"source":"pdf_text","source_observed_at":"2026-08-03T09:21:45.040455Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2601.14192"},"observation_digest":"sha256:75d10414ca1d723d0163b52b111a87a3a548c08f19e4f1331299a6f55e2f1df0","observation_id":"2d6e198c-3d5b-43e2-a5b9-11ac55bdbb70","resolution":{"observed_at":"2026-08-03T09:21:45.040455Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2602.06470","last_updated":"2026-06-17T03:51:53Z","snapshot_observed_at":"2026-08-03T04:00:15.683476Z","submitted_at":"2026-02-06T07:55:26Z","title":"Improve Large Language Model Systems with User Logs","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-21T14:34:01.332088Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2602.06470"},"observation_digest":"sha256:bf1c8611a2fc964d099458115585a615aa1b2f95164cab010f32b073489222f3","observation_id":"d2db6c0e-0fc4-41a3-9046-10b6cbba76dd","resolution":{"observed_at":"2026-05-21T14:34:12.932663Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-08-03T04:00:20.314092Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.06470","last_updated":"2026-06-17T03:51:53Z","snapshot_observed_at":"2026-08-03T04:00:15.683476Z","submitted_at":"2026-02-06T07:55:26Z","title":"Improve Large Language Model Systems with User Logs","version":3},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-03T04:00:20.314092Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2602.06470"},"observation_digest":"sha256:cf1996cbed1d221db6eff9ed19f8d140773e143163f0b5723a34d723361298b7","observation_id":"162a938e-68b1-45ea-a0c3-e5c58b3c8886","resolution":{"observed_at":"2026-08-03T04:00:20.314092Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2604.03295","last_updated":"2026-03-27T19:34:23Z","snapshot_observed_at":"2026-08-02T12:48:24.476609Z","submitted_at":"2026-03-27T19:34:23Z","title":"Scaling Teams or Scaling Time? Memory Enabled Lifelong Learning in LLM Multi-Agent Systems","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-14T22:42:43.070265Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2604.03295"},"observation_digest":"sha256:a41d9124ecbcc548eda4e7a86da7e39b43ef736fb8a80d9ba52ac473ac1966d0","observation_id":"34861772-4e01-4d15-8107-7d542ab0be67","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2604.03512","last_updated":"2026-04-09T22:32:31Z","snapshot_observed_at":"2026-08-03T00:32:26.785394Z","submitted_at":"2026-04-03T23:19:11Z","title":"ActionNex: A Virtual Outage Manager for Cloud Computing","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-13T19:18:39.881155Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2604.03512"},"observation_digest":"sha256:8914070ff4e70ac2e60d6d9ce12437f91ab52ef93e62378d1b445c2f4ab56aa2","observation_id":"85f52853-f057-4f59-937f-27754691a28a","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2604.07894","last_updated":"2026-04-09T07:04:19Z","snapshot_observed_at":"2026-08-04T00:00:08.854253Z","submitted_at":"2026-04-09T07:04:19Z","title":"TSUBASA: Improving Long-Horizon Personalization via Evolving Memory and Self-Learning with Context Distillation","version":1},"reference_index":72,"source":"arxiv_source","source_observed_at":"2026-05-10T16:48:53.092395Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2604.07894"},"observation_digest":"sha256:bf316c4e3a38e32d20bd92f169f9d731305b2ab16b76a4b03285471bf595f89f","observation_id":"901c90a3-fe5f-49fd-9bba-0da567080a59","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2604.08216","last_updated":"2026-05-18T11:20:22Z","snapshot_observed_at":"2026-07-31T19:05:02.577530Z","submitted_at":"2026-04-09T13:13:53Z","title":"MemCoT: Test-Time Scaling through Memory-Driven Chain-of-Thought","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-10T17:46:59.581486Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2604.08216"},"observation_digest":"sha256:6d7278b8a7565c18ae29fd802e26301dd124fd06d63783ddfdda6f3467dc6714","observation_id":"0f50ded2-146e-48dd-9d83-fdbbdfd0e440","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2604.08216","last_updated":"2026-05-18T11:20:22Z","snapshot_observed_at":"2026-07-31T19:05:02.577530Z","submitted_at":"2026-04-09T13:13:53Z","title":"MemCoT: Test-Time Scaling through Memory-Driven Chain-of-Thought","version":3},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-21T09:34:16.292323Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2604.08216"},"observation_digest":"sha256:a93eb2b0c78e61fdd667242e3a711611370b77362e773885a21d08e8211ff544","observation_id":"a22e10f1-281c-4305-b4ef-f8f02df92a0e","resolution":{"observed_at":"2026-05-21T09:34:57.214613Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2604.11811","last_updated":"2026-05-23T08:43:43Z","snapshot_observed_at":"2026-07-12T23:39:45.091388Z","submitted_at":"2026-04-10T03:22:26Z","title":"M$^\\star$: Every Task Deserves Its Own Memory Harness","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-10T17:40:12.368228Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2604.11811"},"observation_digest":"sha256:644cd6c687e0aa80061a2040ebbb707f2f74e789d99619c21dd3c6f0d5eb1461","observation_id":"a8af0925-a521-4c93-9d65-d64adcad0345","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2604.14475","last_updated":"2026-04-15T23:12:02Z","snapshot_observed_at":"2026-07-06T23:02:13.885476Z","submitted_at":"2026-04-15T23:12:02Z","title":"Evo-MedAgent: Beyond One-Shot Diagnosis with Agents That Remember, Reflect, and Improve","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-10T12:35:32.915685Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2604.14475"},"observation_digest":"sha256:f0fe7b84d21484c6d21554aa7ed33d6133cb3431ae579fc0c4157a7c72766633","observation_id":"bccacb77-9bac-48a5-8127-bc4143e7a00d","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2604.15097","last_updated":"2026-06-02T07:26:24Z","snapshot_observed_at":"2026-07-12T19:51:47.153234Z","submitted_at":"2026-04-16T14:55:49Z","title":"From Procedural Skills to Strategy Genes: Towards Experience-Driven Test-Time Evolution","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-10T10:52:47.760627Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2604.15097"},"observation_digest":"sha256:57e7ad5558010f808da3ea9adb87ac4ec62db6862bc216bc620c852bd45e7430","observation_id":"4e900222-1b84-4215-827a-f08631de3ea6","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-12T19:51:48.531280Z","title":"Chi, Chi Wang, Shuo Chen, Fernando Pereira, Wang-Cheng Kang, and Derek Zhiyuan Cheng","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2604.15097","last_updated":"2026-06-02T07:26:24Z","snapshot_observed_at":"2026-07-12T19:51:47.153234Z","submitted_at":"2026-04-16T14:55:49Z","title":"From Procedural Skills to Strategy Genes: Towards Experience-Driven Test-Time Evolution","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-07-12T19:51:48.531280Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2604.15097"},"observation_digest":"sha256:e2016d8384d635388c934ee16f7ee07e85642cb54d3fa13a5faad9b2a3a1074b","observation_id":"e60cb321-02d7-44c4-8bfd-a5351ade78fb","resolution":{"observed_at":"2026-07-12T19:51:48.531280Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2604.17308","last_updated":"2026-04-19T07:51:46Z","snapshot_observed_at":"2026-07-06T23:04:28.260463Z","submitted_at":"2026-04-19T07:51:46Z","title":"SkillFlow:Benchmarking Lifelong Skill Discovery and Evolution for Autonomous Agents","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-10T06:13:32.434201Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2604.17308"},"observation_digest":"sha256:5d26b93984f46f05989a62da3436eefbc72dd96c8fcfb66f1ab58ebbe1dae346","observation_id":"542a4428-6f28-43e9-9a91-f6504d871da7","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2604.26805","last_updated":"2026-05-10T15:46:25Z","snapshot_observed_at":"2026-08-03T01:00:46.930498Z","submitted_at":"2026-04-29T15:35:01Z","title":"Bian Que: An Agentic Framework with Flexible Skill Arrangement for Online System Operations","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-05-07T11:17:46.692409Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2604.26805"},"observation_digest":"sha256:de5d1b21144e7bc48e4d9e3e26a248a997ac2d090cfdd9e8584ab0a1372133d8","observation_id":"96ade66f-415d-475f-a0ca-790d84caf1c1","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2604.26805","last_updated":"2026-05-10T15:46:25Z","snapshot_observed_at":"2026-08-03T01:00:46.930498Z","submitted_at":"2026-04-29T15:35:01Z","title":"Bian Que: An Agentic Framework with Flexible Skill Arrangement for Online System Operations","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-05-12T01:50:45.507172Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2604.26805"},"observation_digest":"sha256:72b7e75edcd451a268348aa5d7b8356647518576617d8a722ccfa25e903c4ef6","observation_id":"d54c7b03-1909-474a-be23-6c38d24dff77","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.03808","last_updated":"2026-05-05T14:35:47Z","snapshot_observed_at":"2026-07-06T23:16:40.201701Z","submitted_at":"2026-05-05T14:35:47Z","title":"Agentic-imodels: Evolving agentic interpretability tools via autoresearch","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-05-07T16:37:43.371592Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.03808"},"observation_digest":"sha256:73058a84f25094aca6e78347b0b73b4363c1ab2527413e88dd1845c8fe61b64d","observation_id":"1abc07e5-e386-4813-b3bf-1c4346d8eb2e","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.06130","last_updated":"2026-05-12T10:25:23Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-05-07T12:33:30Z","title":"Skill1: Unified Evolution of Skill-Augmented Agents via Reinforcement Learning","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-05-08T10:23:52.522238Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.06130"},"observation_digest":"sha256:d46445bea08677408dc56bc1ac8e20b70cf0c24e9c17bb2b65d99e7d599e0b6a","observation_id":"6b35f8de-26d3-4610-85c0-9623e6a8232d","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.06130","last_updated":"2026-05-12T10:25:23Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-05-07T12:33:30Z","title":"Skill1: Unified Evolution of Skill-Augmented Agents via Reinforcement Learning","version":2},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-05-11T02:00:00.663355Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.06130"},"observation_digest":"sha256:e694b19676756fc5b5c5a6a565ea0b131edc4ae8bfcbc2b7a87a4f23b5188b48","observation_id":"e9173142-86b6-423b-b790-2cd94e31e416","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.06130","last_updated":"2026-05-12T10:25:23Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-05-07T12:33:30Z","title":"Skill1: Unified Evolution of Skill-Augmented Agents via Reinforcement Learning","version":3},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-05-13T07:17:13.708752Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.06130"},"observation_digest":"sha256:e65f29385ab86658c543ad5f6a33a622550b6c7e536443a1cd55e7cf7344554e","observation_id":"27e3c272-0d5b-46f2-a228-c86c589adea4","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.06365","last_updated":"2026-05-07T14:39:37Z","snapshot_observed_at":"2026-07-06T23:18:51.157048Z","submitted_at":"2026-05-07T14:39:37Z","title":"From Agent Loops to Deterministic Graphs: Execution Lineage for Reproducible AI-Native Work","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-05-08T09:50:30.639962Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.06365"},"observation_digest":"sha256:1ac80fd1de63e0ec8f4fc56c1a305ef387a8307ac4636c00deefa694bfdf5fc5","observation_id":"046f4948-05dc-4eeb-b9cf-00d648ced2ba","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.07358","last_updated":"2026-05-26T05:21:04Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:10:26Z","title":"A Comprehensive Survey on Agent Skills: Taxonomy, Techniques, and Applications","version":1},"reference_index":146,"source":"pdf_text","source_observed_at":"2026-05-11T01:47:39.926540Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.07358"},"observation_digest":"sha256:d8d904a6a842d4733fc27c02aa2e5e706b7e2ff6ea9ea182e013e4cb002815c6","observation_id":"ec3adc7a-e507-48a1-9c81-c7c9d35a3d10","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.07358","last_updated":"2026-05-26T05:21:04Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:10:26Z","title":"A Comprehensive Survey on Agent Skills: Taxonomy, Techniques, and Applications","version":2},"reference_index":148,"source":"pdf_text","source_observed_at":"2026-05-20T23:15:44.550045Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.07358"},"observation_digest":"sha256:4a27e42ba59bb57e9424804d3521f5d3faf92b2a015ca5926079aa83d567252d","observation_id":"35c27b98-50d0-43c0-825f-d6c8348da9bc","resolution":{"observed_at":"2026-05-20T23:19:14.814754Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.07358","last_updated":"2026-05-26T05:21:04Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:10:26Z","title":"A Comprehensive Survey on Agent Skills: Taxonomy, Techniques, and Applications","version":3},"reference_index":140,"source":"pdf_text","source_observed_at":"2026-06-30T23:23:42.883286Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.07358"},"observation_digest":"sha256:9ae2d1aac3990f4d804d22525c71d64ca38060b1abc2f6c2f4561d0356524392","observation_id":"e87966a1-9cdc-4f5e-8d91-a6102df020e7","resolution":{"observed_at":"2026-06-30T23:25:07.289996Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.07594","last_updated":"2026-05-14T02:47:01Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-05-08T11:07:04Z","title":"MemCompiler: Compile, Don't Inject -- State-Conditioned Memory for Embodied Agents","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-11T02:36:53.775557Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.07594"},"observation_digest":"sha256:3fcc2c87d8d2ebb108b3f1a23d4bf8bc377b96a2952d1af987d141c0c3e3db5e","observation_id":"22175fa4-fb63-4439-a326-c859a44c11fd","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.07594","last_updated":"2026-05-14T02:47:01Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-05-08T11:07:04Z","title":"MemCompiler: Compile, Don't Inject -- State-Conditioned Memory for Embodied Agents","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-15T05:57:05.259111Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.07594"},"observation_digest":"sha256:b3375676bdcbe9ec2e78ceaa560ddcc6e0e1f2360f460dbcf013a4cba371b9be","observation_id":"d3c3b01d-8889-4961-abcf-0aaf5839706f","resolution":{"observed_at":"2026-05-15T05:59:48.997096Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.08704","last_updated":"2026-06-26T05:42:46Z","snapshot_observed_at":"2026-07-06T23:20:52.521341Z","submitted_at":"2026-05-09T05:38:21Z","title":"AgentPSO: Evolving Agent Reasoning Skill via Multi-agent Particle Swarm Optimization","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-05-12T00:55:52.965890Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.08704"},"observation_digest":"sha256:00a24792c015d678d04a88ecf24d8ddd2501bf51c8055d3041b9fcedd128b4fb","observation_id":"545925b9-4c11-4566-a109-7185f2096439","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.08704","last_updated":"2026-06-26T05:42:46Z","snapshot_observed_at":"2026-07-06T23:20:52.521341Z","submitted_at":"2026-05-09T05:38:21Z","title":"AgentPSO: Evolving Agent Reasoning Skill via Multi-agent Particle Swarm Optimization","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-06-30T23:29:12.545347Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.08704"},"observation_digest":"sha256:b4236007695fe65133172dc7808cce224c0f435e53a8dc9def7c0aaed0827367","observation_id":"aaccb7a4-5374-4895-88b7-e2c25f86751a","resolution":{"observed_at":"2026-06-30T23:35:07.629183Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.09315","last_updated":"2026-05-10T04:20:24Z","snapshot_observed_at":"2026-07-06T23:21:26.017420Z","submitted_at":"2026-05-10T04:20:24Z","title":"Do Self-Evolving Agents Forget? Capability Degradation and Preservation in Lifelong LLM Agent Adaptation","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-12T04:10:31.784413Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.09315"},"observation_digest":"sha256:338439a7dd78443d9c1665273504d4be7425d7788b8f891152a178474752c1c9","observation_id":"f5c80c84-4e9d-4b40-ab64-e0496befc17b","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.10064","last_updated":"2026-05-11T06:39:51Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-05-11T06:39:51Z","title":"MAGE: Multi-Agent Self-Evolution with Co-Evolutionary Knowledge Graphs","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-12T05:13:28.089038Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.10064"},"observation_digest":"sha256:56d8d9a7eb34be77dceb969eb039c85016aa7ba39a4b7bf725224841fe55b70a","observation_id":"ad4452eb-2f58-4c98-97ac-ad8c2ad2c128","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.10268","last_updated":"2026-05-11T09:30:59Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-05-11T09:30:59Z","title":"MemReread: Enhancing Agentic Long-Context Reasoning via Memory-Guided Rereading","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-12T05:22:43.330891Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.10268"},"observation_digest":"sha256:cf373dc21853c7d9a744af2aab85f5f20d0468204e1b589e8eccda43070c2570","observation_id":"b18eb268-a1be-4705-9d75-605a53e8741f","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.10913","last_updated":"2026-06-24T17:55:07Z","snapshot_observed_at":"2026-08-02T06:55:55.998448Z","submitted_at":"2026-05-11T17:50:51Z","title":"Shepherd: Enabling Programmable Meta-Agents via Reversible Agentic Execution Traces","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-05-12T03:29:33.497561Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.10913"},"observation_digest":"sha256:1dd44e50b588a47fbb462b764e6e919ab15228f67e4b23755e2d623c43ea7d55","observation_id":"ac261d9c-b950-4c1c-8f05-92fc46e95974","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.10913","last_updated":"2026-06-24T17:55:07Z","snapshot_observed_at":"2026-08-02T06:55:55.998448Z","submitted_at":"2026-05-11T17:50:51Z","title":"Shepherd: Enabling Programmable Meta-Agents via Reversible Agentic Execution Traces","version":3},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-06-30T22:24:26.528532Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.10913"},"observation_digest":"sha256:1ce203d5b9b23243ff1f0bdb18c6165c13f3248b84f1c694bafee8dd725b0d5a","observation_id":"298d372a-b8a5-44c8-be60-66a7d322fcb7","resolution":{"observed_at":"2026-06-30T22:25:06.845608Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.13050","last_updated":"2026-05-14T12:47:52Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-05-13T06:15:32Z","title":"Context Training with Active Information Seeking","version":1},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-05-14T20:03:01.429424Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.13050"},"observation_digest":"sha256:861fea4a87e8441f72dd28f7532b6ca23af47e4e5609411a4b80722d106a7b44","observation_id":"7731b378-3142-4b4c-bce8-e214c85c5e38","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.13050","last_updated":"2026-05-14T12:47:52Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-05-13T06:15:32Z","title":"Context Training with Active Information Seeking","version":2},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-05-15T06:07:10.805180Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.13050"},"observation_digest":"sha256:324c00435a325766ffce3b8c46fb05d93eb51758a4c8f4149b8f817b0506a616","observation_id":"cf3aad9b-f80b-4de0-a85f-91308d2ec64d","resolution":{"observed_at":"2026-05-15T06:09:49.991157Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.13542","last_updated":"2026-05-13T13:52:42Z","snapshot_observed_at":"2026-07-06T23:25:07.022065Z","submitted_at":"2026-05-13T13:52:42Z","title":"RealICU: Do LLM Agents Understand Long-Context ICU Data? A Benchmark Beyond Behavior Imitation","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-05-14T18:47:46.239810Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.13542"},"observation_digest":"sha256:3221d4ca006cf960f9eef27e80916bd350bff929c01fdec538d8eae8b1688132","observation_id":"6af1ea49-8558-49c3-a9c3-a49393cbfb23","resolution":{"observed_at":"2026-05-14T23:13:16.650386Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.13941","last_updated":"2026-05-13T17:12:44Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-05-13T17:12:44Z","title":"EvolveMem:Self-Evolving Memory Architecture via AutoResearch for LLM Agents","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-15T04:48:12.883761Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.13941"},"observation_digest":"sha256:2b19ad97cc33945f5453c323a4beea7bbb0d86226286221b790db3d816f71eb4","observation_id":"c19c32c9-dcc7-44d4-80e2-40fa9ea126dd","resolution":{"observed_at":"2026-05-15T04:49:44.079669Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.14477","last_updated":"2026-07-14T21:30:34Z","snapshot_observed_at":"2026-08-02T14:06:39.249902Z","submitted_at":"2026-05-14T07:18:12Z","title":"Test-Time Learning with an Evolving Library","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-15T01:44:16.865595Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.14477"},"observation_digest":"sha256:ab050504e1c77146542f4b57a59bb3ce711f50e87dcd5c71ca20933a1f2688e8","observation_id":"8ec255aa-ff9b-492d-a17e-5345aad3104f","resolution":{"observed_at":"2026-05-15T01:48:28.969956Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-08-02T14:06:40.867144Z","title":"Evo-memory: Benchmarking llm agent test-time learning with self-evolving memory.arXiv preprint arXiv:2511.20857, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.14477","last_updated":"2026-07-14T21:30:34Z","snapshot_observed_at":"2026-08-02T14:06:39.249902Z","submitted_at":"2026-05-14T07:18:12Z","title":"Test-Time Learning with an Evolving Library","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-02T14:06:40.867144Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.14477"},"observation_digest":"sha256:a210e0394dd207140bfaf02a80f779f5587df199403d520aecfa217195cac39d","observation_id":"c86321c1-05b6-446c-b9e4-48b35335e3cb","resolution":{"observed_at":"2026-08-02T14:06:40.867144Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.15384","last_updated":"2026-05-14T20:15:22Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-05-14T20:15:22Z","title":"Is One Score Enough? Rethinking the Evaluation of Sequentially Evolving LLM Memory","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-19T16:43:37.472644Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.15384"},"observation_digest":"sha256:6eb0cabcd5edf3d3082e52b61d2e564d71a5449c32411b047c2d9f7134e4a0a9","observation_id":"1312ebbd-7bd8-470b-86a2-ce2c58c2b570","resolution":{"observed_at":"2026-05-19T16:47:40.489988Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.17721","last_updated":"2026-05-18T00:50:23Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-05-18T00:50:23Z","title":"EXG: Self-Evolving Agents with Experience Graphs","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-19T22:17:10.728289Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.17721"},"observation_digest":"sha256:aa7c7b5d3ff073a879cd89802908c802e8b917c8a01f0682fa1a50c483896c2a","observation_id":"80b0efc5-e7bb-45d7-91c6-bebcd39af8ac","resolution":{"observed_at":"2026-05-19T22:17:49.338215Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.18421","last_updated":"2026-05-18T13:54:38Z","snapshot_observed_at":"2026-08-02T07:00:41.376916Z","submitted_at":"2026-05-18T13:54:38Z","title":"EvoMemBench: Benchmarking Agent Memory from a Self-Evolving Perspective","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-20T11:35:23.892275Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.18421"},"observation_digest":"sha256:7cc941b218c153cea0b72c0fa662aa6beaab9183940808497c44fce1379b57bb","observation_id":"88f441a7-c589-4cc2-b067-fa1df2bb3a91","resolution":{"observed_at":"2026-05-20T11:38:14.541752Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.18747","last_updated":"2026-05-18T17:59:03Z","snapshot_observed_at":"2026-07-06T23:29:33.702647Z","submitted_at":"2026-05-18T17:59:03Z","title":"Code as Agent Harness","version":1},"reference_index":202,"source":"pdf_text","source_observed_at":"2026-05-20T10:54:54.558241Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.18747"},"observation_digest":"sha256:b37943b750cea6b99b3905c1a0f6b826e1ecd95479b55d194d6d09f57f6c490d","observation_id":"5e0676b8-bc48-4c42-b361-08a1b64eb9a0","resolution":{"observed_at":"2026-05-20T10:58:14.364346Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.20616","last_updated":"2026-05-20T02:03:34Z","snapshot_observed_at":"2026-08-01T18:32:22.388727Z","submitted_at":"2026-05-20T02:03:34Z","title":"Auto-Dreamer: Learning Offline Memory Consolidation for Language Agents","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-21T05:36:56.151440Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.20616"},"observation_digest":"sha256:686fd2e83afd3ab88985307db9a470ef3751a3bbaaa4205a09595d5f6f59526b","observation_id":"b1eae535-c00f-4620-aeb5-00429443b89c","resolution":{"observed_at":"2026-05-21T05:39:40.764887Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2605.28773","last_updated":"2026-05-27T17:35:34Z","snapshot_observed_at":"2026-07-06T23:38:19.096349Z","submitted_at":"2026-05-27T17:35:34Z","title":"Rethinking Memory as Continuously Evolving Connectivity","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-06-29T12:25:50.220145Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2605.28773"},"observation_digest":"sha256:fc27bb4d936383c396962f86dd2b6a2cda5204886fc0ea90ff21abe1429ac571","observation_id":"0c82230f-c0e9-49ba-b087-1b26e3407495","resolution":{"observed_at":"2026-06-29T12:33:24.800094Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2606.01223","last_updated":"2026-05-31T13:16:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-05-31T13:16:02Z","title":"Connecting the Dots: Benchmarking Reflective Memory in Long-Horizon Dialogue","version":1},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-06-28T17:41:06.216002Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2606.01223"},"observation_digest":"sha256:6ecd2135bc072ecd08a773073eef6dcffeabb2f311cfbc3228a2d7315249163b","observation_id":"56be5de9-6e7a-4d25-95ff-a44afd8099a3","resolution":{"observed_at":"2026-06-28T17:42:25.950848Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2606.02461","last_updated":"2026-06-02T03:07:54Z","snapshot_observed_at":"2026-08-03T22:24:41.920871Z","submitted_at":"2026-06-01T16:32:59Z","title":"AgentCL: Toward Rigorous Evaluation of Continual Learning in Language Agents","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-06-28T14:52:49.244423Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2606.02461"},"observation_digest":"sha256:66fff28e778086f9fbccaa0f1915e7588a003501af8e5c5bd9527d5550eebcff","observation_id":"f6b69651-616d-4605-aebd-d29360283f78","resolution":{"observed_at":"2026-07-01T22:56:20.164476Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2606.05008","last_updated":"2026-06-03T15:28:57Z","snapshot_observed_at":"2026-07-06T23:45:02.342193Z","submitted_at":"2026-06-03T15:28:57Z","title":"M$^3$Eval: Multi-Modal Memory Evaluation through Cognitively-Grounded Video Tasks","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-06-28T06:16:07.090870Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2606.05008"},"observation_digest":"sha256:a7e7e7cffdc92a048a64ba1d020104163135671dcfb40720dea427e83345881c","observation_id":"18914754-6459-4567-b3c1-236c4e069731","resolution":{"observed_at":"2026-07-02T08:16:47.771817Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2606.05513","last_updated":"2026-06-03T23:40:30Z","snapshot_observed_at":"2026-07-06T23:45:28.379468Z","submitted_at":"2026-06-03T23:40:30Z","title":"EpiEvolve: Self-Evolving Agents for Streaming Pandemic Forecasting under Regime Shifts","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-06-28T05:40:18.458584Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2606.05513"},"observation_digest":"sha256:9a5200759ba2249ea0d85e1f884143a225a0653fbf907d4238205bcbd626598f","observation_id":"b134c7f7-175b-41f6-b55b-2fe0f6311855","resolution":{"observed_at":"2026-07-02T09:06:49.133300Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2606.06448","last_updated":"2026-06-04T17:44:18Z","snapshot_observed_at":"2026-08-02T14:46:30.773077Z","submitted_at":"2026-06-04T17:44:18Z","title":"Agent Memory: Characterization and System Implications of Stateful Long-Horizon Workloads","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-06-28T01:10:02.610738Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2606.06448"},"observation_digest":"sha256:161dd9a6286b8aadbc29d0ccb88bebbd3bd09634199c6fd4b10fee3ad840ed8e","observation_id":"bf09f45e-3139-427f-b4b9-31967b364d20","resolution":{"observed_at":"2026-07-02T13:36:59.295338Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2606.06960","last_updated":"2026-06-05T06:39:16Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-06-05T06:39:16Z","title":"Tree-of-Experience: A Structured Experience-Management Solution for Self-Evolving Agents under Low-Repetition and Implicit-Reward Environments","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-27T22:24:25.502732Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2606.06960"},"observation_digest":"sha256:2cb0a37eb55f1a434ad269754370f7d17cc28fe845da16a3625afb838ae92f93","observation_id":"1343f239-25ba-481f-b88f-1e4a8892f3c6","resolution":{"observed_at":"2026-07-02T16:47:09.749147Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2606.09365","last_updated":"2026-06-10T14:46:00Z","snapshot_observed_at":"2026-08-02T16:59:59.758751Z","submitted_at":"2026-06-08T11:37:01Z","title":"Experience Makes Skillful: Enabling Generalizable Medical Agent Reasoning via Self-Evolving Skill Memory","version":2},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-06-27T16:45:30.431403Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2606.09365"},"observation_digest":"sha256:dba5e3a9228b0f6219e68a6f34239c9a3fac2c9d92716b6441a3d71895b26426","observation_id":"326fa9a8-b64e-4147-a12e-8da722189dbb","resolution":{"observed_at":"2026-07-03T01:07:30.920151Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2606.20475","last_updated":"2026-06-18T16:54:25Z","snapshot_observed_at":"2026-07-06T23:55:38.979306Z","submitted_at":"2026-06-18T16:54:25Z","title":"Marginal Advantage Accumulation for Memory-Driven Agent Self-Evolution","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-06-26T17:45:49.272070Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2606.20475"},"observation_digest":"sha256:0c90278ebf6d263c213fba7755c9d9aa3405b6cb9b9259ea933a5a4f807b1195","observation_id":"fcd865db-df08-447b-9b56-eaabc62d6fce","resolution":{"observed_at":"2026-07-04T03:39:30.783334Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2606.20625","last_updated":"2026-05-26T15:48:09Z","snapshot_observed_at":"2026-07-06T23:55:43.830896Z","submitted_at":"2026-05-26T15:48:09Z","title":"AlphaMemo: Structured Search-Process Memory for Self-Evolving Alpha Mining Agents","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-06-29T16:55:21.649886Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2606.20625"},"observation_digest":"sha256:5f5f17932b07a310275dad65853774d04786753889f1f9f8f96ace00e5f7a1d2","observation_id":"4c24f545-ca42-47ca-8bd6-d0803abc7a13","resolution":{"observed_at":"2026-06-29T17:03:41.262391Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2606.22385","last_updated":"2026-06-21T08:22:23Z","snapshot_observed_at":"2026-08-01T18:23:09.613202Z","submitted_at":"2026-06-21T08:22:23Z","title":"MetaPS: Adaptive Programmatic Strategy Selection for Market Agents","version":1},"reference_index":132,"source":"arxiv_source","source_observed_at":"2026-06-26T11:06:28.690956Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2606.22385"},"observation_digest":"sha256:bd1cfb243cdfb4497abc29ace6fd61bc1c8ddb47d4bf2bd9791508f962a3e51a","observation_id":"d0de9262-ab7b-4fbf-ac70-6d8f9738b9e0","resolution":{"observed_at":"2026-07-04T08:39:42.650845Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2606.24428","last_updated":"2026-06-23T11:05:05Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-06-23T11:05:05Z","title":"Escaping the Self-Confirmation Trap: An Execute-Distill-Verify Paradigm for Agentic Experience Learning","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-06-25T23:49:38.932474Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2606.24428"},"observation_digest":"sha256:403f21ebb80563b35be8c19ccc1498f0d63c72b08fdf3a01a7278be19a0643e9","observation_id":"3847d528-82cd-482e-9e54-6b382803f934","resolution":{"observed_at":"2026-07-04T17:20:00.020604Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2606.31121","last_updated":"2026-06-30T04:33:18Z","snapshot_observed_at":"2026-08-04T05:42:21.520626Z","submitted_at":"2026-06-30T04:33:18Z","title":"The Past Is Prologue: A Plug-in Controller for Selective Updates in Sequentially Evolving LLM Memory","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-07-01T06:00:11.022368Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2606.31121"},"observation_digest":"sha256:bf4f29d972277c0018f64441be8f31640210a7b5b9a097090b26e486b22f03ce","observation_id":"8ec8cff0-f92f-459e-8025-d5ac0a39c4b2","resolution":{"observed_at":"2026-07-01T09:55:41.341437Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2606.31612","last_updated":"2026-07-02T06:42:39Z","snapshot_observed_at":"2026-08-02T11:25:34.061644Z","submitted_at":"2026-06-30T13:01:19Z","title":"What Memory Do GUI Agents Really Need? From Passive Records to Active Task-Driving States","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-07-01T05:26:53.441833Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2606.31612"},"observation_digest":"sha256:5657ed19e54df69652d0e836aea7a80651ab6cf694b337bf0a3cc1dc83c8192e","observation_id":"e4564c04-6bfd-40cc-8caf-5df5fbee1fad","resolution":{"observed_at":"2026-07-01T10:35:41.276404Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":"2511.20857","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-04T17:20:00.018935Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","venue":"cs.CL","work_id":"eafad83d-679f-4020-8413-9451109c58ba","year":2025},"citing_paper":{"arxiv_id":"2606.31612","last_updated":"2026-07-02T06:42:39Z","snapshot_observed_at":"2026-08-02T11:25:34.061644Z","submitted_at":"2026-06-30T13:01:19Z","title":"What Memory Do GUI Agents Really Need? From Passive Records to Active Task-Driving States","version":2},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-07-03T22:03:26.625209Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2606.31612"},"observation_digest":"sha256:1ee82f6c4e2869be8859597c16036f72b44db60faf73bbade8f16b13243fb3cf","observation_id":"657e666d-7df1-4c95-a97b-fc8a2eced198","resolution":{"observed_at":"2026-07-03T22:08:58.667321Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-11T07:57:43.000834Z","title":"Chi, et al","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.05155","last_updated":"2026-07-06T14:39:22Z","snapshot_observed_at":"2026-07-11T07:57:39.948830Z","submitted_at":"2026-07-06T14:39:22Z","title":"EdgeBench: Unveiling Scaling Laws of Learning from Real-World Environments","version":1},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-07-11T07:57:43.000834Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2607.05155"},"observation_digest":"sha256:5d62ce9b1ad83ae357098e8c3b6353dfd11cf200c452a9db41031c0b0219e864","observation_id":"252f20b3-6da8-46dd-b279-166d986c5e88","resolution":{"observed_at":"2026-07-11T07:57:43.000834Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-14T12:26:27.446079Z","title":"Chi, et al","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.10350","last_updated":"2026-07-17T17:27:03Z","snapshot_observed_at":"2026-08-02T07:21:21.791047Z","submitted_at":"2026-07-11T15:24:43Z","title":"ABot-AgentOS: A General Robotic Agent OS with Lifelong Multi-modal Memory","version":1},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-07-14T12:26:27.446079Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2607.10350"},"observation_digest":"sha256:544810cb93896561468d6a23f85c2f68e7898883fcd99b6e6cd4d04ce8e49454","observation_id":"f80d9411-59af-440a-8a82-80771dacb623","resolution":{"observed_at":"2026-07-14T12:26:27.446079Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-08-02T07:21:33.187454Z","title":"Chi, et al","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.10350","last_updated":"2026-07-17T17:27:03Z","snapshot_observed_at":"2026-08-02T07:21:21.791047Z","submitted_at":"2026-07-11T15:24:43Z","title":"ABot-AgentOS: A General Robotic Agent OS with Lifelong Multi-modal Memory","version":3},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-08-02T07:21:33.187454Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2607.10350"},"observation_digest":"sha256:c07b4b37168224076692bb9e8d05c469afb90f56b51a8c88ccb51d780165389c","observation_id":"5fb71451-e324-4a06-aa00-80c198d4a3ea","resolution":{"observed_at":"2026-08-02T07:21:33.187454Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-14T11:04:44.592375Z","title":"arXiv preprint arXiv:2511.20857 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10526","last_updated":"2026-07-27T05:08:50Z","snapshot_observed_at":"2026-08-02T13:09:12.564576Z","submitted_at":"2026-07-12T01:08:59Z","title":"Agents Don't Just Agree, They Remember: Benchmarking Persistent Sycophancy in Stateful Personal Agents","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-07-14T11:04:44.592375Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2607.10526"},"observation_digest":"sha256:7fde2b9ab25145b9ec6aaac82ddeb6c24fe793dd1e2810d2f1480153e7e217b9","observation_id":"1a12b528-c9c4-43c6-8ab8-0e1627436779","resolution":{"observed_at":"2026-07-14T11:04:44.592375Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-08-01T22:44:00.259964Z","title":"Evo-memory: Benchmarking llm agent test-time learning with self-evolving memory.arXiv preprint arXiv:2511.20857, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.15657","last_updated":"2026-07-17T06:05:17Z","snapshot_observed_at":"2026-08-01T22:43:48.102799Z","submitted_at":"2026-07-17T06:05:17Z","title":"Do Agents Dream of False Memories? Black-box Visual Attacks on Long-term Memory in Multimodal AI Agents","version":1},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-01T22:44:00.259964Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2607.15657"},"observation_digest":"sha256:f97d7261572940e7c486ecd425ce780b473c13b4abfb850f3b3499369780575a","observation_id":"483f00e8-5f31-4023-a97f-b7c8743336dc","resolution":{"observed_at":"2026-08-01T22:44:00.259964Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-08-01T11:46:02.563189Z","title":"Chi, Chi Wang, Shuo Chen, Fernando Pereira, Wang-Cheng Kang, and Derek Zhiyuan Cheng","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.26643","last_updated":"2026-07-29T09:05:40Z","snapshot_observed_at":"2026-08-04T07:42:56.265981Z","submitted_at":"2026-07-29T09:05:40Z","title":"Rethinking Self-Evolution: A Constrained Exploration-Exploitation Process for Mitigating Skill Overfitting","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-01T11:46:02.563189Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2607.26643"},"observation_digest":"sha256:eb096cde53ff0839dd0f9acf7603155a8fd9ea5811a82b5521d9983397f7d7dc","observation_id":"67342421-022f-45fd-859c-8120df1aba13","resolution":{"observed_at":"2026-08-01T11:46:02.563189Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-07-31T16:09:14.331637Z","title":"H.; Wang, C.; Chen, S.; Pereira, F.; Kang, W.-C.; and Cheng, D","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28156","last_updated":"2026-07-30T12:58:57Z","snapshot_observed_at":"2026-08-03T00:05:15.219263Z","submitted_at":"2026-07-30T12:58:57Z","title":"RRM: Experience-Driven Reflective Retrieval Memory for Long-Horizon Multimodal Reasoning","version":1},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-07-31T16:09:14.331637Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2607.28156"},"observation_digest":"sha256:a1baa8d1c84b1b6282086a6aa7475de38ee3ce5f91bc3b223ce397d9c348a54c","observation_id":"7e28cc05-454d-47e8-acc0-62ed074bc7e3","resolution":{"observed_at":"2026-07-31T16:09:14.331637Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-08-04T01:12:32.201860Z","title":"Chi, Chi Wang, Shuo Chen, Fernando Pereira, Wang-Cheng Kang, and Derek Zhiyuan Cheng","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.00155","last_updated":"2026-07-31T17:50:08Z","snapshot_observed_at":"2026-08-05T02:53:01.909525Z","submitted_at":"2026-07-31T17:50:08Z","title":"AgentStream: How Well Do Self-Evolving LLM Agents Perform Under Streaming Tasks?","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-04T01:12:32.201860Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2608.00155"},"observation_digest":"sha256:5f530277042f00428f9879b026d5238748923311e0f22c96329f5dcc6f55fc8b","observation_id":"19de0ee3-c2cd-4dac-84f9-1710b859a36c","resolution":{"observed_at":"2026-08-04T01:12:32.201860Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-08-04T21:37:35.836340Z","title":"arXiv preprint arXiv:2511.20857 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.01759","last_updated":"2026-08-03T06:28:21Z","snapshot_observed_at":"2026-08-05T02:44:05.059909Z","submitted_at":"2026-08-03T06:28:21Z","title":"Benign Alone, Harmful Together: Exploiting Experience Composition in Self-Evolving LLM Agents","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-04T21:37:35.836340Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2608.01759"},"observation_digest":"sha256:77756cc2c412d16145ae1ff689d99bfff2585641d6a17222caccc94dbcf679a3","observation_id":"b3809490-a219-4107-ae23-10d6df1a983b","resolution":{"observed_at":"2026-08-04T21:37:35.836340Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2511.20857/citation-record","integrity":"/paper/2511.20857/integrity","json":"/paper/2511.20857/citation-record.json","paper":"/paper/2511.20857"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2402.17753","last_updated":"2024-02-27T18:42:31Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-02-27T18:42:31Z","title":"Evaluating Very Long-Term Conversational Memory of LLM Agents","version":1},"cited_work":{"arxiv_id":"2402.17753","doi":"10.18653/v1/2024.acl-long.747arxiv:2402.17753","metadata_source":"pith","pith_arxiv_id":"2402.17753","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Evaluating Very Long-Term Conversational Memory of LLM Agents","venue":"cs.CL","work_id":"2d8c9fb7-9ace-4925-8fe1-f4e44625d04c","year":2024},"citing_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-21T18:04:48.156727Z"},"links":{"cited_paper":"/paper/2402.17753","citing_paper":"/paper/2511.20857"},"observation_digest":"sha256:c089369dd9757602ab8a1b4101d81ad8bc25de88a69339b6180ea47653df35df","observation_id":"7e2fa6bd-8d0d-4e51-b189-993c0edfe3f7","resolution":{"observed_at":"2026-05-21T18:05:27.147044Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"1,3” or “2-4","venue":null,"work_id":"05ab3b4e-5e34-4b35-8567-9d7379ebadc2","year":2022},"citing_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-21T18:04:48.156727Z"},"links":{"citing_paper":"/paper/2511.20857"},"observation_digest":"sha256:f52fbb84bf8c82ce11e8d2e31ef14126048d5a7ea9b2e13b1370a4dde4f5289d","observation_id":"799b9935-80e0-4020-b1a5-f045028d9abb","resolution":{"observed_at":"2026-05-21T18:05:27.248013Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory"},"reference_resolution":{"displayed":2,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":1,"verified_fuzzy":0},"total_outbound_references":2},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"thesis":"As of 5 August 2026, this Paper Citation Record lists 2 of 2 outbound references and 68 inbound Pith citation observations for arXiv:2511.20857."}