{"as_of":"2026-08-07T23:03:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:836a9f40e76e0eac8f6e131a82ab62620d176bc58978526bc9b85e1f8743bf9e","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":39,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":39,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":39,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":39,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T15:30:50.846655Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"pith","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":2,"observed_at":"2026-08-05T02:28:24.338817Z","source":"pith"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-07T15:30:50.846655Z","title":"Webwalker: Benchmarking llms in web traversal","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.14963","last_updated":"2025-05-20T22:42:33Z","snapshot_observed_at":"2026-08-07T15:24:14.180803Z","submitted_at":"2025-05-20T22:42:33Z","title":"MedBrowseComp: Benchmarking Medical Deep Research and Computer Use","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T15:30:50.846655Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2505.14963"},"observation_digest":"sha256:fad63da72848f68d9f89edf5afc8797e10a7781580e94a965d26e78d7cc2eaf2","observation_id":"db2af963-922e-497b-b305-d71e337269c6","resolution":{"observed_at":"2026-08-07T15:30:50.846655Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-07T15:06:08.136388Z","title":"Webwalker: Benchmarking llms in web traversal","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.16410","last_updated":"2025-05-22T09:00:19Z","snapshot_observed_at":"2026-08-07T20:42:11.938642Z","submitted_at":"2025-05-22T09:00:19Z","title":"Tool-Star: Empowering LLM-Brained Multi-Tool Reasoner via Reinforcement Learning","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-07T15:06:08.136388Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2505.16410"},"observation_digest":"sha256:3e9b1bbb69cea25c5064c13f99c78f87bafac9d48eff98cae3a03e6fda2af6cf","observation_id":"ce961658-b063-4860-8900-c03507313a72","resolution":{"observed_at":"2026-08-07T15:06:08.136388Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-07T14:51:16.385377Z","title":"Webwalker: Benchmarking llms in web traversal,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.17399","last_updated":"2025-05-26T11:15:36Z","snapshot_observed_at":"2026-08-07T14:45:38.156584Z","submitted_at":"2025-05-23T02:16:11Z","title":"FullFront: Benchmarking MLLMs Across the Full Front-End Engineering Workflow","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T14:51:16.385377Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2505.17399"},"observation_digest":"sha256:c65e887bc0620d934151e0a55572d04f80100c5ab58cfd5a85ca8f09488fce5d","observation_id":"4936694c-0f25-4f74-b84e-90d957e77f11","resolution":{"observed_at":"2026-08-07T14:51:16.385377Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-07T13:24:53.713826Z","title":"Webwalker: Benchmarking llms in web traversal","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.22019","last_updated":"2025-06-03T05:28:55Z","snapshot_observed_at":"2026-08-07T13:14:56.558609Z","submitted_at":"2025-05-28T06:30:51Z","title":"VRAG-RL: Empower Vision-Perception-Based RAG for Visually Rich Information Understanding via Iterative Reasoning with Reinforcement Learning","version":2},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-07T13:24:53.713826Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2505.22019"},"observation_digest":"sha256:2d5fd10de0631c17c0de9f51683a35cdcf4b0deb280e7d7b742d0cccdd66e8b7","observation_id":"f15666ee-bdbc-4691-9a14-5246d2f099a3","resolution":{"observed_at":"2026-08-07T13:24:53.713826Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-07T13:12:59.086994Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.22501","last_updated":"2025-05-28T15:50:48Z","snapshot_observed_at":"2026-08-07T13:03:18.495419Z","submitted_at":"2025-05-28T15:50:48Z","title":"EvolveSearch: An Iterative Self-Evolving Search Agent","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-07T13:12:59.086994Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2505.22501"},"observation_digest":"sha256:bf228c893a58fec02ec6505a294bf607c14b8535c751b54d7871c680bad859ba","observation_id":"7585a35d-73f4-4f38-a71c-3e69a3f31b79","resolution":{"observed_at":"2026-08-07T13:12:59.086994Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-07T13:08:56.317112Z","title":"Webwalker: Benchmarking llms in web traversal, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.22648","last_updated":"2025-08-10T06:05:46Z","snapshot_observed_at":"2026-08-07T13:00:07.473210Z","submitted_at":"2025-05-28T17:57:07Z","title":"WebDancer: Towards Autonomous Information Seeking Agency","version":3},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-07T13:08:56.317112Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2505.22648"},"observation_digest":"sha256:119ad2a7b6cb62156cf0026e33d56e06cf606576bbf590a5d913b1ee1a17ab9d","observation_id":"9a078472-a785-4ce2-8e51-1ce3ae9a89af","resolution":{"observed_at":"2026-08-07T13:08:56.317112Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-07T11:34:47.838851Z","title":"Webwalker: Benchmarking llms in web traversal.arXiv preprint arXiv:2501.07572, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.01952","last_updated":"2025-06-02T17:59:45Z","snapshot_observed_at":"2026-08-07T11:28:26.175325Z","submitted_at":"2025-06-02T17:59:45Z","title":"WebChoreArena: Evaluating Web Browsing Agents on Realistic Tedious Web Tasks","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T11:34:47.838851Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2506.01952"},"observation_digest":"sha256:dcdfde446aad17ebeac2dab18ec91bb1bf41b22eb349871e4135aebc6dac34eb","observation_id":"e0d2fc98-8018-4903-9e98-93f790ad3b57","resolution":{"observed_at":"2026-08-07T11:34:47.838851Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":"2501.07572","doi":"10.48550/arxiv.2501.07572","metadata_source":"pith","pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jialong Wu, Baixuan Li, Runnan Fang, Wenbiao Yin, Liwen Zhang, Zhengwei Tao, Dingchu Zhang, Zekun Xi, Gang Fu, Yong Jiang, Pengjun Xie, Fei Huang, and Jingren Zhou","venue":"cs.CL","work_id":"8528e4cd-bcbc-4f57-9bd2-11ce86dd3493","year":2025},"citing_paper":{"arxiv_id":"2506.11763","last_updated":"2025-06-13T13:17:32Z","snapshot_observed_at":"2026-08-07T14:15:18.269910Z","submitted_at":"2025-06-13T13:17:32Z","title":"DeepResearch Bench: A Comprehensive Benchmark for Deep Research Agents","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-16T08:07:39.384613Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2506.11763"},"observation_digest":"sha256:04aa0b247fc03c5e4678a4b292648d29e9845eb1157154085989241a9f481638","observation_id":"51eb62ac-2c37-4665-a577-f7c589bd4f7c","resolution":{"observed_at":"2026-05-16T08:07:39.525756Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-07T00:22:54.000576Z","title":"Webwalker: Benchmarking llms in web traversal, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.14477","last_updated":"2025-06-17T12:50:35Z","snapshot_observed_at":"2026-08-07T00:15:24.463675Z","submitted_at":"2025-06-17T12:50:35Z","title":"GUI-Robust: A Comprehensive Dataset for Testing GUI Agent Robustness in Real-World Anomalies","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T00:22:54.000576Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2506.14477"},"observation_digest":"sha256:9662bcd27ef49cbbd71ae16c2345c21b0c14419df1e4b69566caa4a6f8e6d306","observation_id":"2ab43e56-f2f4-44c0-9293-8da1f2c88dc4","resolution":{"observed_at":"2026-08-07T00:22:54.000576Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-06T23:26:56.963706Z","title":"Webwalker: Benchmarking llms in web traversal","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.18096","last_updated":"2025-09-03T15:32:23Z","snapshot_observed_at":"2026-08-06T23:20:46.744749Z","submitted_at":"2025-06-22T16:52:48Z","title":"Deep Research Agents: A Systematic Examination And Roadmap","version":2},"reference_index":121,"source":"pdf_text","source_observed_at":"2026-08-06T23:26:56.963706Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2506.18096"},"observation_digest":"sha256:56b1362cc598ae3ea403617ad87ae774f3cf28d5b739b410e26cf1893b21a75c","observation_id":"86d8ba32-188d-4ba1-9ab3-ca63af754fa2","resolution":{"observed_at":"2026-08-06T23:26:56.963706Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-06T22:29:27.956083Z","title":"Webwalker: Benchmarking llms in web traversal","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.21506","last_updated":"2025-07-03T15:47:40Z","snapshot_observed_at":"2026-08-07T20:30:36.688999Z","submitted_at":"2025-06-26T17:32:50Z","title":"Mind2Web 2: Evaluating Agentic Search with Agent-as-a-Judge","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T22:29:27.956083Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2506.21506"},"observation_digest":"sha256:174847e35eea2e9af4ace33e57376aa0f932a591d82fac9126c1a6cf21e6168b","observation_id":"67ef4bc2-ea12-4f28-8d2c-495d0b49cedc","resolution":{"observed_at":"2026-08-06T22:29:27.956083Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-06T21:11:18.419497Z","title":"Webwalker: Benchmarking llms in web traver- sal,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.00841","last_updated":"2025-07-01T15:10:00Z","snapshot_observed_at":"2026-08-07T06:45:54.453824Z","submitted_at":"2025-07-01T15:10:00Z","title":"SafeMobile: Chain-level Jailbreak Detection and Automated Evaluation for Multimodal Mobile Agents","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-06T21:11:18.419497Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2507.00841"},"observation_digest":"sha256:1533c755667f5057be376b130a847323f2a573bdab592f6140b98b5762adc755","observation_id":"7e132520-222a-414d-82c6-cbf8af697334","resolution":{"observed_at":"2026-08-06T21:11:18.419497Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":"2501.07572","doi":"10.48550/arxiv.2501.07572","metadata_source":"pith","pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jialong Wu, Baixuan Li, Runnan Fang, Wenbiao Yin, Liwen Zhang, Zhengwei Tao, Dingchu Zhang, Zekun Xi, Gang Fu, Yong Jiang, Pengjun Xie, Fei Huang, and Jingren Zhou","venue":"cs.CL","work_id":"8528e4cd-bcbc-4f57-9bd2-11ce86dd3493","year":2025},"citing_paper":{"arxiv_id":"2507.21046","last_updated":"2026-01-16T20:59:08Z","snapshot_observed_at":"2026-08-01T06:32:44.461162Z","submitted_at":"2025-07-28T17:59:05Z","title":"A Survey of Self-Evolving Agents: What, When, How, and Where to Evolve on the Path to Artificial Super Intelligence","version":4},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-05-14T22:23:14.621091Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2507.21046"},"observation_digest":"sha256:6ebd66450395b70c1dc20acbfa3bea395a96842d256c8aaf78269ba63bec10fa","observation_id":"3dd944c4-a22e-4e0a-b053-96929108e8f8","resolution":{"observed_at":"2026-05-14T22:23:14.784570Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-06T10:17:13.691426Z","title":"Webwalker: Benchmarking llms in web traversal","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.00271","last_updated":"2025-09-01T02:48:41Z","snapshot_observed_at":"2026-08-07T10:28:46.635998Z","submitted_at":"2025-08-01T02:30:32Z","title":"MetaAgent: Toward Self-Evolving Agent via Tool Meta-Learning","version":2},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-06T10:17:13.691426Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2508.00271"},"observation_digest":"sha256:dfeff112a20c5ce48df9819013c9f5ccaa67a100024728d3468f8033771bbf31","observation_id":"c50c4945-572c-4664-a2a5-57796646c867","resolution":{"observed_at":"2026-08-06T10:17:13.691426Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":"2501.07572","doi":"10.48550/arxiv.2501.07572","metadata_source":"pith","pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jialong Wu, Baixuan Li, Runnan Fang, Wenbiao Yin, Liwen Zhang, Zhengwei Tao, Dingchu Zhang, Zekun Xi, Gang Fu, Yong Jiang, Pengjun Xie, Fei Huang, and Jingren Zhou","venue":"cs.CL","work_id":"8528e4cd-bcbc-4f57-9bd2-11ce86dd3493","year":2025},"citing_paper":{"arxiv_id":"2508.00414","last_updated":"2026-04-22T08:14:17Z","snapshot_observed_at":"2026-08-02T16:42:25.825514Z","submitted_at":"2025-08-01T08:11:31Z","title":"Cognitive Kernel-Pro: A Framework for Deep Research Agents and Agent Foundation Models Training","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-19T01:50:30.626249Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2508.00414"},"observation_digest":"sha256:011fb06d4edebb2573d53143f54aab2eb1fbf7f56360cb91036c6354824c4d6f","observation_id":"e7b6cfcf-5813-4a4d-95e9-51d39c9b2157","resolution":{"observed_at":"2026-05-19T01:51:57.838814Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T13:45:13.203875Z","title":"Webdancer: Towards autonomous information seeking agency, 2025a","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2509.00375","last_updated":"2025-08-30T06:02:56Z","snapshot_observed_at":"2026-08-07T12:20:14.962171Z","submitted_at":"2025-08-30T06:02:56Z","title":"Open Data Synthesis For Deep Research","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-05T13:45:13.203875Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2509.00375"},"observation_digest":"sha256:1c9ce710398549aa48b6ecce672a1e0ed56d65786133d4dc77cb8d28812c0006","observation_id":"4acf394a-80e4-4910-87b1-0967ce9e5168","resolution":{"observed_at":"2026-08-05T13:45:13.203875Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-04T19:29:00.015869Z","title":"Webwalker: Benchmarking llms in web traversal","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.09265","last_updated":"2025-09-11T08:50:01Z","snapshot_observed_at":"2026-08-07T09:00:49.598825Z","submitted_at":"2025-09-11T08:50:01Z","title":"Harnessing Uncertainty: Entropy-Modulated Policy Gradients for Long-Horizon LLM Agents","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-04T19:29:00.015869Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2509.09265"},"observation_digest":"sha256:80db941205281431f26ccef0f54e9abc57a4e1fb3f6237846c5543902e5fa276","observation_id":"9524d2bd-fded-4fd1-aa92-d028f2277706","resolution":{"observed_at":"2026-08-04T19:29:00.015869Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-04T17:57:52.512404Z","title":"Webwalker: Benchmarking llms in web traversal","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.14257","last_updated":"2026-07-23T12:32:53Z","snapshot_observed_at":"2026-08-05T11:19:48.609021Z","submitted_at":"2025-09-12T15:34:07Z","title":"Student-Centered Distillation Narrows the Agentic Gap Between Small and Large LLMs","version":3},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-04T17:57:52.512404Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2509.14257"},"observation_digest":"sha256:aaa34650a2f5a32a4405576b16a3b3af8fc4c22b6cfbf083bc0289dde80983d9","observation_id":"5dad2497-2441-41ef-86a2-ce2b2e9472e2","resolution":{"observed_at":"2026-08-04T17:57:52.512404Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":"2501.07572","doi":"10.48550/arxiv.2501.07572","metadata_source":"pith","pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jialong Wu, Baixuan Li, Runnan Fang, Wenbiao Yin, Liwen Zhang, Zhengwei Tao, Dingchu Zhang, Zekun Xi, Gang Fu, Yong Jiang, Pengjun Xie, Fei Huang, and Jingren Zhou","venue":"cs.CL","work_id":"8528e4cd-bcbc-4f57-9bd2-11ce86dd3493","year":2025},"citing_paper":{"arxiv_id":"2511.11793","last_updated":"2026-04-21T16:02:11Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-14T18:52:07Z","title":"MiroThinker: Pushing the Performance Boundaries of Open-Source Research Agents via Model, Context, and Interactive Scaling","version":3},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-17T21:44:18.744201Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2511.11793"},"observation_digest":"sha256:584bea30e300b9463c1cb473a859552cbcc2c945d132b0b53e0cc7d3b62ff3f4","observation_id":"e31a7e2e-05c8-4923-842f-0e992185d5c1","resolution":{"observed_at":"2026-05-17T21:45:17.854810Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":"2501.07572","doi":"10.48550/arxiv.2501.07572","metadata_source":"pith","pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jialong Wu, Baixuan Li, Runnan Fang, Wenbiao Yin, Liwen Zhang, Zhengwei Tao, Dingchu Zhang, Zekun Xi, Gang Fu, Yong Jiang, Pengjun Xie, Fei Huang, and Jingren Zhou","venue":"cs.CL","work_id":"8528e4cd-bcbc-4f57-9bd2-11ce86dd3493","year":2025},"citing_paper":{"arxiv_id":"2601.06943","last_updated":"2026-05-18T06:58:21Z","snapshot_observed_at":"2026-08-02T18:32:02.969277Z","submitted_at":"2026-01-11T15:07:37Z","title":"Watching, Reasoning, and Searching: A Video Deep Research Benchmark on Open Web for Agentic Video Reasoning","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-21T16:35:24.557809Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2601.06943"},"observation_digest":"sha256:0982416911b1f45422dca64722b314a5dae4ca5e7835382d2b053ad616b9b885","observation_id":"09e60d6c-21b9-4aa7-8585-732c7c5d4411","resolution":{"observed_at":"2026-05-21T16:40:22.779383Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":"2501.07572","doi":"10.48550/arxiv.2501.07572","metadata_source":"pith","pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jialong Wu, Baixuan Li, Runnan Fang, Wenbiao Yin, Liwen Zhang, Zhengwei Tao, Dingchu Zhang, Zekun Xi, Gang Fu, Yong Jiang, Pengjun Xie, Fei Huang, and Jingren Zhou","venue":"cs.CL","work_id":"8528e4cd-bcbc-4f57-9bd2-11ce86dd3493","year":2025},"citing_paper":{"arxiv_id":"2603.04751","last_updated":"2026-04-27T08:53:25Z","snapshot_observed_at":"2026-07-31T05:56:51.116141Z","submitted_at":"2026-03-05T02:56:42Z","title":"Evaluating the Search Agent in a Parallel World","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-15T17:02:10.963839Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2603.04751"},"observation_digest":"sha256:5d8e329b4643789d6f2a1c11631e939cc2ae4003265a5cca950535d70e584dff","observation_id":"9f3870a8-9548-4161-b2ca-782b39f61d76","resolution":{"observed_at":"2026-05-15T17:06:19.348499Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":"2501.07572","doi":"10.48550/arxiv.2501.07572","metadata_source":"pith","pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jialong Wu, Baixuan Li, Runnan Fang, Wenbiao Yin, Liwen Zhang, Zhengwei Tao, Dingchu Zhang, Zekun Xi, Gang Fu, Yong Jiang, Pengjun Xie, Fei Huang, and Jingren Zhou","venue":"cs.CL","work_id":"8528e4cd-bcbc-4f57-9bd2-11ce86dd3493","year":2025},"citing_paper":{"arxiv_id":"2603.25342","last_updated":"2026-04-29T12:00:40Z","snapshot_observed_at":"2026-07-06T22:50:37.468743Z","submitted_at":"2026-03-26T11:37:26Z","title":"From Intent to Evidence: A Categorical Approach for Structural Evaluation of Deep Research Agents","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-15T00:25:17.671344Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2603.25342"},"observation_digest":"sha256:18442d184107fc61fbf6efb170bed79b72a0671b13a00f600bc13dca32a0437c","observation_id":"eee2d20b-f330-4a19-82e5-16d691ec033f","resolution":{"observed_at":"2026-05-15T00:28:23.594867Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":"2501.07572","doi":"10.48550/arxiv.2501.07572","metadata_source":"pith","pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jialong Wu, Baixuan Li, Runnan Fang, Wenbiao Yin, Liwen Zhang, Zhengwei Tao, Dingchu Zhang, Zekun Xi, Gang Fu, Yong Jiang, Pengjun Xie, Fei Huang, and Jingren Zhou","venue":"cs.CL","work_id":"8528e4cd-bcbc-4f57-9bd2-11ce86dd3493","year":2025},"citing_paper":{"arxiv_id":"2604.03679","last_updated":"2026-04-04T10:46:09Z","snapshot_observed_at":"2026-07-06T22:52:48.277339Z","submitted_at":"2026-04-04T10:46:09Z","title":"LightThinker++: From Reasoning Compression to Memory Management","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-05-13T17:25:28.432170Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2604.03679"},"observation_digest":"sha256:e2d602b7b7b0fe7b727d3f7ebf03fb5cdacaaf82f4ba42151a1bae21340b147b","observation_id":"51e16ac6-ecbc-411c-ac6b-8480a01d5ce4","resolution":{"observed_at":"2026-05-13T17:28:02.490629Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":"2501.07572","doi":"10.48550/arxiv.2501.07572","metadata_source":"pith","pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jialong Wu, Baixuan Li, Runnan Fang, Wenbiao Yin, Liwen Zhang, Zhengwei Tao, Dingchu Zhang, Zekun Xi, Gang Fu, Yong Jiang, Pengjun Xie, Fei Huang, and Jingren Zhou","venue":"cs.CL","work_id":"8528e4cd-bcbc-4f57-9bd2-11ce86dd3493","year":2025},"citing_paper":{"arxiv_id":"2604.07927","last_updated":"2026-04-28T06:09:15Z","snapshot_observed_at":"2026-07-06T22:57:09.435331Z","submitted_at":"2026-04-09T07:47:31Z","title":"EigentSearch-Q+: Enhancing Deep Research Agents with Structured Reasoning Tools","version":3},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-10T18:36:45.297356Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2604.07927"},"observation_digest":"sha256:ce83f2b6d1778e5ac132029731cbb199192fd2c0ae5ecfeb6ba412c0cddb0e64","observation_id":"52fbe3a7-a5e8-4bc4-a7ac-c2ed0bd5b645","resolution":{"observed_at":"2026-05-11T00:15:55.614890Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":"2501.07572","doi":"10.48550/arxiv.2501.07572","metadata_source":"pith","pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jialong Wu, Baixuan Li, Runnan Fang, Wenbiao Yin, Liwen Zhang, Zhengwei Tao, Dingchu Zhang, Zekun Xi, Gang Fu, Yong Jiang, Pengjun Xie, Fei Huang, and Jingren Zhou","venue":"cs.CL","work_id":"8528e4cd-bcbc-4f57-9bd2-11ce86dd3493","year":2025},"citing_paper":{"arxiv_id":"2604.17931","last_updated":"2026-07-01T16:44:57Z","snapshot_observed_at":"2026-07-06T23:04:55.190916Z","submitted_at":"2026-04-20T08:11:09Z","title":"LiteResearcher: A Scalable Agentic RL Training Framework for Deep Research Agent","version":4},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-07-05T15:03:50.420072Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2604.17931"},"observation_digest":"sha256:3d095142ab7ec5536b187cb049b97eee2a21233b12db3b7b97611fe5072d7bb3","observation_id":"83f43b55-3511-4922-b5c3-4d183c76bdd6","resolution":{"observed_at":"2026-07-05T15:11:10.913704Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":"2501.07572","doi":"10.48550/arxiv.2501.07572","metadata_source":"pith","pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jialong Wu, Baixuan Li, Runnan Fang, Wenbiao Yin, Liwen Zhang, Zhengwei Tao, Dingchu Zhang, Zekun Xi, Gang Fu, Yong Jiang, Pengjun Xie, Fei Huang, and Jingren Zhou","venue":"cs.CL","work_id":"8528e4cd-bcbc-4f57-9bd2-11ce86dd3493","year":2025},"citing_paper":{"arxiv_id":"2604.18131","last_updated":"2026-04-20T11:54:20Z","snapshot_observed_at":"2026-07-06T23:05:04.279268Z","submitted_at":"2026-04-20T11:54:20Z","title":"Training LLM Agents for Spontaneous, Reward-Free Self-Evolution via World Knowledge Exploration","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-10T04:36:27.381942Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2604.18131"},"observation_digest":"sha256:4d16843e65b7665cfc6cdfb2550ba10fed26cef6094803fa4549248046734781","observation_id":"c14ae87d-840d-4b97-af62-f824bc72ab98","resolution":{"observed_at":"2026-05-10T12:10:23.405556Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":"2501.07572","doi":"10.48550/arxiv.2501.07572","metadata_source":"pith","pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jialong Wu, Baixuan Li, Runnan Fang, Wenbiao Yin, Liwen Zhang, Zhengwei Tao, Dingchu Zhang, Zekun Xi, Gang Fu, Yong Jiang, Pengjun Xie, Fei Huang, and Jingren Zhou","venue":"cs.CL","work_id":"8528e4cd-bcbc-4f57-9bd2-11ce86dd3493","year":2025},"citing_paper":{"arxiv_id":"2604.18292","last_updated":"2026-04-20T14:01:10Z","snapshot_observed_at":"2026-07-06T23:05:13.178333Z","submitted_at":"2026-04-20T14:01:10Z","title":"Agent-World: Scaling Real-World Environment Synthesis for Evolving General Agent Intelligence","version":1},"reference_index":106,"source":"pdf_text","source_observed_at":"2026-05-10T05:24:00.503836Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2604.18292"},"observation_digest":"sha256:56497f21bbfe26c944416177f56284035dc60ab631cc54ece462b275c8a4203a","observation_id":"382bab25-5023-4df6-8feb-db1eca045751","resolution":{"observed_at":"2026-05-10T05:25:54.613490Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":"2501.07572","doi":"10.48550/arxiv.2501.07572","metadata_source":"pith","pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jialong Wu, Baixuan Li, Runnan Fang, Wenbiao Yin, Liwen Zhang, Zhengwei Tao, Dingchu Zhang, Zekun Xi, Gang Fu, Yong Jiang, Pengjun Xie, Fei Huang, and Jingren Zhou","venue":"cs.CL","work_id":"8528e4cd-bcbc-4f57-9bd2-11ce86dd3493","year":2025},"citing_paper":{"arxiv_id":"2604.18292","last_updated":"2026-04-20T14:01:10Z","snapshot_observed_at":"2026-07-06T23:05:13.178333Z","submitted_at":"2026-04-20T14:01:10Z","title":"Agent-World: Scaling Real-World Environment Synthesis for Evolving General Agent Intelligence","version":1},"reference_index":107,"source":"pdf_text","source_observed_at":"2026-05-10T05:24:00.503836Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2604.18292"},"observation_digest":"sha256:cd1106a0019b885364a03a5147f29e7ba545f29b1a7f9cda5aa58005e0ae939b","observation_id":"4283de78-81fe-47ea-9ba2-3592a8bc7a96","resolution":{"observed_at":"2026-05-10T05:25:54.251167Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":"2501.07572","doi":"10.48550/arxiv.2501.07572","metadata_source":"pith","pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jialong Wu, Baixuan Li, Runnan Fang, Wenbiao Yin, Liwen Zhang, Zhengwei Tao, Dingchu Zhang, Zekun Xi, Gang Fu, Yong Jiang, Pengjun Xie, Fei Huang, and Jingren Zhou","venue":"cs.CL","work_id":"8528e4cd-bcbc-4f57-9bd2-11ce86dd3493","year":2025},"citing_paper":{"arxiv_id":"2604.20486","last_updated":"2026-04-22T12:20:46Z","snapshot_observed_at":"2026-08-06T20:38:43.258492Z","submitted_at":"2026-04-22T12:20:46Z","title":"ProMMSearchAgent: A Generalizable Multimodal Search Agent Trained with Process-Oriented Rewards","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-10T01:12:17.469552Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2604.20486"},"observation_digest":"sha256:d4204454caf1a23bdcd15dde11978ef5fb8b704ff6a831c0a1759611de8abf0f","observation_id":"1af9ed58-292a-4b10-a8a6-14e1ef88c4b5","resolution":{"observed_at":"2026-05-11T13:41:09.901542Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":"2501.07572","doi":"10.48550/arxiv.2501.07572","metadata_source":"pith","pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jialong Wu, Baixuan Li, Runnan Fang, Wenbiao Yin, Liwen Zhang, Zhengwei Tao, Dingchu Zhang, Zekun Xi, Gang Fu, Yong Jiang, Pengjun Xie, Fei Huang, and Jingren Zhou","venue":"cs.CL","work_id":"8528e4cd-bcbc-4f57-9bd2-11ce86dd3493","year":2025},"citing_paper":{"arxiv_id":"2605.01489","last_updated":"2026-05-26T13:12:52Z","snapshot_observed_at":"2026-07-06T23:14:43.034336Z","submitted_at":"2026-05-02T15:26:45Z","title":"SciResearcher: Scaling Deep Research Agents for Frontier Scientific Reasoning","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-05-09T14:18:14.048230Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2605.01489"},"observation_digest":"sha256:5e983761ff8c833689854fe464bad5139b5bb77b978ef3996f7cc13affce3da4","observation_id":"84dc6e30-60c9-4849-beab-4785b157c31b","resolution":{"observed_at":"2026-05-11T17:01:08.439416Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":"2501.07572","doi":"10.48550/arxiv.2501.07572","metadata_source":"pith","pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jialong Wu, Baixuan Li, Runnan Fang, Wenbiao Yin, Liwen Zhang, Zhengwei Tao, Dingchu Zhang, Zekun Xi, Gang Fu, Yong Jiang, Pengjun Xie, Fei Huang, and Jingren Zhou","venue":"cs.CL","work_id":"8528e4cd-bcbc-4f57-9bd2-11ce86dd3493","year":2025},"citing_paper":{"arxiv_id":"2605.08762","last_updated":"2026-05-09T07:47:42Z","snapshot_observed_at":"2026-07-06T23:20:57.084438Z","submitted_at":"2026-05-09T07:47:42Z","title":"Omni-DeepSearch: A Benchmark for Audio-Driven Omni-Modal Deep Search","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-12T03:21:52.972107Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2605.08762"},"observation_digest":"sha256:6debcd101399764bf3d924362c45af4615cb86b53318c67876b79f250e2b08ae","observation_id":"6f510158-9be6-4e56-adf8-c2853ac9fbaf","resolution":{"observed_at":"2026-05-12T07:26:25.909056Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":"2501.07572","doi":"10.48550/arxiv.2501.07572","metadata_source":"pith","pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jialong Wu, Baixuan Li, Runnan Fang, Wenbiao Yin, Liwen Zhang, Zhengwei Tao, Dingchu Zhang, Zekun Xi, Gang Fu, Yong Jiang, Pengjun Xie, Fei Huang, and Jingren Zhou","venue":"cs.CL","work_id":"8528e4cd-bcbc-4f57-9bd2-11ce86dd3493","year":2025},"citing_paper":{"arxiv_id":"2605.12004","last_updated":"2026-05-12T11:54:23Z","snapshot_observed_at":"2026-08-07T08:55:02.984200Z","submitted_at":"2026-05-12T11:54:23Z","title":"Learning Agentic Policy from Action Guidance","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-05-13T05:02:49.206053Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2605.12004"},"observation_digest":"sha256:c509a9e1a8db72d05e009b156029936cdf6f0f5205d40b7379bf712aeaad4491","observation_id":"f4c1fd2f-5f2e-44e6-af5b-fb0dce07b088","resolution":{"observed_at":"2026-05-13T05:07:17.791143Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":"2501.07572","doi":"10.48550/arxiv.2501.07572","metadata_source":"pith","pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jialong Wu, Baixuan Li, Runnan Fang, Wenbiao Yin, Liwen Zhang, Zhengwei Tao, Dingchu Zhang, Zekun Xi, Gang Fu, Yong Jiang, Pengjun Xie, Fei Huang, and Jingren Zhou","venue":"cs.CL","work_id":"8528e4cd-bcbc-4f57-9bd2-11ce86dd3493","year":2025},"citing_paper":{"arxiv_id":"2606.07689","last_updated":"2026-06-05T06:25:11Z","snapshot_observed_at":"2026-07-06T23:47:19.773490Z","submitted_at":"2026-06-05T06:25:11Z","title":"Struct-Searcher: Agentic Structural Thinking Advances Multimodal Deep Information Seeking","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-06-27T22:20:53.187362Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2606.07689"},"observation_digest":"sha256:7ab22f956cff04d7b3b4f3324127e015abd4e9ec2ba9b8adf35d988361a60b13","observation_id":"dbe14a35-2c87-46be-b4f1-7364b791aef6","resolution":{"observed_at":"2026-07-02T16:47:10.203346Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":"2501.07572","doi":"10.48550/arxiv.2501.07572","metadata_source":"pith","pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Jialong Wu, Baixuan Li, Runnan Fang, Wenbiao Yin, Liwen Zhang, Zhengwei Tao, Dingchu Zhang, Zekun Xi, Gang Fu, Yong Jiang, Pengjun Xie, Fei Huang, and Jingren Zhou","venue":"cs.CL","work_id":"8528e4cd-bcbc-4f57-9bd2-11ce86dd3493","year":2025},"citing_paper":{"arxiv_id":"2606.31504","last_updated":"2026-06-30T11:22:24Z","snapshot_observed_at":"2026-08-02T16:40:34.907828Z","submitted_at":"2026-06-30T11:22:24Z","title":"SimpleSearch-VL: A Simple Recipe for Multimodal Agentic Deep Search","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-07-01T06:02:48.532478Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2606.31504"},"observation_digest":"sha256:44ef4760a848399d83aaaa8a5b8038a64f5fd5b34f136b3fcf51fdff58f13c62","observation_id":"10e6e54c-0a4f-4786-994d-c0078531566f","resolution":{"observed_at":"2026-07-01T09:55:41.075928Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-02T04:50:31.648085Z","title":"WebWalker: Benchmarking llms in web traversal.arXiv preprint arXiv:2501.07572, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.13591","last_updated":"2026-07-15T08:32:40Z","snapshot_observed_at":"2026-08-06T21:05:54.197003Z","submitted_at":"2026-07-15T08:32:40Z","title":"Memory as a Controlled Process: Learned Adaptive Memory Management for LLM Agents","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-02T04:50:31.648085Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2607.13591"},"observation_digest":"sha256:3d2ce62771cdd274072605af988d41022838513b608c5d7b58ba710f21664907","observation_id":"4c94a364-327a-4d94-a830-b187464b95eb","resolution":{"observed_at":"2026-08-02T04:50:31.648085Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-07-31T20:06:14.128623Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28026","last_updated":"2026-07-30T11:14:11Z","snapshot_observed_at":"2026-08-03T00:02:04.534763Z","submitted_at":"2026-07-30T11:14:11Z","title":"Contrastive Reinforced Policy Optimization via Privileged Self-Distillation","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-07-31T20:06:14.128623Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2607.28026"},"observation_digest":"sha256:70527deb7ad3cb434f6d46d044360dd11008ff3d8d3f77b3e233d6bc699d56ef","observation_id":"5c68a7ed-1f71-49fa-b8e9-d74a21fbd5aa","resolution":{"observed_at":"2026-07-31T20:06:14.128623Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-07-31T20:06:17.267411Z","title":"CoRR , volume =","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28026","last_updated":"2026-07-30T11:14:11Z","snapshot_observed_at":"2026-08-03T00:02:04.534763Z","submitted_at":"2026-07-30T11:14:11Z","title":"Contrastive Reinforced Policy Optimization via Privileged Self-Distillation","version":1},"reference_index":86,"source":"arxiv_source","source_observed_at":"2026-07-31T20:06:17.267411Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2607.28026"},"observation_digest":"sha256:c800652f5a3c9665520d98ec1f067bbc68595cd40b7cacbf21119a1373208737","observation_id":"263eb880-823c-416d-8b72-850373b39ef7","resolution":{"observed_at":"2026-07-31T20:06:17.267411Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-04T15:12:55.811934Z","title":"2025 , archivePrefix =","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.02097","last_updated":"2026-08-03T11:58:21Z","snapshot_observed_at":"2026-08-07T22:15:33.180420Z","submitted_at":"2026-08-03T11:58:21Z","title":"Fetch-then-Explore: Decoupling Selection from Extraction over a Persistent Workspace for Search Agents","version":1},"reference_index":117,"source":"arxiv_source","source_observed_at":"2026-08-04T15:12:55.811934Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2608.02097"},"observation_digest":"sha256:c9ced83fc2456c1df24c27d99bcfa43a4a26dd72d07b4dbc252c5bb5f327b1af","observation_id":"de9f7ef5-cd71-4eab-bd58-170c6954da30","resolution":{"observed_at":"2026-08-04T15:12:55.811934Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.07572","snapshot_observed_at":"2026-08-04T13:31:24.190907Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.02163","last_updated":"2026-08-03T12:43:00Z","snapshot_observed_at":"2026-08-07T17:38:02.246950Z","submitted_at":"2026-08-03T12:43:00Z","title":"From Simple QA to Deep Research: A Verifiable Benchmark Constructed through Iterative Task Evolution","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-04T13:31:24.190907Z"},"links":{"cited_paper":"/paper/2501.07572","citing_paper":"/paper/2608.02163"},"observation_digest":"sha256:c45c5d4e73ba03289d83b791b8092ec217bc9ea82515b0eabae661ae12a6af5b","observation_id":"7dc2973b-826c-4975-8efb-e55a894676be","resolution":{"observed_at":"2026-08-04T13:31:24.190907Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2501.07572/citation-record","integrity":"/paper/2501.07572/integrity","json":"/paper/2501.07572/citation-record.json","paper":"/paper/2501.07572"},"outbound":[],"paper":{"arxiv_id":"2501.07572","last_updated":"2025-08-10T05:59:20Z","latest_version":3,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-04T14:46:05.418102Z","submitted_at":"2025-01-13T18:58:07Z","title":"WebWalker: Benchmarking LLMs in Web Traversal"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 39 inbound Pith citation observations for arXiv:2501.07572."}