{"as_of":"2026-08-10T18:44:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:f857698c631fa5d9b65c7d963321f970686c2be0be6e057ee66d2d0e047321d2","coverage":[{"denominator":2,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-12T23:36:55.286437Z","state":"measured"},{"denominator":5,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":5,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":3,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":3,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T13:09:06.712982Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-07T23:54:18.647304Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2604.08988","last_updated":"2026-05-24T11:57:07Z","snapshot_observed_at":"2026-08-07T20:51:33.776054Z","submitted_at":"2026-04-10T05:49:50Z","title":"SEA-Eval: A Benchmark for Evaluating Self-Evolving Agents Beyond Episodic Assessment","version":3},"cited_work":{"arxiv_id":"2604.08988","doi":null,"metadata_source":"pith","pith_arxiv_id":"2604.08988","snapshot_observed_at":"2026-07-07T23:54:18.647304Z","title":"SEA-Eval: A Benchmark for Evaluating Self-Evolving Agents Beyond Episodic Assessment","venue":"cs.AI","work_id":"2989332a-9af5-4c60-a66d-c9fa9247df41","year":2026},"citing_paper":{"arxiv_id":"2607.05202","last_updated":"2026-07-06T15:17:09Z","snapshot_observed_at":"2026-08-01T16:59:54.701989Z","submitted_at":"2026-07-06T15:17:09Z","title":"EvoAgentBench: Benchmarking Agent Self-Evolution via Ability Transfer","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-07-07T23:48:16.309103Z"},"links":{"cited_paper":"/paper/2604.08988","citing_paper":"/paper/2607.05202"},"observation_digest":"sha256:7be7787fc01a7009310ff6e09459a69079e1c4932a68f7a3c4577d7e5c0a1c75","observation_id":"e621c63e-2a99-46be-9728-1efdb61d4f12","resolution":{"observed_at":"2026-07-07T23:54:18.649142Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2604.08988","last_updated":"2026-05-24T11:57:07Z","snapshot_observed_at":"2026-08-07T20:51:33.776054Z","submitted_at":"2026-04-10T05:49:50Z","title":"SEA-Eval: A Benchmark for Evaluating Self-Evolving Agents Beyond Episodic Assessment","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2604.08988","snapshot_observed_at":"2026-08-05T04:26:29.286560Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.02636","last_updated":"2026-07-31T04:02:51Z","snapshot_observed_at":"2026-08-09T14:34:24.393371Z","submitted_at":"2026-07-31T04:02:51Z","title":"Rethinking Self-Evolving Agent Skills: Feedback Dynamics over Multiple Rounds","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-05T04:26:29.286560Z"},"links":{"cited_paper":"/paper/2604.08988","citing_paper":"/paper/2608.02636"},"observation_digest":"sha256:5c538e25361b92125033a1347670c0bf85645a7adddcb214f65c9d62bced9ce1","observation_id":"0fed06f9-0529-4d2f-9cd8-3020561fa34e","resolution":{"observed_at":"2026-08-05T04:26:29.286560Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2604.08988","last_updated":"2026-05-24T11:57:07Z","snapshot_observed_at":"2026-08-07T20:51:33.776054Z","submitted_at":"2026-04-10T05:49:50Z","title":"SEA-Eval: A Benchmark for Evaluating Self-Evolving Agents Beyond Episodic Assessment","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2604.08988","snapshot_observed_at":"2026-08-05T13:09:06.712982Z","title":"SEA-Eval: A benchmark for evaluating self-evolving agents beyond episodic assessment.CoRR, abs/2604.08988, 2026","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.03764","last_updated":"2026-08-04T14:51:56Z","snapshot_observed_at":"2026-08-10T16:07:02.243649Z","submitted_at":"2026-08-04T14:51:56Z","title":"GDPevo: Evaluating Agent Self-Evolution on Real Business Tasks","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-05T13:09:06.712982Z"},"links":{"cited_paper":"/paper/2604.08988","citing_paper":"/paper/2608.03764"},"observation_digest":"sha256:3b8cfaafd8f40f7ca3e2e49daac5cb870185fbbcf14efc1df53e01f0a6de25ca","observation_id":"0f7172b8-3e07-4067-818d-4c13b2c0a6df","resolution":{"observed_at":"2026-08-05T13:09:06.712982Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2604.08988/citation-record","integrity":"/paper/2604.08988/integrity","json":"/paper/2604.08988/citation-record.json","paper":"/paper/2604.08988"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T23:36:55.286437Z","title":"Yuhong Cao, Jeric Lew, Jingsong Liang, Jin Cheng, and Guillaume Sartoretti","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2604.08988","last_updated":"2026-05-24T11:57:07Z","snapshot_observed_at":"2026-08-07T20:51:33.776054Z","submitted_at":"2026-04-10T05:49:50Z","title":"SEA-Eval: A Benchmark for Evaluating Self-Evolving Agents Beyond Episodic Assessment","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-07-12T23:36:55.286437Z"},"links":{"citing_paper":"/paper/2604.08988"},"observation_digest":"sha256:86ac84809e891278befda35459b5f5f4b8ce3af5a1bde9d46dd663f65745963b","observation_id":"d45ae89d-29f8-4bb4-9f28-63673416edb9","resolution":{"observed_at":"2026-07-12T23:36:55.286437Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07972","last_updated":"2024-05-30T08:55:12Z","snapshot_observed_at":"2026-08-09T01:39:11.956106Z","submitted_at":"2024-04-11T17:56:05Z","title":"OSWorld: Benchmarking Multimodal Agents for Open-Ended Tasks in Real Computer Environments","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07972","snapshot_observed_at":"2026-07-12T23:36:55.286437Z","title":"InFind- ings of the Association for Computational Linguistics: ACL 2025, pages 3350–3376","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2604.08988","last_updated":"2026-05-24T11:57:07Z","snapshot_observed_at":"2026-08-07T20:51:33.776054Z","submitted_at":"2026-04-10T05:49:50Z","title":"SEA-Eval: A Benchmark for Evaluating Self-Evolving Agents Beyond Episodic Assessment","version":3},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-12T23:36:55.286437Z"},"links":{"cited_paper":"/paper/2404.07972","citing_paper":"/paper/2604.08988"},"observation_digest":"sha256:266c81339c0376ae7a3434006fd6ff078f8bc8c75f11ebe57f11fcab623628a9","observation_id":"ecc082ca-81bb-41b2-aee2-eddc131da08c","resolution":{"observed_at":"2026-07-12T23:36:55.286437Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2604.08988","last_updated":"2026-05-24T11:57:07Z","latest_version":3,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-07T20:51:33.776054Z","submitted_at":"2026-04-10T05:49:50Z","title":"SEA-Eval: A Benchmark for Evaluating Self-Evolving Agents Beyond Episodic Assessment"},"reference_resolution":{"displayed":2,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":2,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":2},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 2 of 2 outbound references and 3 inbound Pith citation observations for arXiv:2604.08988."}