{"as_of":"2026-08-12T03:06:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:6d1a3e43d80a4e6d2867afe1e7373d10f8df2a267daa11dbf1363bc307350771","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":10,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":10,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-11T06:34:44.6726+00:00","state":"measured"},{"denominator":10,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":10,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-10T23:24:14.619683Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-10T06:15:00.866473Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"1709.04326","last_updated":"2018-09-19T19:22:48Z","snapshot_observed_at":"2026-07-06T05:59:28.465615Z","submitted_at":"2017-09-13T13:42:15Z","title":"Learning with Opponent-Learning Awareness","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1709.04326","snapshot_observed_at":"2026-08-10T23:24:14.619683Z","title":"Udaya Ghai, Udari Madhushani, Naomi Leonard, and Elad Hazan","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.20523","last_updated":"2024-12-29T17:15:40Z","snapshot_observed_at":"2026-08-10T23:16:43.216129Z","submitted_at":"2024-12-29T17:15:40Z","title":"Game Theory and Multi-Agent Reinforcement Learning : From Nash Equilibria to Evolutionary Dynamics","version":1},"reference_index":2017,"source":"pdf_text","source_observed_at":"2026-08-10T23:24:14.619683Z"},"links":{"cited_paper":"/paper/1709.04326","citing_paper":"/paper/2412.20523"},"observation_digest":"sha256:6a86989f0d61981a9e676f62bae33813ff879e937f7f4652d2d41216bb7be88a","observation_id":"7eb79500-a344-4e33-b335-a7b6fbca3df5","resolution":{"observed_at":"2026-08-10T23:24:14.619683Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1709.04326","last_updated":"2018-09-19T19:22:48Z","snapshot_observed_at":"2026-07-06T05:59:28.465615Z","submitted_at":"2017-09-13T13:42:15Z","title":"Learning with Opponent-Learning Awareness","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1709.04326","snapshot_observed_at":"2026-08-07T13:57:39.358832Z","title":"Learning with opponent-learning awareness","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2505.20579","last_updated":"2026-07-18T20:07:20Z","snapshot_observed_at":"2026-08-07T13:49:07.327912Z","submitted_at":"2025-05-26T23:28:52Z","title":"The challenge of hidden gifts in multi-agent reinforcement learning","version":7},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-07T13:57:39.358832Z"},"links":{"cited_paper":"/paper/1709.04326","citing_paper":"/paper/2505.20579"},"observation_digest":"sha256:fca71e6db909024a2ce79c6ade92fd59350a475b05797fe01671b70186715caf","observation_id":"1d218df6-f17b-4151-b312-7a51599cfd3e","resolution":{"observed_at":"2026-08-07T13:57:39.358832Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1709.04326","last_updated":"2018-09-19T19:22:48Z","snapshot_observed_at":"2026-07-06T05:59:28.465615Z","submitted_at":"2017-09-13T13:42:15Z","title":"Learning with Opponent-Learning Awareness","version":4},"cited_work":{"arxiv_id":"1709.04326","doi":"10.48550/arxiv.1709.04326","metadata_source":"pith","pith_arxiv_id":"1709.04326","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Learning with Opponent-Learning Awareness","venue":"cs.AI","work_id":"61ef30eb-2271-4eee-9b6a-e01c6a4bcad7","year":2017},"citing_paper":{"arxiv_id":"2506.02387","last_updated":"2026-04-13T08:26:12Z","snapshot_observed_at":"2026-08-11T01:40:24.202082Z","submitted_at":"2025-06-03T02:57:38Z","title":"VS-Bench: Evaluating VLMs for Strategic Abilities in Multi-Agent Environments","version":3},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-19T11:57:08.314088Z"},"links":{"cited_paper":"/paper/1709.04326","citing_paper":"/paper/2506.02387"},"observation_digest":"sha256:0451430162bf70e9da10f3fa252050c488838cd270a9e0a6420f33495dae6d4a","observation_id":"1f983252-eb2c-41f9-aa30-8982c065f9ec","resolution":{"observed_at":"2026-05-19T11:57:16.263774Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1709.04326","last_updated":"2018-09-19T19:22:48Z","snapshot_observed_at":"2026-07-06T05:59:28.465615Z","submitted_at":"2017-09-13T13:42:15Z","title":"Learning with Opponent-Learning Awareness","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1709.04326","snapshot_observed_at":"2026-08-07T00:16:42.264738Z","title":"Learning with opponent-learning awareness.arXiv preprint arXiv:1709.04326, 2017","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.14990","last_updated":"2026-06-18T15:11:38Z","snapshot_observed_at":"2026-08-11T13:04:12.483486Z","submitted_at":"2025-06-17T21:50:04Z","title":"MEAL: A Benchmark for Continual Multi-Agent Reinforcement Learning","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T00:16:42.264738Z"},"links":{"cited_paper":"/paper/1709.04326","citing_paper":"/paper/2506.14990"},"observation_digest":"sha256:2cbb7f46978e020f0ec975d0208b8f9f91c53af0cebd25c2e9a6bb38b0a08d20","observation_id":"a0296929-eb95-4f4a-a644-f1c206dab12d","resolution":{"observed_at":"2026-08-07T00:16:42.264738Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1709.04326","last_updated":"2018-09-19T19:22:48Z","snapshot_observed_at":"2026-07-06T05:59:28.465615Z","submitted_at":"2017-09-13T13:42:15Z","title":"Learning with Opponent-Learning Awareness","version":4},"cited_work":{"arxiv_id":"1709.04326","doi":"10.48550/arxiv.1709.04326","metadata_source":"pith","pith_arxiv_id":"1709.04326","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Learning with Opponent-Learning Awareness","venue":"cs.AI","work_id":"61ef30eb-2271-4eee-9b6a-e01c6a4bcad7","year":2017},"citing_paper":{"arxiv_id":"2605.12655","last_updated":"2026-06-10T15:03:44Z","snapshot_observed_at":"2026-08-08T21:43:49.970250Z","submitted_at":"2026-05-12T19:01:16Z","title":"Robust Instruction Compliance in Cooperative Multi-Agent Reinforcement Learning","version":3},"reference_index":125,"source":"arxiv_source","source_observed_at":"2026-06-30T22:11:35.277901Z"},"links":{"cited_paper":"/paper/1709.04326","citing_paper":"/paper/2605.12655"},"observation_digest":"sha256:18c70aea98bfd8c6aa7275af179d6c20c1b09cd3027af8e029f5c9160f7e844d","observation_id":"7a123442-5a63-44ac-9744-68331814acfd","resolution":{"observed_at":"2026-06-30T22:15:05.882560Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1709.04326","last_updated":"2018-09-19T19:22:48Z","snapshot_observed_at":"2026-07-06T05:59:28.465615Z","submitted_at":"2017-09-13T13:42:15Z","title":"Learning with Opponent-Learning Awareness","version":4},"cited_work":{"arxiv_id":"1709.04326","doi":"10.48550/arxiv.1709.04326","metadata_source":"pith","pith_arxiv_id":"1709.04326","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Learning with Opponent-Learning Awareness","venue":"cs.AI","work_id":"61ef30eb-2271-4eee-9b6a-e01c6a4bcad7","year":2017},"citing_paper":{"arxiv_id":"2605.20348","last_updated":"2026-05-19T18:03:48Z","snapshot_observed_at":"2026-08-02T15:33:17.310428Z","submitted_at":"2026-05-19T18:03:48Z","title":"Memory-Induced Supra-Competitive Outcomes Between Deep Reinforcement Learning Agents in Optimal Trade Execution","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-05-21T07:05:07.816803Z"},"links":{"cited_paper":"/paper/1709.04326","citing_paper":"/paper/2605.20348"},"observation_digest":"sha256:1289d3157156f2130a1a19d68b25e8c0d3057155f283b3488fbc60d883b323f7","observation_id":"3393aa44-4f0f-415a-96b0-8c3068bd786c","resolution":{"observed_at":"2026-05-21T07:09:46.474984Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1709.04326","last_updated":"2018-09-19T19:22:48Z","snapshot_observed_at":"2026-07-06T05:59:28.465615Z","submitted_at":"2017-09-13T13:42:15Z","title":"Learning with Opponent-Learning Awareness","version":4},"cited_work":{"arxiv_id":"1709.04326","doi":"10.48550/arxiv.1709.04326","metadata_source":"pith","pith_arxiv_id":"1709.04326","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Learning with Opponent-Learning Awareness","venue":"cs.AI","work_id":"61ef30eb-2271-4eee-9b6a-e01c6a4bcad7","year":2017},"citing_paper":{"arxiv_id":"2606.04359","last_updated":"2026-06-03T02:20:02Z","snapshot_observed_at":"2026-07-06T23:44:28.265478Z","submitted_at":"2026-06-03T02:20:02Z","title":"Learning to cooperate with emergent reputation via multi-agent reinforcement learning","version":1},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-06-28T04:23:39.844854Z"},"links":{"cited_paper":"/paper/1709.04326","citing_paper":"/paper/2606.04359"},"observation_digest":"sha256:f1c4eb8f6f42f61d8e7df0f531ff2fe51954757c2ff4cac2149ea01676f00e2d","observation_id":"8a8376b1-c190-4030-969f-293bda598d69","resolution":{"observed_at":"2026-07-02T11:16:53.408728Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1709.04326","last_updated":"2018-09-19T19:22:48Z","snapshot_observed_at":"2026-07-06T05:59:28.465615Z","submitted_at":"2017-09-13T13:42:15Z","title":"Learning with Opponent-Learning Awareness","version":4},"cited_work":{"arxiv_id":"1709.04326","doi":"10.48550/arxiv.1709.04326","metadata_source":"pith","pith_arxiv_id":"1709.04326","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Learning with Opponent-Learning Awareness","venue":"cs.AI","work_id":"61ef30eb-2271-4eee-9b6a-e01c6a4bcad7","year":2017},"citing_paper":{"arxiv_id":"2606.06486","last_updated":"2026-06-04T17:59:08Z","snapshot_observed_at":"2026-08-09T05:09:04.512605Z","submitted_at":"2026-06-04T17:59:08Z","title":"Regret Minimization with Adaptive Opponents in Repeated Games","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-06-28T02:02:27.201726Z"},"links":{"cited_paper":"/paper/1709.04326","citing_paper":"/paper/2606.06486"},"observation_digest":"sha256:3a7b2be77847967c3380180816816377fdc8539ae37c97e865d9ddb25bcbeea3","observation_id":"bf5e7661-810d-4023-a78b-be1ee30f7f41","resolution":{"observed_at":"2026-07-02T12:36:56.452029Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1709.04326","last_updated":"2018-09-19T19:22:48Z","snapshot_observed_at":"2026-07-06T05:59:28.465615Z","submitted_at":"2017-09-13T13:42:15Z","title":"Learning with Opponent-Learning Awareness","version":4},"cited_work":{"arxiv_id":"1709.04326","doi":"10.48550/arxiv.1709.04326","metadata_source":"pith","pith_arxiv_id":"1709.04326","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Learning with Opponent-Learning Awareness","venue":"cs.AI","work_id":"61ef30eb-2271-4eee-9b6a-e01c6a4bcad7","year":2017},"citing_paper":{"arxiv_id":"2606.27068","last_updated":"2026-06-25T14:14:24Z","snapshot_observed_at":"2026-08-07T22:43:21.947586Z","submitted_at":"2026-06-25T14:14:24Z","title":"Parametric Open Source Games","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-06-26T01:59:35.625359Z"},"links":{"cited_paper":"/paper/1709.04326","citing_paper":"/paper/2606.27068"},"observation_digest":"sha256:fe9681ff5d926dcab31de2c9eed9f33d4186e573147e417add615d3357362af5","observation_id":"479acccb-ec12-47a4-8455-bb95bf6b01ad","resolution":{"observed_at":"2026-06-26T02:08:54.893123Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1709.04326","last_updated":"2018-09-19T19:22:48Z","snapshot_observed_at":"2026-07-06T05:59:28.465615Z","submitted_at":"2017-09-13T13:42:15Z","title":"Learning with Opponent-Learning Awareness","version":4},"cited_work":{"arxiv_id":"1709.04326","doi":"10.48550/arxiv.1709.04326","metadata_source":"pith","pith_arxiv_id":"1709.04326","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Learning with Opponent-Learning Awareness","venue":"cs.AI","work_id":"61ef30eb-2271-4eee-9b6a-e01c6a4bcad7","year":2017},"citing_paper":{"arxiv_id":"2607.02292","last_updated":"2026-07-02T15:09:50Z","snapshot_observed_at":"2026-07-29T19:07:27.779678Z","submitted_at":"2026-07-02T15:09:50Z","title":"One More Time: Revisiting Neural Quantum States from a Reinforcement Learning Perspective","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-07-03T16:36:22.888448Z"},"links":{"cited_paper":"/paper/1709.04326","citing_paper":"/paper/2607.02292"},"observation_digest":"sha256:d027a86c81aa583750cb32e907345323658bf82f1202b12a4515cf41c3cdbe8e","observation_id":"71c0b7c2-12c0-4e17-8ec0-33074dfdb5cf","resolution":{"observed_at":"2026-07-03T16:38:39.668340Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/1709.04326/citation-record","integrity":"/paper/1709.04326/integrity","json":"/paper/1709.04326/citation-record.json","paper":"/paper/1709.04326"},"outbound":[],"paper":{"arxiv_id":"1709.04326","last_updated":"2018-09-19T19:22:48Z","latest_version":4,"primary_category":"cs.AI","snapshot_observed_at":"2026-07-06T05:59:28.465615Z","submitted_at":"2017-09-13T13:42:15Z","title":"Learning with Opponent-Learning Awareness"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"thesis":"As of 12 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 10 inbound Pith citation observations for arXiv:1709.04326."}