{"as_of":"2026-08-19T09:29:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:ed66f368a0a8858f57e45dc7f9a9aac305734c4a79b10f4b3500c92fee4787d4","coverage":[{"denominator":39,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":39,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T23:08:08.171638Z","state":"measured"},{"denominator":39,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":39,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-19T06:32:44.657259+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2505.06319/citation-record","integrity":"/paper/2505.06319/integrity","json":"/paper/2505.06319/citation-record.json","paper":"/paper/2505.06319"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.670833Z","title":"To compute or not to compute? adaptive smart sensing in resource-constrained edge computing,","venue":null,"work_id":"25fef14b-1421-4b10-9c58-aabdcab74d57","year":2022},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.014883Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:fdaf13a1176e0b5b52db58c7029df9ecdac21ff0db7e4fbc37c693612d5a6f7a","observation_id":"8078cd8b-08cc-44a7-9467-8884261c3028","resolution":{"observed_at":"2026-08-15T23:08:08.674467Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.660890Z","title":"Energy-efficient industrial internet of things in green 6g networks,","venue":null,"work_id":"7ee09423-7f60-4b19-970b-1e6054f6d904","year":2024},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.019771Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:7e9322375b47905b9b81ebe4db946562d851dac626ae88e999794e242b078c62","observation_id":"939c7a0f-ec29-480b-9094-e18c48269276","resolution":{"observed_at":"2026-08-15T23:08:08.664296Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.650070Z","title":"Cooperative air and ground surveillance,","venue":null,"work_id":"139c5d1a-c9c5-4db2-8492-f14415969875","year":2006},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.024170Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:9657761fa3e9f6883098fce5fcaf20c905d0549b776fa5521b6b05cbb87b35b1","observation_id":"92e02646-5344-4a42-8e83-22b0be4081f5","resolution":{"observed_at":"2026-08-15T23:08:08.654140Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.639697Z","title":"Planning for opportunistic surveil- lance with multiple robots,","venue":null,"work_id":"03d665ec-1825-4f02-bb84-cc7a3641e8d7","year":2013},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.030294Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:9d42edfc31f484a22c4c32c42b293df1442572985c33c3621d482d077050b5a0","observation_id":"897b99e2-0e79-4175-aa7e-4e233672cc3a","resolution":{"observed_at":"2026-08-15T23:08:08.643254Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2007.02204","last_updated":"2021-03-04T13:34:22Z","snapshot_observed_at":"2026-08-19T08:53:39.558164Z","submitted_at":"2020-07-04T23:03:38Z","title":"Failure-Resilient Coverage Maximization with Multiple Robots","version":3},"cited_work":{"arxiv_id":"2007.02204","doi":null,"metadata_source":"pith","pith_arxiv_id":"2007.02204","snapshot_observed_at":"2026-08-15T23:08:08.373367Z","title":"Failure-Resilient Coverage Maximization with Multiple Robots","venue":"cs.RO","work_id":"eca53b38-cc43-4eda-895f-9ecfeec88e68","year":2020},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.034622Z"},"links":{"cited_paper":"/paper/2007.02204","citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:5c791dcd4e51378ff797579df8463fe9409f8f1e0928c8e2da88a96ed16fb0f8","observation_id":"6f7d76b4-4bf8-4932-82db-bb58bca6678e","resolution":{"observed_at":"2026-08-15T23:08:08.377435Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.630566Z","title":"Game tree search for minimizing detectability and maximizing visibility,","venue":null,"work_id":"5ed03797-ada6-4a96-91f5-f591c4d2d699","year":2021},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.039164Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:9e48c6d725de33fc26485a91d2e26960386384db3503589c52f4651cab08332b","observation_id":"93f08083-574f-446c-9c40-e8518394f866","resolution":{"observed_at":"2026-08-15T23:08:08.633842Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.620202Z","title":"Stackelberg game approaches for anti-jamming defence in wireless networks,","venue":null,"work_id":"e34ab0b3-89f6-4c27-960f-e38af2faddcb","year":2018},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.043089Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:6f148755538262b9a666e482344c6a3e0e21cec9e4ec3bfc6ff71efc375fefda","observation_id":"6f4b1447-15ae-44b3-aa00-aa8ac68b5b1f","resolution":{"observed_at":"2026-08-15T23:08:08.623534Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.19868","last_updated":"2024-03-28T22:37:02Z","snapshot_observed_at":"2026-08-19T06:44:38.666514Z","submitted_at":"2024-03-28T22:37:02Z","title":"Jamming Intrusions in Extreme Bandwidth Communication: A Comprehensive Overview","version":1},"cited_work":{"arxiv_id":"2403.19868","doi":null,"metadata_source":"pith","pith_arxiv_id":"2403.19868","snapshot_observed_at":"2026-08-15T23:08:08.358664Z","title":"Jamming Intrusions in Extreme Bandwidth Communication: A Comprehensive Overview","venue":"cs.IT","work_id":"21a455f3-4aef-49b4-b4f1-fa43e94e9377","year":2024},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.047719Z"},"links":{"cited_paper":"/paper/2403.19868","citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:3c83fe0e3f1027d4b5b55871a60a0284522a79d64fb3d6da3926cc6b1e4f7065","observation_id":"80ed427f-98d2-4f22-b9ff-7f5cdc4c11c6","resolution":{"observed_at":"2026-08-15T23:08:08.362407Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.609205Z","title":"Equilibrium approximating and online learning for anti-jamming game of satellite communication power allocation,","venue":null,"work_id":"77795652-d34b-46b1-8826-a9eb157ef668","year":2022},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.052130Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:9531420dd7aa75bb303850159148616f879ddb4bdf74df85e2af2415e59dc918","observation_id":"d704a264-0bf7-40eb-9a90-4f1a5d858e2b","resolution":{"observed_at":"2026-08-15T23:08:08.612875Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1111.0380","last_updated":"2012-06-19T04:24:33Z","snapshot_observed_at":"2026-07-06T02:36:49.729256Z","submitted_at":"2011-11-02T04:23:41Z","title":"An Efficient Security Mechanism for High-Integrity Wireless Sensor Networks","version":2},"cited_work":{"arxiv_id":"1111.0380","doi":null,"metadata_source":"pith","pith_arxiv_id":"1111.0380","snapshot_observed_at":"2026-08-15T23:08:08.342854Z","title":"An Efficient Security Mechanism for High-Integrity Wireless Sensor Networks","venue":"cs.CR","work_id":"48e9da63-dbf3-4925-a40e-d312c98abb7a","year":2011},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.056591Z"},"links":{"cited_paper":"/paper/1111.0380","citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:b9202b6619247edafcad19212fc58a33b9a093f81b597bb66135c6b4cd30d4e1","observation_id":"fb770bfb-672e-494b-b57b-0865c3741079","resolution":{"observed_at":"2026-08-15T23:08:08.347710Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.13594","last_updated":"2022-10-08T20:42:13Z","snapshot_observed_at":"2026-08-16T16:50:00.482245Z","submitted_at":"2022-06-27T19:20:00Z","title":"Cyber Network Resilience against Self-Propagating Malware Attacks","version":2},"cited_work":{"arxiv_id":"2206.13594","doi":null,"metadata_source":"pith","pith_arxiv_id":"2206.13594","snapshot_observed_at":"2026-08-15T23:08:08.325203Z","title":"Cyber Network Resilience against Self-Propagating Malware Attacks","venue":"cs.CR","work_id":"b63fcd04-6758-4fad-8628-fb18a7cb37a5","year":2022},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.060878Z"},"links":{"cited_paper":"/paper/2206.13594","citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:99cd0860f03479ae1f3c817d98e54d123463fc73f8e79241c9161df985ef526e","observation_id":"8f9fd8c4-2f55-4f66-b725-05e9026881fd","resolution":{"observed_at":"2026-08-15T23:08:08.332156Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.586579Z","title":"Hibid: A cross-channel constrained bidding system with budget allocation by hierarchical offline deep reinforcement learning,","venue":null,"work_id":"6e3abd9c-8b5c-4c55-bec7-bbf6d612fd23","year":2023},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.069667Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:b1522811616bdc03eef078f95a756da6eea7c690068a55a8434b220c06539e9a","observation_id":"44f94000-91df-4697-a9ec-f1cc64622b34","resolution":{"observed_at":"2026-08-15T23:08:08.590537Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.573968Z","title":"Fairness-aware competitive bidding influence maximization in social networks,","venue":null,"work_id":"ca8e3b33-6fd7-4934-a0f3-f0c1a8e816a3","year":null},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.073400Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:57cbf4e7698323388542d75c11d08970985562d5244cb385367c925eb7f1ae68","observation_id":"a1da3851-32a0-49b4-bb08-36ee46e17f8c","resolution":{"observed_at":"2026-08-15T23:08:08.578748Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.551294Z","title":"La th ´eorie du jeu et les ´equations int ´egralesa noyau sym´etrique,","venue":null,"work_id":"b310b1a4-26f1-4a9f-a93a-70a958989ba0","year":1921},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.081295Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:d0d43653d448c0b1caafdc182e80d7c5d62bcad2e2151b47dcaf5c50d27b9c84","observation_id":"90a71a41-534b-4833-80cb-ee7bbaec63ee","resolution":{"observed_at":"2026-08-15T23:08:08.556110Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.540566Z","title":"The theory of play and integral equations with skew symmetric kernels,","venue":null,"work_id":"4c73a254-ee72-496e-a54d-ea6d5d512be0","year":1953},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.085626Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:a8df808b0a7dbf255be36163198a8cb28b01c2ffe814739dc2e1267734cb754a","observation_id":"fbab8c00-e0fc-4c8f-9294-a87a8c5e5e58","resolution":{"observed_at":"2026-08-15T23:08:08.544859Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.529720Z","title":"Dynamic defender- attacker blotto game,","venue":null,"work_id":"44dde234-fac4-4b70-bcdc-f1ebe108f8af","year":2022},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.089982Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:8c82caaf7a806e428de38ca8aef5f3fcabcf40523427104fa53b318368a0a348","observation_id":"0269ea6b-466e-4bb0-b28e-678193f096b9","resolution":{"observed_at":"2026-08-15T23:08:08.534308Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.517907Z","title":"Double oracle algorithm for game-theoretic robot allocation on graphs,","venue":null,"work_id":"4c3b5b7b-6723-4474-bc44-680d7ef4fc90","year":2025},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.093919Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:c321a35479de903500bfbbf0593b4a9a675b9803570a2d1db900ea52d5608eec","observation_id":"3e290042-4f67-4125-9df8-db2cbb355d4d","resolution":{"observed_at":"2026-08-15T23:08:08.521935Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.097763Z","title":"Mastering atari, go, chess and shogi by planning with a learned model,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.097763Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:85a417cfa034eb60a29316feb56184960d00e4044d9264ee3e54d0a6de960cfc","observation_id":"4ffe2c2f-8c31-4c60-8ff6-4988ede3d50a","resolution":{"observed_at":"2026-08-15T23:08:08.097763Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1712.01815","last_updated":"2017-12-05T18:45:38Z","snapshot_observed_at":"2026-08-14T20:06:37.819179Z","submitted_at":"2017-12-05T18:45:38Z","title":"Mastering Chess and Shogi by Self-Play with a General Reinforcement Learning Algorithm","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1712.01815","snapshot_observed_at":"2026-08-15T23:08:08.101217Z","title":"Mastering chess and shogi by self-play with a general reinforcement learning algorithm,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.101217Z"},"links":{"cited_paper":"/paper/1712.01815","citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:190c2f12f3ef82cf64109f09b55bc39533070fc919be10db24f591158d059522","observation_id":"b75b929a-43d5-427b-bc57-70cda0174ceb","resolution":{"observed_at":"2026-08-15T23:08:08.101217Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.500647Z","title":"A general reinforcement learning algorithm that masters chess, shogi, and go through self-play,","venue":null,"work_id":"551b4383-86da-4e7b-85d6-fbd9f12a43de","year":2018},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.105759Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:b8be412cde163adb8490d79880b6884b37aedc7ab1eabba0351e2879e4165fa7","observation_id":"5153838b-5f59-46e7-8e5e-9bebb326ec4c","resolution":{"observed_at":"2026-08-15T23:08:08.504809Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1509.01549","last_updated":"2015-09-14T15:42:35Z","snapshot_observed_at":"2026-08-14T22:33:43.656609Z","submitted_at":"2015-09-04T18:21:52Z","title":"Giraffe: Using Deep Reinforcement Learning to Play Chess","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1509.01549","snapshot_observed_at":"2026-08-15T23:08:08.109339Z","title":"Giraffe: Using deep reinforcement learning to play chess,","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.109339Z"},"links":{"cited_paper":"/paper/1509.01549","citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:ca0cb9ae1fe14292c3e0851059cf200b2ef14badb75c6d20842d918b41d09292","observation_id":"e67b1dd5-1492-4b3c-a99a-8347a9f43a4c","resolution":{"observed_at":"2026-08-15T23:08:08.109339Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.490619Z","title":"A game theory–reinforcement learning (gt–rl) method to develop optimal operation policies for multi-operator reservoir systems,","venue":null,"work_id":"6becf0f1-a638-43b7-ae4e-b5f61c9ff07f","year":2014},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.113343Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:eb9f8b41d87c4fdfd97d500c828af337b420d9dd3129c084625faa0ca42f4259","observation_id":"1e6da1ba-1760-4db5-97ab-67c07324542d","resolution":{"observed_at":"2026-08-15T23:08:08.493896Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.480700Z","title":"Applying reinforcement learning to small scale combat in the real-time strategy game starcraft: Broodwar,","venue":null,"work_id":"1621e1cb-0090-492c-9997-eca93a4784c3","year":2012},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.117411Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:61408fc85d6e9b43b354f0f3daaeae3076ad39879969c9b4272882be6bf7be1d","observation_id":"81e7a3ec-b841-4f19-9685-5efb691043d9","resolution":{"observed_at":"2026-08-15T23:08:08.484217Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.469720Z","title":"Towards playing full moba games with deep reinforcement learning,","venue":null,"work_id":"96a795a3-e250-43d4-9546-928bebc1ecc5","year":2020},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.121015Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:645b6db3fef2eb029f1d8f43b1b9cb5d5d619a3446c55882c9bf23d13aa18d7d","observation_id":"53875c94-760e-4ace-a62f-7c26bf90d8f2","resolution":{"observed_at":"2026-08-15T23:08:08.473868Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.459349Z","title":"Deep reinforcement learning based resource allocation for v2v communications,","venue":null,"work_id":"4b7345d7-2925-43fe-85ab-96e319c22b95","year":2019},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.124959Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:4b46dcca924dd8b46c4bd291c777e7c145f2e81aff766894943f417b9e3db43a","observation_id":"f355b079-dd60-4acd-be8e-c52b7d4e9edb","resolution":{"observed_at":"2026-08-15T23:08:08.462835Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.130152Z","title":"Multi-agent reinforcement learning- based resource allocation for uav networks,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.130152Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:7623a56825d5fcbd6eb88a7db41d2bc6e4c7fc5a7d555230c8a0c033d154f536","observation_id":"ff870b95-f205-4027-970f-fd5569c5cd51","resolution":{"observed_at":"2026-08-15T23:08:08.130152Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.441704Z","title":"Dynamical resource allocation in edge for trustable internet-of-things systems: A reinforcement learning method,","venue":null,"work_id":"4d7ee8c7-b8fd-4778-8e07-29c5675b07cd","year":2020},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.133838Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:db87ef3ba18117d924394fe6a9eca42e7c949f703228c24bdbf60bc26b1561fa","observation_id":"6db803ae-ab35-4a7b-a46c-0d9a105497de","resolution":{"observed_at":"2026-08-15T23:08:08.446406Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2204.02785","last_updated":"2022-04-04T16:18:01Z","snapshot_observed_at":"2026-08-16T17:09:24.978243Z","submitted_at":"2022-04-04T16:18:01Z","title":"Reinforcement Learning Agents in Colonel Blotto","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2204.02785","snapshot_observed_at":"2026-08-15T23:08:08.137476Z","title":"Reinforcement learning agents in colonel blotto,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.137476Z"},"links":{"cited_paper":"/paper/2204.02785","citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:7f159db6aa6038fb1dda4fe60b0956694c4abe37f037bca9da8b689b4fde271b","observation_id":"0a9a8c3c-cab8-4576-842b-aa2f2f5ec857","resolution":{"observed_at":"2026-08-15T23:08:08.137476Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.141179Z","title":"Human-level control through deep reinforcement learning,","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.141179Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:f8f327fa3893c71520f56b1ae4d9159c017268a7ba1c2a8b0b61c73bf2fb40e4","observation_id":"a759cba1-17ae-410f-9d72-47ee71356d3d","resolution":{"observed_at":"2026-08-15T23:08:08.141179Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06347","last_updated":"2017-08-28T09:20:06Z","snapshot_observed_at":"2026-08-15T20:26:32.102285Z","submitted_at":"2017-07-20T02:32:33Z","title":"Proximal Policy Optimization Algorithms","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.06347","snapshot_observed_at":"2026-08-15T23:08:08.145081Z","title":"Prox- imal policy optimization algorithms,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.145081Z"},"links":{"cited_paper":"/paper/1707.06347","citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:22334f45300fcc3119b67fef21b2c2b16c1c75233fbf616b86ab1012295d1fbb","observation_id":"c75d7822-df32-4894-b5dc-8364427f41ad","resolution":{"observed_at":"2026-08-15T23:08:08.145081Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.425020Z","title":"Adaptiveε-greedy exploration in reinforcement learning based on value differences,","venue":null,"work_id":"fed44a62-ba0f-45c5-bb9d-b277ffe695f0","year":2010},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.149584Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:a7272438f1a473cd690b2a2cdaabc06b3a9aa6b8466a6a2762b0f70d7f5bc6ed","observation_id":"1b7e7431-f07b-445f-9c2b-dc37b80c9221","resolution":{"observed_at":"2026-08-15T23:08:08.429305Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.152793Z","title":"Pettingzoo: Gym for multi-agent reinforcement learning,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.152793Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:8e6a3918dbbcd20c46bbb0b71742255a6529b967082a1d0ddf514024c4de59c1","observation_id":"52cff6be-cdf6-43d5-9918-c3f4c6a2182d","resolution":{"observed_at":"2026-08-15T23:08:08.152793Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.407106Z","title":"Tianshou: A highly modularized deep reinforcement learning library,","venue":null,"work_id":"f12d38fb-672c-436d-98e3-40e85d7f766d","year":2022},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.156637Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:b52b48caf4360c6655dead31aaa20a3de254a0d506a51fcd6aa57b47f3928d3e","observation_id":"59fa546e-b335-4723-938e-d3ae2f01d882","resolution":{"observed_at":"2026-08-15T23:08:08.411755Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.160193Z","title":"Deep reinforcement learning with double q-learning,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.160193Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:ceee80c47cdd2c9e672fb25bd94611569237a130e59197286f48e733b55bfa67","observation_id":"7bd26e77-355a-4ca4-a60e-437b3da98711","resolution":{"observed_at":"2026-08-15T23:08:08.160193Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.163929Z","title":"Dueling network architectures for deep reinforcement learning,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.163929Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:1d5d50290a19399771ca590f68279d982600ce3059e4bd0a79ae3ff47c8009e7","observation_id":"acc89d56-5f90-412d-9416-ca5a7eab9575","resolution":{"observed_at":"2026-08-15T23:08:08.163929Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1511.05952","last_updated":"2016-02-25T17:55:31Z","snapshot_observed_at":"2026-08-15T17:33:32.559096Z","submitted_at":"2015-11-18T20:54:44Z","title":"Prioritized Experience Replay","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1511.05952","snapshot_observed_at":"2026-08-15T23:08:08.167982Z","title":"Prioritized experience replay,","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.167982Z"},"links":{"cited_paper":"/paper/1511.05952","citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:d0205ebe53288dd20eec545922065c597e98751cdb0f16022849736d0cdbacc4","observation_id":"c271e3ab-3e1e-425d-889b-599a6be06987","resolution":{"observed_at":"2026-08-15T23:08:08.167982Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.384553Z","title":"A novel ddpg method with prioritized experience replay,","venue":null,"work_id":"e1d64460-fef4-4c99-be17-edecb15af680","year":2017},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.171638Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:8af044cea468110fe8212ce4f2bbdd37185e2cc25ed1c40e32cd7a9762c8c0b1","observation_id":"23c9ba41-f192-47a2-b6f7-cea203da4edd","resolution":{"observed_at":"2026-08-15T23:08:08.388455Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.598182Z","title":"Available: https://api.semanticscholar.org/CorpusId: 250089045","venue":null,"work_id":"0a84aa98-3553-4801-9cfc-8217cd12ff43","year":null},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.065671Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:d1cb6fc4652f38fbc49bd93cb8bba908656c1b6a87bf637524f0c239216cfec7","observation_id":"d6515bbe-c9aa-47cc-90b1-1b5db21f028a","resolution":{"observed_at":"2026-08-15T23:08:08.601483Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T23:08:08.562510Z","title":"Available: https://api.semanticscholar.org/CorpusId: 259539732","venue":null,"work_id":"805343ae-bdd8-42ab-a944-d855cef43def","year":null},"citing_paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-15T23:08:08.077867Z"},"links":{"citing_paper":"/paper/2505.06319"},"observation_digest":"sha256:1cdb0e1165f1135d499e603cc347824ec4d93ca129d49a04bb9361a8c9507def","observation_id":"635f8787-57cd-45b5-9530-6505252b6d04","resolution":{"observed_at":"2026-08-15T23:08:08.566601Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2505.06319","last_updated":"2025-05-08T21:12:34Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-15T22:58:03.957252Z","submitted_at":"2025-05-08T21:12:34Z","title":"Reinforcement Learning for Game-Theoretic Resource Allocation on Graphs"},"reference_resolution":{"displayed":39,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":11,"verified_exact":4,"verified_fuzzy":24},"total_outbound_references":39},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"thesis":"As of 19 August 2026, this Paper Citation Record lists 39 of 39 outbound references and 0 inbound Pith citation observations for arXiv:2505.06319."}