{"as_of":"2026-08-11T04:20:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:625cfef94492de6544153b3f68612599c728d0adcb1b772b8436fa9e3fe442dc","coverage":[{"denominator":29,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":29,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T20:27:04.535925Z","state":"measured"},{"denominator":29,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":29,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2508.10804/citation-record","integrity":"/paper/2508.10804/integrity","json":"/paper/2508.10804/citation-record.json","paper":"/paper/2508.10804"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:09.985356Z","title":"Whittle’s index policy for a multi-class queueing system with convex holding costs","venue":null,"work_id":"b4bdcc96-a551-422b-a6b2-ed7e99227fa7","year":2003},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:01.997873Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:3444caf1a31ecaec54f36ef566f268800ed0e3c741b94584d4de2616044c6062","observation_id":"e219ae47-4f7b-4d24-bdd4-15a4dfab993d","resolution":{"observed_at":"2026-08-05T20:27:10.088833Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:09.716010Z","title":"Dynamic allocation indices for restless projects and queueing admission control: a polyhedral approach","venue":null,"work_id":"6fba17ef-dfe8-4ed5-8cd3-a82ba5381d7b","year":2002},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.074608Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:836225d4675e59f6c1049a88ff577fa3de9c0ce7069f2cbf41a2ebd764e92f1c","observation_id":"57191e67-757e-47ab-a81c-20927c5d9c1f","resolution":{"observed_at":"2026-08-05T20:27:09.843742Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:09.479525Z","title":"Field study in deploying restless multi- armed bandits: Assisting non-profits in improving maternal and child health","venue":null,"work_id":"f723d386-1170-46a4-afd6-386604a6c0ec","year":2022},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.136881Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:de54a5615a421b3deced65291341b425c489e80200a4b493b2d569d00d431dd6","observation_id":"61adbc05-dd6d-496b-83aa-e821ae97042b","resolution":{"observed_at":"2026-08-05T20:27:09.581931Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:09.270612Z","title":"Collapsing bandits and their application to public health intervention","venue":null,"work_id":"26660420-3e32-4369-89dd-681c07f0d4b0","year":2020},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.205380Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:cba5ac1ab444c7b705689fdc9a2be9580640245133a6de67df0fb1107c7e09d8","observation_id":"5634b63f-55e1-4272-afdb-a302923553bc","resolution":{"observed_at":"2026-08-05T20:27:09.368472Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:09.018491Z","title":"Distributed optimal relay selection in wireless cooperative networks with finite-state markov channels","venue":null,"work_id":"36744545-133b-4338-b702-98965850ff3e","year":2010},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.269087Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:0bbb3bfd25c40386dbe49fdb9958d80ff50d25a7f4e8d0ebc55c669211df8a26","observation_id":"36e713a6-2481-4cb5-b13c-0b44269523c4","resolution":{"observed_at":"2026-08-05T20:27:09.143708Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:08.836333Z","title":"Cell association with user behavior awareness in heterogeneous cellular networks","venue":null,"work_id":"da1b66b4-b1f7-4963-a827-7395a57c7b3a","year":2018},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.338129Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:1feadd112f8e55d259ffb8a2e0f8587f9045b6d4cfed55d1b1e4209214bfc68b","observation_id":"6531ab99-6dca-4f7b-bccf-825758c92033","resolution":{"observed_at":"2026-08-05T20:27:08.917036Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:08.623770Z","title":"Optimistic whittle index policy: Online learning for restless bandits","venue":null,"work_id":"452b05e6-8ddd-418f-bcc7-e413f63be622","year":2023},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.408278Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:38a9de2656212c13a32047c17336f6aaf32411a7bbc3ee25f28d299ca874eea1","observation_id":"a863df58-7f4a-49c3-a485-c81a1ac39673","resolution":{"observed_at":"2026-08-05T20:27:08.723592Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:08.320710Z","title":"Reinforcement learning for non- stationary markov decision processes: The blessing of (more) optimism","venue":null,"work_id":"f1a69b91-eb31-494d-abca-249f17f2b2ac","year":2020},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.473859Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:6dce4196ac001eaaaefeee2ac0621428639b3197a9ebc8a7b9f55907665e4fc1","observation_id":"489a499a-f7a1-41d8-a302-59a33b08cc3f","resolution":{"observed_at":"2026-08-05T20:27:08.424887Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:08.121769Z","title":"Regret bounds for thompson sampling in episodic restless bandit problems","venue":null,"work_id":"f795ad79-e64e-47b1-8d2b-eacf348bb136","year":2019},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.566224Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:65ece339e66f4202af929313f31c577e858de0009d1ea00474e5aa59ac17206c","observation_id":"1ef18fa3-6690-455b-8141-7c7e02bd0c39","resolution":{"observed_at":"2026-08-05T20:27:08.210494Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1910.05654","last_updated":"2019-10-12T22:30:24Z","snapshot_observed_at":"2026-08-02T20:03:42.135513Z","submitted_at":"2019-10-12T22:30:24Z","title":"Thompson Sampling in Non-Episodic Restless Bandits","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1910.05654","snapshot_observed_at":"2026-08-05T20:27:02.648338Z","title":"Thompson sampling in non-episodic restless bandits","venue":null,"work_id":null,"year":1910},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.648338Z"},"links":{"cited_paper":"/paper/1910.05654","citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:61ba6eb67962d5770efdbd9b6781c6d1e220324620837fda908500116de0afd9","observation_id":"51d51be1-e07d-41b9-9b0f-d96e31b4ed55","resolution":{"observed_at":"2026-08-05T20:27:02.648338Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:02.738561Z","title":"Restless bandits: Activity allocation in a changing world","venue":null,"work_id":null,"year":1988},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.738561Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:838adafff486fd6c4beb1fc68acb596c2afe136b3615ea48657e7f61a36decad","observation_id":"95c17267-f2e2-44d0-9843-112aafd5f792","resolution":{"observed_at":"2026-08-05T20:27:02.738561Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:07.939720Z","title":"Introduction to multi-armed bandits","venue":null,"work_id":"e9a6dc67-c62d-4f9e-886b-d727b443a5d0","year":2019},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.810158Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:430d12909de1808a479c5fbd118574b7d0eccf6c92688c220b47bae9b54622ec","observation_id":"1a471b5c-7ae8-408f-aa8e-197778d1fcc4","resolution":{"observed_at":"2026-08-05T20:27:08.064518Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1402.6028","last_updated":"2014-02-25T01:34:43Z","snapshot_observed_at":"2026-07-06T03:36:51.545768Z","submitted_at":"2014-02-25T01:34:43Z","title":"Algorithms for multi-armed bandit problems","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1402.6028","snapshot_observed_at":"2026-08-05T20:27:02.916543Z","title":"Algorithms for multi-armed bandit problems.arXiv preprint arXiv:1402.6028, 2014","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.916543Z"},"links":{"cited_paper":"/paper/1402.6028","citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:45d23a2b78e123559a9e680b97a7e413cc32594a02bc64ef69711dcd488b4dbb","observation_id":"579c20fe-caee-4318-9d21-6dc70c8324f4","resolution":{"observed_at":"2026-08-05T20:27:02.916543Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:02.998565Z","title":"The complexity of optimal queuing network control","venue":null,"work_id":null,"year":1999},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:02.998565Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:3b6e850f6b9ec0f93a8de0ec3bc113533adfd966e5607936aad7a19083d6feb8","observation_id":"1d4e2415-cdb7-434e-a79d-f3ae2fa2cc28","resolution":{"observed_at":"2026-08-05T20:27:02.998565Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:07.722009Z","title":"On an index policy for restless bandits","venue":null,"work_id":"67ad0e29-f01b-4835-af1d-cd4d6470ca3f","year":1990},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.095698Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:3860fa42a7c8d0d3e0f53dd22e558fba19d4e24cb2d9d72b00c35358f38adc78","observation_id":"41cbc4cd-3155-420d-bd28-8e757f3872a2","resolution":{"observed_at":"2026-08-05T20:27:07.814489Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:07.503997Z","title":"A restless bandit formulation of opportunistic access: Indexablity and index policy","venue":null,"work_id":"5a6ca861-579f-49b9-bf3c-ed235d998ccb","year":2008},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.149894Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:5b827bf52849d15b5237f5355c4cb84818eaac17f384bce9c2eab277cc2aa0b5","observation_id":"fa63806b-c50f-40da-9175-ee81a4a4f7e7","resolution":{"observed_at":"2026-08-05T20:27:07.643212Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:07.231599Z","title":"Indexability of restless bandit problems and optimality of whittle index for dynamic multichannel access.IEEE Transactions on Information Theory, 56(11):5547– 5567, 2010","venue":null,"work_id":"62235149-0607-4221-8dfc-3ea272b371f2","year":2010},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.152865Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:de4934d89173b7396b57fd574df4a89c0633c13011189d05675a613a544a4e22","observation_id":"04aad0be-57c6-4ceb-9d86-78f7a112e136","resolution":{"observed_at":"2026-08-05T20:27:07.327107Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:07.039487Z","title":"Qwi: Q-learning with whittle index","venue":null,"work_id":"4038ff46-9dee-4bbf-8a72-d75977216c82","year":2022},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.186474Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:4d4ae1a7529b54a33cd7923dc5a36c4ca20de0df51525c7b4c26bb403c3763cc","observation_id":"3771e6c2-68db-4240-9077-3604d2a8ee23","resolution":{"observed_at":"2026-08-05T20:27:07.145281Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2105.07965","last_updated":"2021-07-22T22:37:04Z","snapshot_observed_at":"2026-07-06T11:10:10.904262Z","submitted_at":"2021-05-17T15:44:55Z","title":"Learn to Intervene: An Adaptive Learning Policy for Restless Bandits in Application to Preventive Healthcare","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2105.07965","snapshot_observed_at":"2026-08-05T20:27:03.273298Z","title":"Learn to intervene: An adaptive learning policy for restless bandits in application to preventive healthcare","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.273298Z"},"links":{"cited_paper":"/paper/2105.07965","citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:8f4d4f276c667be341e74461fe40f48d5d9b85c4471b9680fc34d30f3c333fe5","observation_id":"f69f6a6f-b8de-421c-a133-d25ac6296dcc","resolution":{"observed_at":"2026-08-05T20:27:03.273298Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:03.350381Z","title":"Towards q-learning the whittle index for restless bandits","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.350381Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:b391aa5488dcaf9a8cd2e09fa9b20c409ba8bf26efb32e74f8b5266c49fb19fa","observation_id":"96ce33c7-4b11-4ccf-9b88-4ea9e493c82a","resolution":{"observed_at":"2026-08-05T20:27:03.350381Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:06.792736Z","title":"Near-optimal regret bounds for reinforcement learning","venue":null,"work_id":"7228532c-b657-410b-a7f8-d9b5337406b3","year":2008},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.411313Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:6b85383e97cb4b1f354aaddf664d92b1fcdef72f43224e59de4d3f12a26f117f","observation_id":"93a38f69-b969-433e-97bb-cf564ef7329d","resolution":{"observed_at":"2026-08-05T20:27:06.923679Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:06.561096Z","title":"Logarithmic online regret bounds for undiscounted reinforcement learning","venue":null,"work_id":"8516cbda-a1e5-4887-8dc8-e706b661a832","year":2006},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.496746Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:6916528bfe0caaa0a9257ed25a39d636c193aac20f55d891fca6dca3a2ccb684","observation_id":"18741a19-acf4-45f8-aa35-a7167068442e","resolution":{"observed_at":"2026-08-05T20:27:06.669105Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1205.2661","last_updated":"2012-05-09T14:47:06Z","snapshot_observed_at":"2026-07-06T02:47:58.266745Z","submitted_at":"2012-05-09T14:47:06Z","title":"REGAL: A Regularization based Algorithm for Reinforcement Learning in Weakly Communicating MDPs","version":1},"cited_work":{"arxiv_id":"1205.2661","doi":null,"metadata_source":"pith","pith_arxiv_id":"1205.2661","snapshot_observed_at":"2026-08-05T20:27:05.152416Z","title":"REGAL: A Regularization based Algorithm for Reinforcement Learning in Weakly Communicating MDPs","venue":"cs.LG","work_id":"edba61d7-2d79-44ad-8daa-52c1358d4fac","year":2012},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.662892Z"},"links":{"cited_paper":"/paper/1205.2661","citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:c42aafffbbe1a890e06466a1caa9f7a425cfb0c505201631456232ecc2a20c22","observation_id":"c88f8d83-b79f-41fa-9a46-20deb7bd4c54","resolution":{"observed_at":"2026-08-05T20:27:05.256092Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:06.337524Z","title":"Non-stationary reinforcement learning without prior knowledge: An optimal black-box approach","venue":null,"work_id":"20eba94e-2d9a-489d-bf14-26f0879b7158","year":2021},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:03.815378Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:667451639b8e98ddecd48a8f857f80bd531b902fd67f372c324dfde609447751","observation_id":"5325003f-4bfd-45f6-9c7d-5bfc62280e21","resolution":{"observed_at":"2026-08-05T20:27:06.481285Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1805.10066","last_updated":"2018-05-25T10:14:20Z","snapshot_observed_at":"2026-07-06T06:41:13.674568Z","submitted_at":"2018-05-25T10:14:20Z","title":"A Sliding-Window Algorithm for Markov Decision Processes with Arbitrarily Changing Rewards and Transitions","version":1},"cited_work":{"arxiv_id":"1805.10066","doi":null,"metadata_source":"pith","pith_arxiv_id":"1805.10066","snapshot_observed_at":"2026-08-05T20:27:04.973913Z","title":"A Sliding-Window Algorithm for Markov Decision Processes with Arbitrarily Changing Rewards and Transitions","venue":"cs.LG","work_id":"2a6b0248-4688-4ec8-8b5f-b0689f71196e","year":2018},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:04.035674Z"},"links":{"cited_paper":"/paper/1805.10066","citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:cbeb31c2328212447c1aeaaff1c8bcac4eb5c9f41672a4d3c3ad04959ebb510a","observation_id":"de96fa00-8325-4043-88b8-9c607e7228b5","resolution":{"observed_at":"2026-08-05T20:27:05.038455Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:05.980928Z","title":"Near-optimal model-free reinforcement learning in non-stationary episodic mdps","venue":null,"work_id":"212e4708-4749-4e12-bcfa-df19509b99e2","year":2021},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:04.203518Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:0762b51f34c17fb52f4a9521281d875081c3ff409568f300a0a891c496e96670","observation_id":"0de02645-306b-4cd3-b529-8919694cc5a7","resolution":{"observed_at":"2026-08-05T20:27:06.110026Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:05.655111Z","title":"Learning in a changing world: Restless multiarmed bandit with unknown dynamics","venue":null,"work_id":"5ed97db3-c355-43f3-8b19-b4783724de84","year":1902},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:04.329643Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:8b50f82ded9182dcaea0ce3e2bcd2af0e2aa61e8c2664ddbc0090f995b737f3e","observation_id":"c11553c4-59ca-472e-ab2e-dfb59c929753","resolution":{"observed_at":"2026-08-05T20:27:05.833904Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2002.05138","last_updated":"2021-05-27T22:32:32Z","snapshot_observed_at":"2026-07-06T08:56:50.671205Z","submitted_at":"2020-02-12T18:30:09Z","title":"Regret Bounds for Discounted MDPs","version":3},"cited_work":{"arxiv_id":"2002.05138","doi":null,"metadata_source":"pith","pith_arxiv_id":"2002.05138","snapshot_observed_at":"2026-08-05T20:27:04.675708Z","title":"Regret Bounds for Discounted MDPs","venue":"cs.LG","work_id":"2b4642f7-bd2e-4ffb-a3e2-295d7f237732","year":2020},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:04.441499Z"},"links":{"cited_paper":"/paper/2002.05138","citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:5e19b3ab24ec24f35587d2a5b04a134d2dcca7d2fffc0b53f838cee62c3684e2","observation_id":"0ba162fe-673c-4ab2-8097-ecca8b779c80","resolution":{"observed_at":"2026-08-05T20:27:04.837991Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T20:27:05.428822Z","title":"X t∈N γt−1 R(st, at) − λ∗ t πt(st)⊤1 − K | s1 = s # (31) ≤ E(s,a)∼(P ,πt)","venue":null,"work_id":"de035725-9ffa-4c4b-a600-e0fe835ecf82","year":1982},"citing_paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-05T20:27:04.535925Z"},"links":{"citing_paper":"/paper/2508.10804"},"observation_digest":"sha256:0a4b086f8b21e3a5689ab0aa13fbefc63a0dbb0cf3c75f21042777cb1dbb826b","observation_id":"d81409f8-2fdf-4991-bbee-084a78a9f3b2","resolution":{"observed_at":"2026-08-05T20:27:05.546835Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2508.10804","last_updated":"2025-08-14T16:26:00Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-10T17:02:50.595747Z","submitted_at":"2025-08-14T16:26:00Z","title":"Non-Stationary Restless Multi-Armed Bandits with Provable Guarantee"},"reference_resolution":{"displayed":29,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":6,"verified_exact":3,"verified_fuzzy":20},"total_outbound_references":29},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 11 August 2026, this Paper Citation Record lists 29 of 29 outbound references and 0 inbound Pith citation observations for arXiv:2508.10804."}