{"as_of":"2026-08-16T21:24:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:22322f8a3032f996e024c20802c64d4caeccee9c4c7e1072c720ea0b8099e3e4","coverage":[{"denominator":57,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":57,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T21:36:19.656167Z","state":"measured"},{"denominator":58,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":58,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-16T06:30:59.297886+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-29T10:31:29.175140Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-06-29T10:33:18.779573Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"cited_work":{"arxiv_id":"2502.09432","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.09432","snapshot_observed_at":"2026-06-29T10:33:18.779573Z","title":"Dual formulation for non-rectangularL p robust Markov decision processes.arXiv preprint arXiv:2502.09432, 2025","venue":null,"work_id":"9d2d3e06-65a4-461b-8ce7-354ba2f43d23","year":2025},"citing_paper":{"arxiv_id":"2605.28706","last_updated":"2026-05-27T16:32:02Z","snapshot_observed_at":"2026-08-12T20:58:54.067670Z","submitted_at":"2026-05-27T16:32:02Z","title":"Robust Markov Decision Processes on Continuous State Spaces","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-06-29T10:31:29.175140Z"},"links":{"cited_paper":"/paper/2502.09432","citing_paper":"/paper/2605.28706"},"observation_digest":"sha256:376af55299fb2d3e53164a84efac96cfd5e1d92966a168978925ed7ba75d0d9b","observation_id":"233a9c31-c3aa-46fd-b70b-debdf5215a43","resolution":{"observed_at":"2026-06-29T10:33:18.781161Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2502.09432/citation-record","integrity":"/paper/2502.09432/integrity","json":"/paper/2502.09432/citation-record.json","paper":"/paper/2502.09432"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.254324Z","title":"Tsitsiklis","venue":null,"work_id":"924aea24-00e2-44f4-a323-172627c24845","year":2004},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.458006Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:1c7f77054e2e4548a4580331a26cb5954022c3e042f8e2b3cf9af74176ead13f","observation_id":"0b6c4991-7123-41f6-87a7-945dd0b2482a","resolution":{"observed_at":"2026-08-07T21:36:20.257433Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.244833Z","title":"Robust data-driven dynamic programming","venue":null,"work_id":"646089be-dd26-4994-9340-8935d5c1b5ad","year":2013},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.462035Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:5a66cf6b7c678352eef1a3f1da5c9c974743d24a9a9c304090016f84995caadd","observation_id":"d068d777-fdf7-4f58-bc22-82cf9e89f98e","resolution":{"observed_at":"2026-08-07T21:36:20.248258Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.235289Z","title":"Scaling up robust mdps using function approxi- mation","venue":null,"work_id":"49248642-fb26-4154-9d27-1fe6dbfbbb41","year":2014},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.465605Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:b08fd09c5ad4b4a1e6dfdbbded098d3d425dd3d3f77c65c029f737ea1c4c093a","observation_id":"aa9b5782-8f78-4873-bde4-cda8b7a91372","resolution":{"observed_at":"2026-08-07T21:36:20.238739Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.224937Z","title":"Robust control of markov decision processes with uncertain transition matrices.Oper","venue":null,"work_id":"2fb42fa6-1869-4f5e-92ea-5ccac01395d1","year":2005},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.469142Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:b05364a23ced8ec8b0de9b451321ec600b946c16070ac5058a21ccae0613741e","observation_id":"afe40254-2806-4775-9ab3-b7a178f8b8ff","resolution":{"observed_at":"2026-08-07T21:36:20.228532Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.215253Z","title":null,"venue":null,"work_id":"902221ca-27bf-4dbb-b429-a6c4c8670b85","year":2005},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.472591Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:8d48673e9eb8a330578318e457cd3f7aceb80798400eb7cc337030e10df50edf","observation_id":"c81ac82f-9975-427a-b930-af27340dcda7","resolution":{"observed_at":"2026-08-07T21:36:20.218580Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.204756Z","title":"Robustness and generalization, 2010","venue":null,"work_id":"ac4d4f85-5a33-4e73-a2bc-c0ade71e7b4e","year":2010},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.476160Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:843d8e90b01cae840276374a50041041312318a88580922fc331ce7cb5403ea0","observation_id":"419da8df-17b9-43d8-a4fb-fb994e518124","resolution":{"observed_at":"2026-08-07T21:36:20.208490Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.195207Z","title":"Hospedales","venue":null,"work_id":"9cdf2926-4975-416f-b4a8-0c10a940c67d","year":2019},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.479899Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:5f169f99f67bd732e80454b49d68b6548bbd487c2a04c76cf57b779f870b1eb2","observation_id":"4a825354-3a52-422d-9f43-e5ae315287a0","resolution":{"observed_at":"2026-08-07T21:36:20.198388Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.185994Z","title":"Assessing generalization in deep reinforcement learning, 2018","venue":null,"work_id":"585fb248-894b-4e29-878b-0baf9b409362","year":2018},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.483166Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:4e9a9c0aba4901d2f8a7442b3c9b3c3b2426bf7466617d589b320d934783119d","observation_id":"5b6896b6-866b-41e5-9487-998733ad90b7","resolution":{"observed_at":"2026-08-07T21:36:20.189184Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.176556Z","title":"Robust markov decision processes","venue":null,"work_id":"bce5c98c-1ecf-4e95-94a1-090c12d237f9","year":2013},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.486850Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:c12ac49ecdb5d0a0afd51756d82b83bfb5e80e31d0ff725d7f667d8595fba496","observation_id":"5775abf2-2ba1-44f5-9c20-2c4a0fda77f2","resolution":{"observed_at":"2026-08-07T21:36:20.179762Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.166823Z","title":"Robust mdps with k-rectangular uncertainty","venue":null,"work_id":"6b217a33-67bf-45ba-a16d-ae6b97fa6661","year":2016},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.490118Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:9a90b8f586a5a2e9dee94e09e621b7e51cd86f938a09cdaf9be20ca27e5fe50d","observation_id":"faf38675-618d-4f52-be99-3259c365fd6e","resolution":{"observed_at":"2026-08-07T21:36:20.170486Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.156106Z","title":"Robust markov decision process: Beyond rectan- gularity, 2018","venue":null,"work_id":"6ecf8951-72d5-4d5d-b676-e9866d1e65f5","year":2018},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.493480Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:18484ac32c49fbecd4ef7ddac5ad561fcf2930be9f92e72fb8e4b0877a26e47b","observation_id":"6958771c-17eb-476e-969f-82f69e30975c","resolution":{"observed_at":"2026-08-07T21:36:20.159650Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.145337Z","title":"Kaufman and Andrew J","venue":null,"work_id":"d43cc9ad-0c1f-4e58-b4ac-73f29e54b796","year":2013},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.496750Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:ed5616bba5b8f67cb192da7302b6a5984bd8396a0527158f5b9dc9c62fce2ca9","observation_id":"50b4c15e-7a13-4973-ab19-404a9d6c2dc2","resolution":{"observed_at":"2026-08-07T21:36:20.149285Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.135083Z","title":"Andrew Bagnell, Andrew Y","venue":null,"work_id":"eb086a33-92aa-4328-a541-1ace4dade09b","year":2001},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.499998Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:888929497eda097ead0e28ec555a8a134f0d0ee00971ec317e75baeac1ae1095","observation_id":"819598ad-e2ed-4b48-bafb-922795164d34","resolution":{"observed_at":"2026-08-07T21:36:20.138685Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.124871Z","title":"Partial policy iteration for l1-robust markov decision processes, 2020","venue":null,"work_id":"828fee52-7404-4ee3-bc67-dc33dd9e96a3","year":2020},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.503457Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:713b6167254b05ef58224b18ef4134b1136c87bb99d64222fe9305fad291e165","observation_id":"bf527111-eaf5-481b-b603-8ffde6c430e2","resolution":{"observed_at":"2026-08-07T21:36:20.128774Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.115129Z","title":"Online robust reinforcement learning with model uncertainty, 2021","venue":null,"work_id":"70602bcf-5c33-4ca3-95c9-bde60e97dc14","year":2021},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.506765Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:35abcd94396ffa78d5b9910e056461af6ccf5f0496a0e6ea38d6d286a8930d47","observation_id":"dcd86abc-f994-4f57-93c5-48c9c9c0cfb2","resolution":{"observed_at":"2026-08-07T21:36:20.118568Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.510193Z","title":"Policy gradient method for robust reinforcement learning, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.510193Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:7602256600ce0f30015bbc1a754a20a19c63f0c005d7f8fcbedac0b2de2507c1","observation_id":"6841709c-70a0-49ff-950e-7ce3ca8ea7ed","resolution":{"observed_at":"2026-08-07T21:36:19.510193Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.099413Z","title":"Policy gradient in robust mdps with global convergence guarantee, 2023","venue":null,"work_id":"fb27e390-1a18-4024-9039-1f33673055a5","year":2023},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.513827Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:1f7eeb459708de30448f8057fe3b96be3b0ce030f7f0ddb3e7bc9ad3d68193f9","observation_id":"0430343d-875d-49b4-a815-0971c3ddfbd0","resolution":{"observed_at":"2026-08-07T21:36:20.102796Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.089916Z","title":"Twice regularized mdps and the equivalence between robustness and regularization, 2021","venue":null,"work_id":"91477ee8-b85f-408e-8c59-d00377e76fc3","year":2021},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.517278Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:533c6bf5a87c64ba2580c018507fd0ac8382808832a638d79922be83f0d855b7","observation_id":"6066d4e6-4f32-416a-b875-dc1bf49c8064","resolution":{"observed_at":"2026-08-07T21:36:20.093368Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.079913Z","title":"Efficient value iteration for s-rectangular robust markov decision processes","venue":null,"work_id":"270e0d41-c4dd-43de-baa4-aa8e71cdf1dd","year":2024},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.520560Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:2b9eef8971c1019a99495314750e782e5cb0852e4e363f85b4bcba13a9e2d8ce","observation_id":"b1e7768d-3351-4017-9ac6-fb82a947672d","resolution":{"observed_at":"2026-08-07T21:36:20.083565Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.069427Z","title":"Pol- icy gradient for rectangular robust markov decision processes","venue":null,"work_id":"e5432607-28e7-48c6-b0f5-8ed295db7d66","year":2023},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.523939Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:393fb5b93ed4a65b3377634215e782afc4a1eb21265bd6e0d57cb4c73426eb16","observation_id":"2fca1ffd-6e9d-4eda-96ab-fcf794ce255d","resolution":{"observed_at":"2026-08-07T21:36:20.073066Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.059745Z","title":"Natural actor-critic for robust reinforcement learning with function approximation","venue":null,"work_id":"1d59ebc2-9c4a-45cd-84ca-de2602283e86","year":2023},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.527335Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:f426843a237254df40ae59aad6376860773a957b87248e1e07aa9993bbccf1ca","observation_id":"580cf6d2-91f7-4995-a432-781b475e0600","resolution":{"observed_at":"2026-08-07T21:36:20.063185Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.049699Z","title":"Robust reinforce- ment learning via adversarial kernel approximation, 2023","venue":null,"work_id":"2fe5915e-d209-45e0-8e1e-8105ef9693b9","year":2023},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.530518Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:28855be7e16311b2451ef4be00bf1996e5daec8e360b68b943f8fd0043ed484b","observation_id":"9f7ec817-9d0e-4894-bcf8-b6677202d6d0","resolution":{"observed_at":"2026-08-07T21:36:20.053423Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.039985Z","title":"Solving non-rectangular reward-robust mdps via frequency regularization, 2023","venue":null,"work_id":"442733f2-b51c-440e-806d-824eb22d6ffb","year":2023},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.533981Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:36f7e63807fa4434eb8a4e2e7cda5e1dd1c879a2db7c26b9c3cb561be5c53cc7","observation_id":"ddc55e57-2f07-4962-95fa-100a13f7354c","resolution":{"observed_at":"2026-08-07T21:36:20.043450Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.030331Z","title":"Smith and Mavina K","venue":null,"work_id":"99ef9a7b-2664-4344-b941-c59512519bc4","year":1989},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.537530Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:a4a94e51af42367f27bccfe524d5422329beefeefefcb664d715f8995077833c","observation_id":"d571b47b-d0d0-4044-b84c-e6ae4990628b","resolution":{"observed_at":"2026-08-07T21:36:20.033705Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.541038Z","title":"Puterman","venue":null,"work_id":null,"year":1994},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.541038Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:da6d26d9f4482ffd8f0489a72539788f43425a68eb1ca3e2b6474d1b63f8b221","observation_id":"fc0ad112-f3c5-4d6d-a40c-cb072c396f7f","resolution":{"observed_at":"2026-08-07T21:36:19.541038Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:20.014803Z","title":"Tractable robust markov decision processes, 2024","venue":null,"work_id":"2f8121ad-97d1-47fe-b755-5bdec50bfae9","year":2024},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.544428Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:28e996fc8ba3d2a8e941da19085a58eb43d6bbd341228978d75d2c9e54e13c52","observation_id":"61788557-f5ed-472d-937f-067bad6b0743","resolution":{"observed_at":"2026-08-07T21:36:20.018254Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.547831Z","title":"Sutton and Andrew G","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.547831Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:ded6ff04ab51013463ff7e486d5412dd6f5860b50ca843af104e2380e9ecd75e","observation_id":"c4da2e3f-ba0f-4e02-88f3-03864ea9b1f8","resolution":{"observed_at":"2026-08-07T21:36:19.547831Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.999235Z","title":"Wasserstein robust reinforcement learning, 2019","venue":null,"work_id":"c86ded1f-52ba-4646-ae07-b376b46c4d5f","year":2019},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.551215Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:d8a7e272b593a9626d534ad7cdbe57570e7ccde27a54406a625077346f7e316a","observation_id":"ebb2794d-4a82-4ad2-b555-fc984540d777","resolution":{"observed_at":"2026-08-07T21:36:20.002741Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.989232Z","title":"Robust $\\phi$-divergence MDPs","venue":null,"work_id":"4b4b680d-3c02-425c-96de-22a7f0db02ed","year":2022},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.554703Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:cd2c536025418b1b94a5bc17c35af07448019e646f97a3b49d63781d150101b1","observation_id":"aaba2e2f-1859-4405-9b5d-cf5a4484971a","resolution":{"observed_at":"2026-08-07T21:36:19.992855Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.978274Z","title":null,"venue":null,"work_id":"69722877-f65e-4f4b-aba0-3c47535789a4","year":1999},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.558054Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:c70717a75f1b44545d92b99be79a78439ac1f4ee3cc45e77f3eae39fb17399bf","observation_id":"7f9ed6a9-a3a2-43c5-aacd-930eab901a5e","resolution":{"observed_at":"2026-08-07T21:36:19.982063Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.561426Z","title":"Bellemare","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.561426Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:fb79ba4a8bbf382aa4c135bee78d306b5d6ff3eabd76895d50aceddab1df451e","observation_id":"64eb255c-ac09-4a74-a887-d4f8e0058f25","resolution":{"observed_at":"2026-08-07T21:36:19.561426Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.960603Z","title":"The geometry of robust value functions","venue":null,"work_id":"ff7a8997-b7de-4104-990f-7d78dbaa24c5","year":2022},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.564840Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:58c18fe4e79acbe589ae419aa3c1e9d398a44ae666c7829697d145dcfb177ba6","observation_id":"039ddbee-795b-408e-a57e-f8aa5fd4e736","resolution":{"observed_at":"2026-08-07T21:36:19.964291Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1206.4643","last_updated":"2012-06-18T15:19:07Z","snapshot_observed_at":"2026-08-16T11:47:01.971177Z","submitted_at":"2012-06-18T15:19:07Z","title":"Lightning Does Not Strike Twice: Robust MDPs with Coupled Uncertainty","version":1},"cited_work":{"arxiv_id":"1206.4643","doi":null,"metadata_source":"pith","pith_arxiv_id":"1206.4643","snapshot_observed_at":"2026-08-07T21:36:19.687000Z","title":"Lightning Does Not Strike Twice: Robust MDPs with Coupled Uncertainty","venue":"cs.LG","work_id":"afb31fdf-8703-49c9-a047-ba1bfce0c0db","year":2012},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.568498Z"},"links":{"cited_paper":"/paper/1206.4643","citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:54baf12cc1dc5c54c4ba4c06dfb19f691d47c19a91def26465536055a3b95703","observation_id":"c2934878-e35f-4b07-b3aa-3d54fe13f949","resolution":{"observed_at":"2026-08-07T21:36:19.692742Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.949205Z","title":null,"venue":null,"work_id":"dba94ccc-2f2c-459c-ae0e-dc0bd47b16ee","year":1951},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.572528Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:0d67c560bce331fecac6c8eb7cd5388526234cced8b525a3972ff00233db86ec","observation_id":"87e096e3-176b-4a56-8a34-29a9890f7309","resolution":{"observed_at":"2026-08-07T21:36:19.952879Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.938444Z","title":"Oliphant, Matt Haberland, Tyler Reddy, David Cournapeau, EvgeniBurovski, PearuPeterson, WarrenWeckesser, JonathanBright, StéfanJ","venue":null,"work_id":"67a2b74a-e8fb-4c31-9f0f-356fa09c65e5","year":2020},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.575826Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:f724d649a5525e11d94f6192a4d5b9bfdd024490215986446f012396a7787baf","observation_id":"fc595b09-c6d7-4e2e-8f9a-2a2ce5757e56","resolution":{"observed_at":"2026-08-07T21:36:19.942183Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.579273Z","title":"Policy gradient methods for reinforcement learning with function approximation","venue":null,"work_id":null,"year":2000},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.579273Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:df7bdccce348df6a9bd2c3b0567dc3b6cd7b752342998e0919e5b57cf05131cf","observation_id":"649bac79-e0ab-4eff-af3e-0a715a06007b","resolution":{"observed_at":"2026-08-07T21:36:19.579273Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.920402Z","title":null,"venue":null,"work_id":"fa378d88-3dcf-42d6-91bc-549498c9f5f6","year":1979},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.583088Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:875172072b6f4bc33eae377d69988d147017d774d350365cff9538a2761cde72","observation_id":"8aa4f6c6-9342-45fa-b733-83f9d17b8dbb","resolution":{"observed_at":"2026-08-07T21:36:19.923924Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.909489Z","title":"Cambridge University Press, March 2004","venue":null,"work_id":"9f5a7a6b-fa5c-498b-a746-d580d1765601","year":2004},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.586820Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:b60a9e2ebed946e8aae3fc4b76dd633a564c905984fb1f0a25f1c92219a4e249","observation_id":"0ce3b136-52f1-4a87-beaa-b93639cba9ad","resolution":{"observed_at":"2026-08-07T21:36:19.913235Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.898428Z","title":"Policy gradient for reinforce- ment learning with general utilities, 2023","venue":null,"work_id":"60935f97-e954-404d-b8c1-f86e1d984db8","year":2023},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.590353Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:931650e5c783cf5144b6b62329d3882b79446fd8aee89ade8211d71faacb3241","observation_id":"dddb15a4-ac21-499f-874e-573f4b4bf369","resolution":{"observed_at":"2026-08-07T21:36:19.902250Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.887787Z","title":"This makes sense, as the more the agent visits states with high uncertainty, the higher is the ability of the adversary to undermine it","venue":null,"work_id":"5aa78691-e3ce-4b5d-90c0-fa8f2d92aaec","year":null},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.594133Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:8f621853f7529a08ef8eb91f960b4de8bc31e808b3a860c77e9713ac6a37a5ca","observation_id":"9def3c86-ab34-4660-9448-3c3caaf52cd8","resolution":{"observed_at":"2026-08-07T21:36:19.891391Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.877379Z","title":"nominal value functionvπ R)","venue":null,"work_id":"408cfd9a-ca88-4950-9933-ed7d8b0f7e21","year":null},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.598039Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:cf6e9d0298cf34db6e3ecea51b1c039c92c7fe1d896604e6da5c582686e8e415","observation_id":"553b68c4-1f7c-4cdb-aac4-0c09dcfa1fac","resolution":{"observed_at":"2026-08-07T21:36:19.880758Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.866547Z","title":"This can be done, by putting negative entries ofk at maximal entries of vπ β, and positive entries ofkT at states which has minimum uncertainty value function","venue":null,"work_id":"94cd6e2b-84e4-4684-91ce-99fbd05bce09","year":null},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.601727Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:4f60afa1f6880d7a22c1f2a5aa281050626b45693c1d63abcda253167d0e3ef9","observation_id":"506d608b-819a-4519-9bb8-740b6421e2f5","resolution":{"observed_at":"2026-08-07T21:36:19.870305Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.856093Z","title":null,"venue":null,"work_id":"c2ad332b-a982-426a-9506-c8d2daa64f17","year":null},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.605716Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:0564a604a7b8c3672315d7d580809b0a09ef24cde366c13ea260b43a8c894058","observation_id":"dd773552-dec4-43b5-89e0-b9bbb083db60","resolution":{"observed_at":"2026-08-07T21:36:19.859571Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.845295Z","title":null,"venue":null,"work_id":"4e0fd0cd-df65-4eab-9944-20ccfb30d012","year":null},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.609273Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:5c244302b6a7b21716fa9267e825b3db256edfa4ad4f1825aff609c60641f77d","observation_id":"422fa3d0-5881-494f-ace0-786c9f567779","resolution":{"observed_at":"2026-08-07T21:36:19.849233Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.834997Z","title":null,"venue":null,"work_id":"b99e88d9-4a4c-468a-aee5-41bf647cf41f","year":null},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.612680Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:21489fc26238003c59b2a8c202c2e1f6ae036af3d6e1bcc8d81d5569ceb0c9db","observation_id":"b138513d-1381-4031-9656-dc3b6ae76802","resolution":{"observed_at":"2026-08-07T21:36:19.838287Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.824807Z","title":null,"venue":null,"work_id":"cdc44af9-ee96-4c90-9bed-bf06f35e3378","year":null},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.617065Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:62ea3b769ab9fb8f44ddc1e882c061c9fa7ba6c91a7f87d5433151ed9cb43d15","observation_id":"2edeacfe-e812-409b-92de-aaa3a4193d7d","resolution":{"observed_at":"2026-08-07T21:36:19.828174Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.813925Z","title":null,"venue":null,"work_id":"e113908d-a50c-4bf0-aa93-2b640344a9f8","year":null},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.620603Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:88261b8017b72de27540f744567cc26a4f469902f9c42dabacbfc3e5aad45821","observation_id":"7c61d827-b01b-4157-b63f-9e7d2d5e4390","resolution":{"observed_at":"2026-08-07T21:36:19.817675Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.803231Z","title":"Robust policy evaluation is proven to be NP-hard for general uncertainty sets defined as intersections of finite hyperplanes [9]","venue":null,"work_id":"c986d9b7-2075-4233-b8f9-0840f30610f3","year":null},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.623783Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:585bd5ec8ea7096c8dcdd960e1448ccda2db5bde334b579db699b5a4328b17e1","observation_id":"6083aa55-ce62-4c26-95b0-2ae338ed70ce","resolution":{"observed_at":"2026-08-07T21:36:19.806953Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.792702Z","title":null,"venue":null,"work_id":"ed8678a8-8eb9-4f80-adf0-9db2416945d6","year":null},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.628677Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:c2c0b5a934eb7a3966ae2e2439760ed9d7f20615fce6ef760f1bea64cc787ac6","observation_id":"a6dc41b7-5135-4578-9cbd-afb4f0b67f3d","resolution":{"observed_at":"2026-08-07T21:36:19.796328Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.781766Z","title":null,"venue":null,"work_id":"02fa319f-3560-4de9-b947-39fde01095c5","year":null},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.632037Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:a62bc769d1fc2eee847a8d7ed412b4612e297f8731805da405362e60e617435e","observation_id":"ec46a7bd-bc95-4f9d-907f-cb128b6f19b9","resolution":{"observed_at":"2026-08-07T21:36:19.785350Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.770227Z","title":"This method is simple to implement but computationally expensive, as it evaluatesA for a large number of randomly generated vectors","venue":null,"work_id":"af9fc89c-6a59-4987-aee6-797d184a11be","year":null},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.635470Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:1c5866a75b19ad267003b2451a76f6f0033aec6d035ca185eb8b226c685a10c3","observation_id":"30afade0-ed22-44d1-97b1-8630032d946e","resolution":{"observed_at":"2026-08-07T21:36:19.774632Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.758856Z","title":null,"venue":null,"work_id":"66fb1346-8ca9-4a47-b059-f8570fb4b7b9","year":null},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.638873Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:b66680926ef6ae8a1767b687aa182e14a8d30cfdf46b214910991ac96316747f","observation_id":"c2615a3f-5157-4667-a01d-4b66477e2297","resolution":{"observed_at":"2026-08-07T21:36:19.762906Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.746758Z","title":null,"venue":null,"work_id":"0f9f5436-2d18-4e20-83b9-4f0f13e902b5","year":null},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.642236Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:fc3daa5c88d24a325237e6b2c6e2d8d77899d54a4d7a06cb05ef7a696f309334","observation_id":"1f70338b-74cb-460f-8826-c216989c711e","resolution":{"observed_at":"2026-08-07T21:36:19.750523Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.735240Z","title":"This method provides the exact solution but is computationally more expensive than the spectral method","venue":null,"work_id":"fc00a9fa-4dd4-4fa9-be97-381f6c2bd5b0","year":null},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.645535Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:cda851bd59dd5bdfbe301def8fd9b1c2c0917647606d3009226b3dcff2d41f86","observation_id":"68957e29-950d-41b5-8977-65ddcfb75bb0","resolution":{"observed_at":"2026-08-07T21:36:19.739160Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.723929Z","title":null,"venue":null,"work_id":"8dfb694d-4cdc-433d-95f8-c673cc0cc565","year":null},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.648971Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:371c8e4116576a006baa3a4df664bdf73dd70e882b1342746e05b2fb88a3e147","observation_id":"78aee9ed-86aa-46ea-97d7-7c98b435ce26","resolution":{"observed_at":"2026-08-07T21:36:19.727520Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.712524Z","title":null,"venue":null,"work_id":"6a00dbb5-69e0-44fd-b9fb-de21c9e57936","year":null},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.652497Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:6b795f41a4a8f7e475912604aa49f263cd71e0996963f27429ea8a2c82436c24","observation_id":"9235abd9-1024-4bd0-b9ab-3a162974a240","resolution":{"observed_at":"2026-08-07T21:36:19.716000Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T21:36:19.701227Z","title":"The results, including the optimal values and computational times, are recorded for each method","venue":null,"work_id":"03294bdf-6afc-4ddf-8092-4f0816b08a5d","year":2023},"citing_paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-07T21:36:19.656167Z"},"links":{"citing_paper":"/paper/2502.09432"},"observation_digest":"sha256:93b9efd38dea0af9e40980ef70b7c047a7aca20c3cd782c854d52ea9d7e11f50","observation_id":"a0231bc6-ba6b-48ac-9920-1d824e750f52","resolution":{"observed_at":"2026-08-07T21:36:19.705050Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2502.09432","last_updated":"2025-02-13T15:55:00Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-14T08:29:03.785552Z","submitted_at":"2025-02-13T15:55:00Z","title":"Dual Formulation for Non-Rectangular Lp Robust Markov Decision Processes"},"reference_resolution":{"displayed":57,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":20,"verified_exact":1,"verified_fuzzy":36},"total_outbound_references":57},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"thesis":"As of 16 August 2026, this Paper Citation Record lists 57 of 57 outbound references and 1 inbound Pith citation observation for arXiv:2502.09432."}