{"as_of":"2026-08-16T15:01:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:06e3b9d8f4d114f0b95e764c46d234ea8374bbc7183cccaccee522f284d2f6ed","coverage":[{"denominator":49,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":49,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T20:36:20.447781Z","state":"measured"},{"denominator":49,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":49,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-16T06:30:59.297886+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2608.10634/citation-record","integrity":"/paper/2608.10634/integrity","json":"/paper/2608.10634/citation-record.json","paper":"/paper/2608.10634"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:21.278290Z","title":null,"venue":null,"work_id":"d2351caf-7639-42da-bb05-71fe7e5b5f4e","year":1998},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.178335Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:f2d2b3c26e17aaf1b388bf764f7ce43d9f475b103b7899e291785808ad6a4644","observation_id":"28b27ca2-c7ee-4184-83a1-6292f33be2ce","resolution":{"observed_at":"2026-08-12T20:36:21.284401Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:21.258962Z","title":"Integrated architectures for learning, planning, and reacting based on approximating dynamic programming,","venue":null,"work_id":"f6e21c6e-186b-42d6-8992-377a658892ec","year":1990},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.184095Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:43865c7d7fab58f69eb24d0773dc73089b8ec4333fdc1b2027d2e058e0f0a64c","observation_id":"f02e2f1a-cf4c-44f8-b1a4-addaef8ec3be","resolution":{"observed_at":"2026-08-12T20:36:21.265235Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.191188Z","title":"A survey on model-based reinforcement learning,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.191188Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:dc6156f74a924d89650f264353d1218d938196a3a3f6cb1ac426fae0ee0edb4c","observation_id":"45a925a3-33fc-4e69-9eb4-b160bd8f136a","resolution":{"observed_at":"2026-08-12T20:36:20.191188Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:21.226752Z","title":"Pilco: A model-based and data-efficient approach to policy search,","venue":null,"work_id":"0b30ba4f-f5ea-460d-88f4-d2327cb9c098","year":2011},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.198063Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:eda732ffd2d19f5007b6028f71a7136e5f557f68472bf22d2bdabff4fbebd4cf","observation_id":"58e7f571-3543-461a-9baf-56f751164edc","resolution":{"observed_at":"2026-08-12T20:36:21.232670Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:21.208003Z","title":"Guided policy search,","venue":null,"work_id":"d6cb4d48-9278-4bea-a428-1fffc5d3841b","year":2013},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.204289Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:71444532cd970ef857dce939237f82954f45a571121289c2eb9954390045a98a","observation_id":"7f34b775-ac32-4beb-ae72-ca2256f3d397","resolution":{"observed_at":"2026-08-12T20:36:21.213586Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.210676Z","title":"Dream to control: Learning behaviors by latent imagination,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.210676Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:ea1631b7d423ab87d7004d5e1b5c306b29f193f1914e73ab1221ed599c3c8a62","observation_id":"995b08b3-c80c-41cf-af88-210166c92694","resolution":{"observed_at":"2026-08-12T20:36:20.210676Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:21.177100Z","title":"Model- based reinforcement learning: A survey,","venue":null,"work_id":"7625fab4-9aea-4664-a565-5fd4d5d9d17d","year":2023},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.217167Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:7dd4dd7a1c788bf52286cbddebc3ba522d4f778e41ded8123d233ce2ff0c114a","observation_id":"e75ac4ca-c976-4cdf-9b96-beb1d4f35868","resolution":{"observed_at":"2026-08-12T20:36:21.183384Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.223410Z","title":"Mastering atari, go, chess and shogi by planning with a learned model,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.223410Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:5c78323e2a4f1790d9a99a9d8308a69e6cce7cef0abdbbdd499936641ed0d196","observation_id":"c896dfa0-ba87-44d1-b083-777145438253","resolution":{"observed_at":"2026-08-12T20:36:20.223410Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.229013Z","title":"When to trust your model: Model-based policy optimization,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.229013Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:ac8ab53c5e62dd7e50e059e25c711368f5037d3497d2eb552c86df887b267bb3","observation_id":"1cb31a0f-9544-4027-83ae-23b71454b247","resolution":{"observed_at":"2026-08-12T20:36:20.229013Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.234867Z","title":"Dyna, an integrated architecture for learning, planning, and reacting,","venue":null,"work_id":null,"year":1991},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.234867Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:6b58d288a216f6b3e3ca6f9252cf2d5d070e00ac4f2f3b9f61ba798d180cfd44","observation_id":"504de25b-14ba-4c70-8b52-e3f9caa80379","resolution":{"observed_at":"2026-08-12T20:36:20.234867Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.240120Z","title":"Soft actor-critic: Off- policy maximum entropy deep reinforcement learning with a stochastic actor,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.240120Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:e87b47508d6bfc7c34bc9c1fdb6d5f8908d1f80c4ce017882c8553c4f1a2101a","observation_id":"61dce01c-3da3-498c-a710-d932218e9254","resolution":{"observed_at":"2026-08-12T20:36:20.240120Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:21.109761Z","title":"Model- ensemble trust-region policy optimization,","venue":null,"work_id":"d118e5b3-cac5-47da-8dcc-15ec6cca3ed4","year":2018},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.246483Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:7f9c243d7a3d378b41caa572487c288df27db349ef61f05cee47e8afefdb2bf9","observation_id":"b1ada517-e2f2-4c66-86a7-cbdf16c094a1","resolution":{"observed_at":"2026-08-12T20:36:21.115480Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:21.092888Z","title":"Improving multi-step prediction of learned time series models,","venue":null,"work_id":"46648aa6-4ffd-4535-99ba-64aa24aa8b77","year":2015},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.251838Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:50e92c403828e91009388e89a34eef290587ead8344bffb80267a8432bbe999d","observation_id":"a9f9fddf-8cce-4702-b00e-a2e106782360","resolution":{"observed_at":"2026-08-12T20:36:21.098144Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:21.075969Z","title":"Lipschitz continuity in model- based reinforcement learning,","venue":null,"work_id":"58c211c8-14ab-42d1-83d9-fc0d2965be0a","year":2018},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.256813Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:f5c4e3ef9e7c0a809cdc7a31b4e1f5c69d48a8975b45c2d2d4b019196e8b58a9","observation_id":"73b0b102-65bc-4404-8cc3-7ad5440de09b","resolution":{"observed_at":"2026-08-12T20:36:21.081203Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:21.058490Z","title":"Adversarial counterfactual environment model learning,","venue":null,"work_id":"47267ec7-fc40-43d7-b8d2-8714feffa53d","year":2023},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.261848Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:7f61d5da0cab35dd86ad8cb39fe1284f8e16c336247cf855736bff6653aa7b52","observation_id":"dc98699a-6ec9-4557-9e59-8a6ba853c707","resolution":{"observed_at":"2026-08-12T20:36:21.063766Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.267264Z","title":"Generative adversarial imitation learning,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.267264Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:553be9b594820ff2ff045b71681f82f79bd1f1227dfbf70b912c579352fc4e1e","observation_id":"22692b15-6104-4750-a73e-c7be37d9e644","resolution":{"observed_at":"2026-08-12T20:36:20.267264Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.272388Z","title":"Off-policy deep reinforcement learning without exploration,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.272388Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:40b05d278315a2cb0abaae15c9085eb4f9d1bd08221a54b1118e5ea430415795","observation_id":"97c20a73-58a5-4de2-b387-a296aca14ff0","resolution":{"observed_at":"2026-08-12T20:36:20.272388Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.277428Z","title":"Conservative q-learning for offline reinforcement learning,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.277428Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:b7287f81ae5331992408a6f1fd62cc4662ee845288bc8bc0ecfa37c2d1f396dc","observation_id":"e693515d-4dda-4898-bf30-8221e08ae0ac","resolution":{"observed_at":"2026-08-12T20:36:20.277428Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:21.004872Z","title":"Trust the model where it trusts itself - model-based actor-critic with uncertainty-aware rollout adaption,","venue":null,"work_id":"6cf6820f-a1ec-422f-8970-af435bfd3233","year":2024},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.283264Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:66b8ddd5255f3aa85dc7bf0288c44417503907712eda4a77344aea12ab886f7d","observation_id":"314fda7a-e4ba-4643-aac9-b0173e31dfa1","resolution":{"observed_at":"2026-08-12T20:36:21.010586Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.986826Z","title":null,"venue":null,"work_id":"9e813d93-d407-48cb-b423-08eedba047df","year":2024},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.288600Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:86b35150dc72b1cfe0d43213def54a6ea08be910de761d7c112185d7417f2d3b","observation_id":"14afb53a-9d66-40e3-a29a-eb81221c133a","resolution":{"observed_at":"2026-08-12T20:36:20.992278Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.966815Z","title":null,"venue":null,"work_id":"3a1fe775-026e-4b26-a5ac-c0e91a5da8d5","year":2011},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.294600Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:6270bd030c602a91a1752e411da9ecd11e75e8e1f381caca0a6fea65d8f4354a","observation_id":"3c790d0c-c3f7-4ad6-afb5-d207ba74c0f7","resolution":{"observed_at":"2026-08-12T20:36:20.973562Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.948227Z","title":"Continuous deep q- learning with model-based acceleration,","venue":null,"work_id":"cb1cb030-b52a-4bd9-9e25-a7cd10ac1512","year":2016},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.301596Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:5a0329fc4013404eb13388d65840f826d7ecbbdcc7436df0c4072ac5997836a8","observation_id":"e95ccb00-28bd-4029-9523-8a1b3dd2eb89","resolution":{"observed_at":"2026-08-12T20:36:20.954421Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.308186Z","title":"Neural network dynamics for model-based deep reinforcement learning with model-free fine-tuning,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.308186Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:b633c1c310544fc44b695bc9b52ff3b585d902b2c69c53ad256aeeac31c36db3","observation_id":"c59a48b8-f357-49e2-99f9-07049c16dea1","resolution":{"observed_at":"2026-08-12T20:36:20.308186Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.917507Z","title":"Improving pilco with bayesian neural network dynamics models,","venue":null,"work_id":"86c54c3b-4e4f-4781-a2fa-5d3311328004","year":2016},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.313769Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:762dfd33c45d88f8722ee4ba81961147c3aa2274977d918d6ae01d464e700a26","observation_id":"4bb7dbe4-2ed9-4c24-8d2f-f9d310e9897f","resolution":{"observed_at":"2026-08-12T20:36:20.923589Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.899294Z","title":"Sample- efficient reinforcement learning with stochastic ensemble value expan- sion,","venue":null,"work_id":"1da3b9e3-a133-45b2-9876-7b48fe858188","year":2018},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.320167Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:b3859f887c5a62c2a3cb5bba2c9cba2022f1cbaf07d84f7f232c581032285713","observation_id":"7eb79ea2-afe9-4084-b0e6-db74fac64df7","resolution":{"observed_at":"2026-08-12T20:36:20.905022Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.325577Z","title":"Deep reinforce- ment learning in a handful of trials using probabilistic dynamics models,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.325577Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:f3d61df22d9a6562d415435b70441bbe99cc63295e6783cc751d7f320f0cff26","observation_id":"94d0768a-f953-4617-aa0e-781c14e8bf8c","resolution":{"observed_at":"2026-08-12T20:36:20.325577Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2005.01643","last_updated":"2020-11-01T23:50:25Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-05-04T17:00:15Z","title":"Offline Reinforcement Learning: Tutorial, Review, and Perspectives on Open Problems","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2005.01643","snapshot_observed_at":"2026-08-12T20:36:20.331828Z","title":"Offline reinforcement learning: Tutorial, review, and perspectives on open problems,","venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.331828Z"},"links":{"cited_paper":"/paper/2005.01643","citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:26d28e899a855ea47333319871e7f9e37163473e9cd07baf1e67efd18aa607b2","observation_id":"6fccfbbf-34c9-44cf-94fe-1276027f7efa","resolution":{"observed_at":"2026-08-12T20:36:20.331828Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.868343Z","title":"Uncertainty-based offline reinforcement learning with diversified q-ensemble,","venue":null,"work_id":"84da7c44-c356-4112-b2c1-5c261a7f840c","year":2021},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.337551Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:5f7267178d9772d0feb037db0079672f7fdf1f8c685c4e8534ce547d8162e442","observation_id":"ecf56e25-1b4e-4f77-bd5a-d4481f4a5756","resolution":{"observed_at":"2026-08-12T20:36:20.875012Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.850248Z","title":"Adversarially trained actor critic for offline reinforcement learning,","venue":null,"work_id":"3f4ac90c-18b5-41bc-832f-84227a19a04d","year":2022},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.342694Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:07ce09a4aa893c2320d83fe61b1027c92599e28030d0e3c4e2342b875a3bca6a","observation_id":"0f3ad808-4563-454e-ac38-4d5f8161f3c7","resolution":{"observed_at":"2026-08-12T20:36:20.855922Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.347942Z","title":"Mopo: Model-based offline policy optimization,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.347942Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:69948034cfe080a4207fe6acd05f0d9e67a6f2bf2b7647fad456301f795d47fd","observation_id":"1d78be10-6b49-4a70-b76b-eddf036d538c","resolution":{"observed_at":"2026-08-12T20:36:20.347942Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.353342Z","title":"Morel: Model-based offline reinforcement learning,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.353342Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:5ceed3d9f4d1cdb8e19f595f6bf6dac149d0228661bf097c525c4a3474708243","observation_id":"27a348e7-2fe2-4a81-b5dc-366b7c4c7b68","resolution":{"observed_at":"2026-08-12T20:36:20.353342Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.808180Z","title":"A review of off-policy evaluation in reinforcement learning,","venue":null,"work_id":"eabf8cca-3cdd-4da9-9a4d-f5dc2a906e78","year":2025},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.358661Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:ed4a68258f27082b008a3761135c9751fecc42848f9f17765f45d33da33172dc","observation_id":"5e7970f8-90ce-4342-a844-1e31d6093bab","resolution":{"observed_at":"2026-08-12T20:36:20.814039Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.790999Z","title":"Marginal mean models for dynamic regimes,","venue":null,"work_id":"847bb845-000f-498a-a1d2-340d8d85dc2d","year":2001},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.364511Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:4b54b066708e575f51a2397fc6ebf56424ec0313501b3b1a132dbc86e92cc84b","observation_id":"5133181f-e575-49db-866f-1bf5c76af332","resolution":{"observed_at":"2026-08-12T20:36:20.796441Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.773425Z","title":"Doubly robust policy evaluation and learning,","venue":null,"work_id":"a8e9f784-be62-44a0-8b07-3f41ff3ee47a","year":2011},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.370281Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:1cf2ac74cc71e9efa74b683572c00666197c64143b3668b3ae590cbb5ecd09c2","observation_id":"36ecd89f-e326-4b22-8755-ff5bae80124c","resolution":{"observed_at":"2026-08-12T20:36:20.778990Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.755950Z","title":"Estimation of the causal effect of a time-varying exposure on the marginal mean of a repeated bi- nary outcome,","venue":null,"work_id":"421e73a8-11b1-4200-abfe-42ff337d303e","year":1999},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.376012Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:8e9b934550ee38eefcfc09e46d6f63d984687d7a9fc8ebd3862a765fb20e307e","observation_id":"9a8dd063-e27a-41ba-b611-d4efe158a887","resolution":{"observed_at":"2026-08-12T20:36:20.761862Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.739187Z","title":"Eligibility traces for off-policy policy evaluation,","venue":null,"work_id":"e9b22297-b509-4e7b-b653-eff6d35f7ee9","year":2000},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.381373Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:c9f956c3e9323770ac85a812ef8920f264c4fd01384da34d62bc14673272405b","observation_id":"22d8fadb-4df1-4764-bec9-ca285aca10bb","resolution":{"observed_at":"2026-08-12T20:36:20.744467Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.386672Z","title":"The central role of the propensity score in observational studies for causal effects,","venue":null,"work_id":null,"year":1983},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.386672Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:18c62e8eb0a8197fdf0700a3704d59748f8543473134db89fc619db7cd1367f2","observation_id":"89e7a820-645e-4e3a-ae7e-97ac8674db67","resolution":{"observed_at":"2026-08-12T20:36:20.386672Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.708488Z","title":"More robust doubly robust off-policy evaluation,","venue":null,"work_id":"d92d827c-3b4c-4c8f-9d3c-0587a8ba42d0","year":2018},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.391792Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:f26135a8ebf714b36e85b9d5bc51a1b3441d416872d49d51108e657804ebf595","observation_id":"87d1898e-f1be-4eac-94b9-215836f69379","resolution":{"observed_at":"2026-08-12T20:36:20.713942Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.690566Z","title":"Estimation of regression coefficients when some regressors are not always observed,","venue":null,"work_id":"3d6cee23-4ea1-4ba2-980f-3b3454c4ae4b","year":1994},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.396837Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:45cd9cf66a5ab0d661e522caa97222cbad4f381985a91002f88ab37fa28c83fb","observation_id":"1f4f5dbf-f02b-49ab-845a-0d4c9e93738a","resolution":{"observed_at":"2026-08-12T20:36:20.696146Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.670730Z","title":"Doubly robust off-policy value evaluation for rein- forcement learning,","venue":null,"work_id":"a88c01a4-20dd-463b-837b-07aafef55179","year":2016},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.402188Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:a3d8b62b572c81062bca7844976c5f6c30198facb40185adb0c8145cc774ca1d","observation_id":"ec4179e0-81b1-4508-af59-1e4256f2068d","resolution":{"observed_at":"2026-08-12T20:36:20.677203Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.651808Z","title":"Data-efficient off-policy policy evaluation for reinforcement learning,","venue":null,"work_id":"77b624cf-646c-44df-ab9e-0338d8f867cf","year":2016},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.407069Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:f93eda1ddbbc00eeb90b5a95289d94fbc3ab6d76acade09a699b53d66e0cbe23","observation_id":"adfa28ff-b8ec-4853-b12e-1f6b7155db90","resolution":{"observed_at":"2026-08-12T20:36:20.657617Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.632795Z","title":"More efficient off-policy evaluation through regularized targeted learning,","venue":null,"work_id":"e5b42d55-4665-4dde-885a-bcb347cead11","year":2021},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.412056Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:5a5acb3f6e730fa5980d93503eb564af104d49044fa25381acce700934d218bb","observation_id":"2edb2bd2-2073-4406-9598-1b8861cbb514","resolution":{"observed_at":"2026-08-12T20:36:20.638625Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.612126Z","title":"Targeted maximum likelihood learning,","venue":null,"work_id":"4e86e2d6-689b-4761-8864-436be91a4200","year":2006},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.416898Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:e6cbd81efc282fe464efeed472adc680192c865709a2c3150deb44a3a47fee95","observation_id":"283a0426-3aa0-47bd-8719-76c3367c1a39","resolution":{"observed_at":"2026-08-12T20:36:20.619129Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.592163Z","title":"Policy gradi- ent methods for reinforcement learning with function approximation,","venue":null,"work_id":"d36a1792-5e3c-44d5-9ac1-3727dc24ca98","year":1999},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.421984Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:b0613e6fa45516cc1a724000298a314a2e187e00cc930bee751d027d4bf2b336","observation_id":"31e55c0a-d63c-4407-9ee6-4597563f95c0","resolution":{"observed_at":"2026-08-12T20:36:20.598157Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.572044Z","title":"Nonlinear ica using auxiliary variables and generalized contrastive learning,","venue":null,"work_id":"298210a1-5521-4a48-a002-ac575bc2fed6","year":2019},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.427101Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:8eaf57d0ae1891b459757a3b7f773b93d11c34bef449da32b8911ce1f2594f6f","observation_id":"1e8fa356-5a1b-4b89-81a7-cd9078f35a96","resolution":{"observed_at":"2026-08-12T20:36:20.578764Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.432128Z","title":"Mujoco: A physics engine for model- based control,","venue":null,"work_id":null,"year":2012},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.432128Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:dab56ce3511fc5125d0f77c065eadafaa41158292d90fbfba58260488116e3ed","observation_id":"45284cf3-0553-4120-9fdf-7823d1c72afd","resolution":{"observed_at":"2026-08-12T20:36:20.432128Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06347","last_updated":"2017-08-28T09:20:06Z","snapshot_observed_at":"2026-08-15T20:26:32.102285Z","submitted_at":"2017-07-20T02:32:33Z","title":"Proximal Policy Optimization Algorithms","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.06347","snapshot_observed_at":"2026-08-12T20:36:20.437034Z","title":"Prox- imal policy optimization algorithms,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.437034Z"},"links":{"cited_paper":"/paper/1707.06347","citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:0efa56bd3f4659616c5ca45eb2dfd0a86c935c2d885f2698904556e9ef217de2","observation_id":"231f4d62-5cee-4d60-b890-520334e26baf","resolution":{"observed_at":"2026-08-12T20:36:20.437034Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T20:36:20.539872Z","title":"Algorithmic framework for model-based deep reinforcement learning with theoretical guarantees,","venue":null,"work_id":"a3de64a2-c3d2-43cd-b100-75d904bbb1b7","year":2019},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.442539Z"},"links":{"citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:63b1affb29558585b3e1ade8b55a9785e7999d00bd7ba948c97b16fe519e643f","observation_id":"91d19a9d-e8ec-49b8-af2b-ef3abefebaf2","resolution":{"observed_at":"2026-08-12T20:36:20.547305Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2103.07861","last_updated":"2021-03-14T07:37:28Z","snapshot_observed_at":"2026-07-06T10:49:42.752119Z","submitted_at":"2021-03-14T07:37:28Z","title":"VCNet and Functional Targeted Regularization For Learning Causal Effects of Continuous Treatments","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2103.07861","snapshot_observed_at":"2026-08-12T20:36:20.447781Z","title":"Vcnet and functional targeted regularization for learning causal effects of continuous treatments,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-12T20:36:20.447781Z"},"links":{"cited_paper":"/paper/2103.07861","citing_paper":"/paper/2608.10634"},"observation_digest":"sha256:124357a8a0259585cac2e364ab4f2acb20f5be4c8e56c7c2a394debf8cd9b383","observation_id":"b8a0b7bc-1b84-4683-97c3-a69109487b55","resolution":{"observed_at":"2026-08-12T20:36:20.447781Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2608.10634","last_updated":"2026-08-11T08:20:02Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-14T23:10:59.817898Z","submitted_at":"2026-08-11T08:20:02Z","title":"IADD-TR: Intervention-Aware Dynamics Decoupling with Targeted Regularization for Model-Based Reinforcement Learning"},"reference_resolution":{"displayed":49,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":21,"verified_exact":0,"verified_fuzzy":28},"total_outbound_references":49},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"thesis":"As of 16 August 2026, this Paper Citation Record lists 49 of 49 outbound references and 0 inbound Pith citation observations for arXiv:2608.10634."}