{"as_of":"2026-08-09T16:30:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:e8a2001efe985f3012542e3c44826788477d6c95495c0e84dc2a79c7f4c55460","coverage":[{"denominator":34,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":34,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T01:10:58.285736Z","state":"measured"},{"denominator":34,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":34,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2506.11912/citation-record","integrity":"/paper/2506.11912/integrity","json":"/paper/2506.11912/citation-record.json","paper":"/paper/2506.11912"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.510215Z","title":null,"venue":null,"work_id":"177ae22e-810a-4e11-93f9-8ddaa1b467b5","year":1994},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:55.203368Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:1f5df8674196f8aed99bad120b9ba6675a23da55315dc3d56bc5a7b7aab2cce6","observation_id":"46dbb151-be58-4093-b9e6-3df113f9e72b","resolution":{"observed_at":"2026-08-07T01:11:01.513026Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.500967Z","title":null,"venue":null,"work_id":"2ccbdac1-9be0-41b6-9646-0d15f1f974b2","year":2001},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:55.291277Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:5270e0b28eac0e0a8aee2bd3bafceb3f6e9ca1537f57059c9ea8fe74c799b98f","observation_id":"7fa74118-4a6f-4053-85ca-dbe902c715d9","resolution":{"observed_at":"2026-08-07T01:11:01.503874Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.491129Z","title":null,"venue":null,"work_id":"9ad7ea5a-c0b4-44ed-bb94-8b0e543a0fe3","year":1999},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:55.396925Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:dd9030ebcac8d810ca0f2cf467753906b64c337fc6bcd1ac5e7ec71d14acb0b0","observation_id":"a48bfeab-5b40-44cf-afa0-0288a6cf3763","resolution":{"observed_at":"2026-08-07T01:11:01.494683Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.481542Z","title":"C., and Le Roux, N","venue":null,"work_id":"9f6e1ed4-bb1d-4f78-bd4a-ad537201d97f","year":2021},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:55.440340Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:ab3c0df860c2c5f75aa0080b7674722bccb17c7908b60bc3ca9409d9e44e963e","observation_id":"472a5fd2-fc47-4978-b20a-eb5ec2ce3451","resolution":{"observed_at":"2026-08-07T01:11:01.484545Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.472409Z","title":"and Vicente, R","venue":null,"work_id":"5f2eb479-dea6-4774-be79-d59ecf55353b","year":2022},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:55.553966Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:3465c5d5db1ad0fbc162bb8c8a3a4ba4af7b9173e15b80a3833fc9b3afd8f303","observation_id":"774533e8-5f8b-42e7-8098-6098615d70c0","resolution":{"observed_at":"2026-08-07T01:11:01.475565Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.462661Z","title":null,"venue":null,"work_id":"251ea4fa-dbb9-4a52-bb7d-916abda65d5a","year":2019},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:55.654464Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:027561ec4e21909564d448aad39f531938cb31f3735bd7b7b4d45dd922f4c546","observation_id":"17b54c4f-3a44-4c3d-9591-9fb04aa295dd","resolution":{"observed_at":"2026-08-07T01:11:01.465631Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.452675Z","title":null,"venue":null,"work_id":"e1327bd4-a4a9-45bf-81db-4c1fc5df5880","year":2023},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:55.748944Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:c19ebd3a9d3aea52587e56aa2bf74283bebef49ca5ad7092212ce3f389a45488","observation_id":"008efe3d-cae2-4101-bb9c-a4cd6fb11ebf","resolution":{"observed_at":"2026-08-07T01:11:01.455751Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.442660Z","title":"G., and Pineau, J","venue":null,"work_id":"4e1ccf9d-fc8f-4978-9d9e-6518b321fb70","year":2018},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:55.860801Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:2ea1a6d0d40da0adf5188b617f97fcde3d9eb2f389bb414a66ffae076a18f857","observation_id":"0d2581ef-e31a-46d5-97d7-76db9bcee0f1","resolution":{"observed_at":"2026-08-07T01:11:01.446104Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.433140Z","title":null,"venue":null,"work_id":"2e82aa65-2851-464d-be4b-26d03ce7241c","year":2001},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:55.947805Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:120e3dc23326eeff1c8997e1e9c9410d9d54204d28a77b052e3ba0a2525c7806","observation_id":"72333d80-22e4-48dc-be10-5598b5e2e5f2","resolution":{"observed_at":"2026-08-07T01:11:01.436266Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.423770Z","title":"M., de Vries, J","venue":null,"work_id":"452186ae-b64c-4e9f-834b-86b1b0066d9d","year":2024},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.026758Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:75dd4b5be7ec6e108e257f78fca072d1f93646e77a8ab0b573934e5fede1515f","observation_id":"0cffefbb-941c-42ae-b72d-98555e55efcd","resolution":{"observed_at":"2026-08-07T01:11:01.426839Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.413392Z","title":null,"venue":null,"work_id":"2f6bab60-badb-43a3-872b-3c0d1acf37d4","year":2017},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.120102Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:fd8bfd24686a0bf3bf01534e74b983afdff2322534fa68388f778317801809f7","observation_id":"0cb95045-92ff-4689-ad7b-f637f7d868e0","resolution":{"observed_at":"2026-08-07T01:11:01.416975Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.402628Z","title":null,"venue":null,"work_id":"350c5ab0-46d3-4713-8678-3af98f49f876","year":2023},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.240654Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:4686f046c4aca873ccbb576bc52c2005da71267d97b90bd2b9c855976e4b92f9","observation_id":"79f77830-87ff-4e15-a2d8-af1c07d25c23","resolution":{"observed_at":"2026-08-07T01:11:01.406013Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.392632Z","title":null,"venue":null,"work_id":"ae94af07-062b-4e2d-b348-733ad6584158","year":2023},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.297251Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:30cb3d3f55b28057bcb64ed227ca1db25148312741e5eac6905ebeb61a1e9449","observation_id":"1e060f05-7894-4cfe-9ce1-b6b35a67ce2b","resolution":{"observed_at":"2026-08-07T01:11:01.396050Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.381649Z","title":"C., Bellemare, M","venue":null,"work_id":"081fda7c-be4f-47b2-93e5-5d8671913862","year":2018},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.401160Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:adaf2f92d7c49f0efda14b31ad7f598ede37e90ac1a214361de76ad98c1bc364","observation_id":"4f4a59dc-a71d-4fe4-99b9-c11ff7f435a2","resolution":{"observed_at":"2026-08-07T01:11:01.385105Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.371622Z","title":null,"venue":null,"work_id":"6f43bdd9-d83c-4ec9-a52a-7b7a6ea1aa12","year":2017},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.520833Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:b5d76c5f64e14f0205a28d7abc61f43a6b9765a96021e0d67edbd463dd5dc336","observation_id":"48fc7080-b5d6-4796-a17a-bb22c0c8f1d5","resolution":{"observed_at":"2026-08-07T01:11:01.374703Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.359605Z","title":"and Tsitsiklis, J","venue":null,"work_id":"61481729-a676-478c-bda6-b6968c884488","year":1999},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.607417Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:997fd59717cf3de44be2e9433857260749a8e618775b9e6f8f2ba58d7734f0b2","observation_id":"0b379814-440b-41ee-880b-86b21970ceff","resolution":{"observed_at":"2026-08-07T01:11:01.363244Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.083368Z","title":null,"venue":null,"work_id":"d8c5ff14-ee94-40ea-99ec-aaec058fc546","year":1995},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.649485Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:80d4fb201c3a0fc7105afc1dcf9327e15460ca438f82c2cf2f61fc63c119b059","observation_id":"039bf625-d68f-45b1-9c64-7387d9bc8109","resolution":{"observed_at":"2026-08-07T01:11:01.230910Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:00.800353Z","title":null,"venue":null,"work_id":"f25904bc-ad48-4b42-bc14-e94ff4c56419","year":2022},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.755411Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:952e7292a6c94ed1682ee802246a4c0ba4ac6c5fe09335ad43db0498782920eb","observation_id":"626055d5-cb09-4b3d-8de2-b6c8b386213b","resolution":{"observed_at":"2026-08-07T01:11:00.918276Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:00.550564Z","title":"u rtler, N., Neitz, A., and Sch \\","venue":null,"work_id":"17f29777-3351-42f4-ad5e-eec18dfcc6e1","year":2022},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.868913Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:1a331df686363e818ee96f89e97eb9b08d27f9ba1d3db2b991d7026bb189a030","observation_id":"d49e9354-15ef-4124-9211-41b03af524b9","resolution":{"observed_at":"2026-08-07T01:11:00.669845Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:00.324269Z","title":"and Sch \\\"o lkopf, B","venue":null,"work_id":"9a80fad2-78ac-4e48-89da-4c8918284acb","year":2024},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.938340Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:7bb64e369fc30fc8c2e0cb2f5b91708dbbd7966f7d19b96a31bafc787f4c1010","observation_id":"e56ff290-4238-4c07-af89-ff79ed58c5b3","resolution":{"observed_at":"2026-08-07T01:11:00.441155Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:00.141659Z","title":null,"venue":null,"work_id":"abbecb45-715c-46cf-8b08-645200974019","year":2016},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.021373Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:649030a978a3a66c6453018d51371d776e4ec3b0891e26c32f655f882242d5f9","observation_id":"5f39d1e9-b2f7-4a8e-be39-4eba7c95d96b","resolution":{"observed_at":"2026-08-07T01:11:00.240560Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:57.146209Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.146209Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:24bb4c35466a7beab7be1b634c3b2e399b2791b8ef790c60a1eabbeb5b837599","observation_id":"8aa180b3-4a14-459b-9713-3c19c869c926","resolution":{"observed_at":"2026-08-07T01:10:57.146209Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:59.918876Z","title":"and Fergus, R","venue":null,"work_id":"c4729972-4f65-41a0-a677-4c8d9599bb9c","year":2021},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.288542Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:0bc97a95f1e55e446ad0243a985b787913a944d4526e709d9eff1e679b7f6b63","observation_id":"41569b00-485d-4119-b71a-356d554fc190","resolution":{"observed_at":"2026-08-07T01:11:00.030784Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06347","last_updated":"2017-08-28T09:20:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2017-07-20T02:32:33Z","title":"Proximal Policy Optimization Algorithms","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.06347","snapshot_observed_at":"2026-08-07T01:10:57.370198Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.370198Z"},"links":{"cited_paper":"/paper/1707.06347","citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:91f371f76993989c424b0858f661c695e65fed5e2a90daa34f1c1011c79a7a5e","observation_id":"7591ef08-5a22-4c0a-adc0-723e66f4365e","resolution":{"observed_at":"2026-08-07T01:10:57.370198Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:59.741513Z","title":null,"venue":null,"work_id":"e28c9ae2-2ce4-452b-aa73-7f45011e58c7","year":2020},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.503345Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:963ffa44581ae0b080088b1441704667bd16b1550a5da90e07d36c9774250a2f","observation_id":"97447029-0d80-4da3-8459-ab750a3cc7ba","resolution":{"observed_at":"2026-08-07T01:10:59.838039Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:59.522454Z","title":null,"venue":null,"work_id":"ef6dd7d2-6ad5-481d-91eb-5adad625b356","year":2024},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.574995Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:ecb918f59e4506482056387b5aff31c1933e7b71640e1cdcad437b883716ee11","observation_id":"87f7db3e-fe1f-44e8-989c-25d07d12a3e9","resolution":{"observed_at":"2026-08-07T01:10:59.624083Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:57.647845Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.647845Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:5b953074935ca7214bf0d549155972b539290740b120ee585c79b9036fade2d1","observation_id":"396052ca-a9e3-49f0-b9cf-d44422daf8f3","resolution":{"observed_at":"2026-08-07T01:10:57.647845Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:59.230776Z","title":"S., McAllester, D., Singh, S., and Mansour, Y","venue":null,"work_id":"974821b0-2394-476b-87f2-2b55b8fa4752","year":1999},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.732333Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:8ad63110c80fe67e2205a8922a0a1eb9a5d82cdc1cb7c2e5b0a0abf8351c2915","observation_id":"bb8be8b2-ca3f-430f-a7c7-428885eb974d","resolution":{"observed_at":"2026-08-07T01:10:59.373570Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:59.016376Z","title":"and Parr, R","venue":null,"work_id":"7e93fab5-09a4-4852-9815-52d7d835f736","year":2009},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.787899Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:00bf8aab8a34a761d4314b44553b23f52bb0669957d2b3ba42d81e608e4fcdcf","observation_id":"0f2b8c9d-558a-43f8-9073-42b7b6129aa6","resolution":{"observed_at":"2026-08-07T01:10:59.149889Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2110.06539","last_updated":"2021-10-13T07:31:31Z","snapshot_observed_at":"2026-07-06T11:57:21.548494Z","submitted_at":"2021-10-13T07:31:31Z","title":"On Covariate Shift of Latent Confounders in Imitation and Reinforcement Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2110.06539","snapshot_observed_at":"2026-08-07T01:10:57.900276Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.900276Z"},"links":{"cited_paper":"/paper/2110.06539","citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:a8e8c8b3b0578392541380545d0e9fcc95d0a9faa43bcbc63ef9398dd252d64a","observation_id":"91a90401-e57a-4cc7-be2d-1249da362344","resolution":{"observed_at":"2026-08-07T01:10:57.900276Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.03565","last_updated":"2026-07-13T10:54:37Z","snapshot_observed_at":"2026-07-16T23:19:23.659963Z","submitted_at":"2024-10-04T16:15:31Z","title":"Training on Irrelevant States Implies Data Augmentation: Generalization in Contextual MDPs","version":4},"cited_work":{"arxiv_id":"2410.03565","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.03565","snapshot_observed_at":"2026-08-07T01:10:58.387718Z","title":"Training on Irrelevant States Implies Data Augmentation: Generalization in Contextual MDPs","venue":"cs.LG","work_id":"b5b5094e-6af4-45b2-ba1e-a03d1b7d4a6c","year":2024},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:58.002804Z"},"links":{"cited_paper":"/paper/2410.03565","citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:bdd8783f5b6a5a6bb9780b5d3e3ad668e5456aba3bed0cdfa2acd22774f87fda","observation_id":"a0f92a0c-aec2-4371-a52b-d10353729270","resolution":{"observed_at":"2026-08-07T01:10:58.465709Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:58.106871Z","title":null,"venue":null,"work_id":null,"year":1992},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:58.106871Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:9dd0525c1d8ddc0b1f530f1547aef0b153cc7331a421d6df2846d2a1bcfa341c","observation_id":"91cc7418-1aae-4ac9-addb-042e5ee8f2ce","resolution":{"observed_at":"2026-08-07T01:10:58.106871Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:58.775263Z","title":null,"venue":null,"work_id":"9e004dea-b842-49bb-acbb-458168caa45d","year":2020},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:58.218410Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:d6c415db30f509543662e68403ae0cb27101fee0231f60828b37fa0c7103aab0","observation_id":"d9267978-3bed-40d2-b067-32ec803fd44e","resolution":{"observed_at":"2026-08-07T01:10:58.896922Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:58.597029Z","title":null,"venue":null,"work_id":"dec52605-b577-4a00-bb45-1caf2c513eae","year":2020},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:58.285736Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:bd04dc714636649f77d8f382942867fd2098a1abfe3274908dc0989942905524","observation_id":"1367568e-2ac0-45e7-9829-ca83dbfddbc0","resolution":{"observed_at":"2026-08-07T01:10:58.684426Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-07T00:58:25.532315Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations"},"reference_resolution":{"displayed":34,"state_counts":{"malformed_identifier":0,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":22,"verified_exact":0,"verified_fuzzy":11},"total_outbound_references":34},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 34 of 34 outbound references and 0 inbound Pith citation observations for arXiv:2506.11912."}