{"as_of":"2026-08-08T04:42:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:5d9e8fe2d1635e583107d060de61949da4410169138054ca4c8542a2c94302e9","coverage":[{"denominator":95,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":95,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T22:12:50.278134Z","state":"measured"},{"denominator":96,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":96,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-27T01:13:11.483599Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-03T20:38:55.959708Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"cited_work":{"arxiv_id":"2502.09297","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.09297","snapshot_observed_at":"2026-07-03T20:38:55.959708Z","title":"arXiv preprint arXiv:2502.09297 , year=","venue":null,"work_id":"2530813b-fade-4c33-b5da-fa4631363e51","year":null},"citing_paper":{"arxiv_id":"2606.18089","last_updated":"2026-07-05T17:40:26Z","snapshot_observed_at":"2026-07-12T13:34:39.011240Z","submitted_at":"2026-06-16T15:55:28Z","title":"From Reasoning Traces to Reusable Modules: Understanding Compositional Generalization in Language Model Reasoning","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-06-27T01:13:11.483599Z"},"links":{"cited_paper":"/paper/2502.09297","citing_paper":"/paper/2606.18089"},"observation_digest":"sha256:7faefacc245816e479c166ba8bb4f1522666db202a9585721cc241dd56290ac5","observation_id":"4da73e55-874a-4baf-8256-557b114f231d","resolution":{"observed_at":"2026-07-03T20:38:55.961175Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2502.09297/citation-record","integrity":"/paper/2502.09297/integrity","json":"/paper/2502.09297/citation-record.json","paper":"/paper/2502.09297"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:49.898591Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.898591Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:7db98583a1a44b2e6ec94474b238f68223f8bfa2a90a0f20d1cc5b4496ef65cf","observation_id":"b3c4a923-7ae7-4449-b595-054fb0bf1cf6","resolution":{"observed_at":"2026-08-07T22:12:49.898591Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:49.904503Z","title":"Generalization on the unseen, logic reasoning and degree curriculum","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.904503Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:27dd0a8fa5b05ea6ed12962a8d6252325264dbbaf4c5c0073b1f95d85867ff38","observation_id":"324013d2-b817-4747-ae34-1047fb26e0ca","resolution":{"observed_at":"2026-08-07T22:12:49.904503Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:49.908933Z","title":"Interventional causal representation learning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.908933Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:ffa64973e34fdb859c4a6f8d1c02d9f34a4c8191dc314edb0bc17bf8996c7272","observation_id":"cf025ddc-bea3-4353-a8e9-7385ee01b7e9","resolution":{"observed_at":"2026-08-07T22:12:49.908933Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:49.913105Z","title":"and Li, Y","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.913105Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:b9bcf2663ed1f36654cabb04ef28f751048db9632ea72e3e46af3aea9dedeb8e","observation_id":"027e45ae-e114-4713-b6d0-c2d904f2fa6f","resolution":{"observed_at":"2026-08-07T22:12:49.913105Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:49.917187Z","title":"V., Pillaud-Vivien, L., and Flammarion, N","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.917187Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:b08718c71b2d05c0e22e765f752ddd61a6b9fd9d1a674eef3753f49200985eb2","observation_id":"c4597850-61a1-44f1-823e-9cdd86b2d903","resolution":{"observed_at":"2026-08-07T22:12:49.917187Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:49.921275Z","title":"Exploring length generalization in large language models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.921275Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:6a6f8cfd9a94fbf7674131a38a3af07a0115680208f9bcb2271bb82546773ecb","observation_id":"ccb3f2de-8324-4560-96bb-053d86448332","resolution":{"observed_at":"2026-08-07T22:12:49.921275Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1907.02893","last_updated":"2020-03-27T19:07:58Z","snapshot_observed_at":"2026-07-06T08:05:24.076802Z","submitted_at":"2019-07-05T15:26:26Z","title":"Invariant Risk Minimization","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1907.02893","snapshot_observed_at":"2026-08-07T22:12:49.925365Z","title":"Invariant risk minimization","venue":null,"work_id":null,"year":1907},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.925365Z"},"links":{"cited_paper":"/paper/1907.02893","citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:5a0dd1e2fdaaeb596d26d23a506a63c2a28c1d4f7733b993fc95b3efad496f0c","observation_id":"9555662b-8e41-4491-a3ad-a5df4b31bed8","resolution":{"observed_at":"2026-08-07T22:12:49.925365Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:49.930254Z","title":"A theoretical analysis of contrastive unsupervised representation learning","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.930254Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:e6fcab78e280dcd71d7208181864d7a0899abbbe7fa03cf97e7760b94c0114cc","observation_id":"a9ba07bd-3734-410a-be3b-713cde8ed9a2","resolution":{"observed_at":"2026-08-07T22:12:49.930254Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:49.934349Z","title":null,"venue":null,"work_id":null,"year":2002},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.934349Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:87fedc86ffcce3ce9562039214f019e2c708771b6034a40341a3d99477ed43d8","observation_id":"18ca27b8-30d6-4f23-8440-e4fec2e649b8","resolution":{"observed_at":"2026-08-07T22:12:49.934349Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:49.938639Z","title":"L., Foster, D","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.938639Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:5d7018d1b3572ac390f6499d90cb290780d9f52440c46ab8f52cefa1b76afd2e","observation_id":"f569303d-eeaa-4229-8fa4-6a6f9dc49606","resolution":{"observed_at":"2026-08-07T22:12:49.938639Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:49.942469Z","title":"L., Long, P","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.942469Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:abd9d37a054893d04812fb99946056c57c3f8873e712570ef062e8f0bde43d5b","observation_id":"e1624327-b057-4ea5-a1ca-3bf78d351c9b","resolution":{"observed_at":"2026-08-07T22:12:49.942469Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:49.946350Z","title":"M., Gebru, T., McMillan-Major, A., and Shmitchell, S","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.946350Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:d8f61d598da6a1a05e0ea699ec7556b9dba4ab4680ccee575897436938c91ded","observation_id":"bd61ba73-c495-4874-978c-f51a93bb8332","resolution":{"observed_at":"2026-08-07T22:12:49.946350Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:49.950312Z","title":"Representation learning: A review and new perspectives","venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.950312Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:edb12ec8d4106e83a3d10194e4e9c85418572a4b5f9ec70f552593feaa4fcb35","observation_id":"2b20971b-59a5-4820-90d3-9e8d3406bf53","resolution":{"observed_at":"2026-08-07T22:12:49.950312Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:49.954065Z","title":"fail to learn ``b is a","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.954065Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:d3d0059eec78c840d2b60174457d96b3ad65d4fcab8accf54bd1e45bb4a6b45a","observation_id":"027e513c-07d8-4e5d-8a5e-4c3a248950c4","resolution":{"observed_at":"2026-08-07T22:12:49.954065Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.12316","last_updated":"2023-07-10T17:15:23Z","snapshot_observed_at":"2026-08-03T17:40:24.585237Z","submitted_at":"2022-11-22T15:10:48Z","title":"Simplicity Bias in Transformers and their Ability to Learn Sparse Boolean Functions","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.12316","snapshot_observed_at":"2026-08-07T22:12:49.957609Z","title":"Simplicity bias in transformers and their ability to learn sparse boolean functions","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.957609Z"},"links":{"cited_paper":"/paper/2211.12316","citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:08e77e39f8f74a44b1389262a847625f63902b6b7e99f64ff8a112167940c177","observation_id":"651f7b2d-6cce-49cd-a35d-cc05da07108b","resolution":{"observed_at":"2026-08-07T22:12:49.957609Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:49.962139Z","title":"Towards monosemanticity: Decomposing language models with dictionary learning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.962139Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:0c6e0cb4ee6490826821731c9aa96c1ba739a1a5e883c6ac9fe98b953361218b","observation_id":"6e0fc7f8-41b2-4f4f-9cd3-a48ff9e18fdc","resolution":{"observed_at":"2026-08-07T22:12:49.962139Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:49.965870Z","title":"Video generation models as world simulators, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.965870Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:ec0d068c5cd77eb1322ff92f3d95ba2db492341b0bec47333bf7535c1955aba6","observation_id":"6205c7eb-3770-4a57-beb8-eb9a7dc38ce1","resolution":{"observed_at":"2026-08-07T22:12:49.965870Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.479184Z","title":null,"venue":null,"work_id":"8c484651-06ae-434e-aff9-173d4d047424","year":2020},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.969809Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:2c6837da7788c35760e904c10f2f7ff4b521d6d9e524776d0f08f686e1f38df5","observation_id":"81f8f6ff-baec-490d-9c74-07595a22c44d","resolution":{"observed_at":"2026-08-07T22:12:51.484143Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2409.12446","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.735502Z","title":"and Sudijono, T","venue":null,"work_id":"ce1013cb-2268-4824-b1f8-2322698734d8","year":2024},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.973561Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:0e75fd94fbc0568adb00785bd7d499d76eac17742e913fbaab3a4d442334bcfb","observation_id":"23111ca0-3328-4835-b79e-4975c07cb657","resolution":{"observed_at":"2026-08-07T22:12:50.742168Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.466587Z","title":"Implicit bias of gradient descent for wide two-layer neural networks trained with the logistic loss","venue":null,"work_id":"13500d4f-72b3-4dd0-8673-42d1dd632a19","year":2020},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.977266Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:1f849114baffed3f616b84d73bce2c8bf3b6181212895b0cda860c1e54ee844a","observation_id":"0fdeaba2-0fd7-4e6e-99b8-c3232bfd77d1","resolution":{"observed_at":"2026-08-07T22:12:51.471050Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.454863Z","title":null,"venue":null,"work_id":"975cb53f-46f1-4b0e-bb5a-7e9a5df590c6","year":1967},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.980848Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:56964d284da78df06e10f6d79ccd7d2223205412949c39cef9777335c2d296ba","observation_id":"f9d1dcfe-4a50-49e9-88cd-6cfcb761fd56","resolution":{"observed_at":"2026-08-07T22:12:51.458513Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.444133Z","title":"Q., and Louis, A","venue":null,"work_id":"c1725688-86f1-4a76-aeab-961218594cdb","year":2018},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.985044Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:5e23130d6666cac182dfdd755b940d37640bed707fc09c6e7d63c529430b75aa","observation_id":"b2975aa4-1285-45e4-9310-9b3fb8afa323","resolution":{"observed_at":"2026-08-07T22:12:51.447850Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.433394Z","title":"An introduction to latent variable models","venue":null,"work_id":"1962978e-9074-4c4c-8561-4b3baed19aa7","year":2013},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.988801Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:cf48110e966006a4dcb59339d748e2c628f1fbd1da43cd6b766d4c4f6abbcfa7","observation_id":"5d026e3e-52bc-4b6f-8453-598850233a20","resolution":{"observed_at":"2026-08-07T22:12:51.437078Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.421815Z","title":"J., Nagai, Y., Taniguchi, T., Gomi, H., and Tenenbaum, J","venue":null,"work_id":"ae1e66f1-50b5-463b-888c-0ae80f3e2748","year":2021},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.992380Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:1847c1ebeffbeadf63a7e1cb12c2ae995ffecc13d378fca64f5e4d2e9bf76d88","observation_id":"4f9669c4-fd26-4da4-826e-8a103282b660","resolution":{"observed_at":"2026-08-07T22:12:51.425459Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.411064Z","title":"On the approximate realization of continuous mappings by neural networks","venue":null,"work_id":"ca8c4132-1321-4342-adac-3a99920c3103","year":1989},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:49.996138Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:1d4ae1bc07945b7bde8010f2e5e657bd47155786721c3bf909f353b1e95c6630","observation_id":"e00d8db7-8ccf-423a-bf02-57ed19b69658","resolution":{"observed_at":"2026-08-07T22:12:51.414924Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.399610Z","title":"A., and Brendel, W","venue":null,"work_id":"ed4a8127-28fc-4356-a3c7-3183855a9270","year":2019},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.000167Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:b73d9be1aa57674241ee5b12062d2ebfea9a4b77d04a021dcaf3ef062e727618","observation_id":"0434e89f-ff80-485e-b5ab-40c1380b4ce3","resolution":{"observed_at":"2026-08-07T22:12:51.404072Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2103.13630","last_updated":"2021-06-21T21:01:12Z","snapshot_observed_at":"2026-07-06T10:53:11.986981Z","submitted_at":"2021-03-25T06:57:11Z","title":"A Survey of Quantization Methods for Efficient Neural Network Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2103.13630","snapshot_observed_at":"2026-08-07T22:12:50.004131Z","title":"W., and Keutzer, K","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.004131Z"},"links":{"cited_paper":"/paper/2103.13630","citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:412cc163370362f3348cba2618255d63ed5083727b8f55e2184c86025b9dad38","observation_id":"6de8f12d-cb44-47f3-9775-b6c5d410df9c","resolution":{"observed_at":"2026-08-07T22:12:50.004131Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.387509Z","title":null,"venue":null,"work_id":"414d6d0b-13cc-4502-b6fc-3814370e986f","year":2024},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.008325Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:2c911690d93972f61f6d407c026277f2139c17818a0dca63c98d1426a9cd71ad","observation_id":"565f5320-396a-476c-b25b-801ddf0a0a10","resolution":{"observed_at":"2026-08-07T22:12:51.391614Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2011.15091","last_updated":"2022-08-01T13:58:56Z","snapshot_observed_at":"2026-07-06T10:19:12.615247Z","submitted_at":"2020-11-30T18:29:25Z","title":"Inductive Biases for Deep Learning of Higher-Level Cognition","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2011.15091","snapshot_observed_at":"2026-08-07T22:12:50.012673Z","title":"and Bengio, Y","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.012673Z"},"links":{"cited_paper":"/paper/2011.15091","citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:38d0c325259409ff26ea8d1a37d597d62b51e186a89a06c722fde496a64ac9b3","observation_id":"2daaa960-61e1-4fd6-8eb9-80278208fb19","resolution":{"observed_at":"2026-08-07T22:12:50.012673Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.375106Z","title":"Characterizing implicit bias in terms of optimization geometry","venue":null,"work_id":"5c9b5c9e-505b-4cf0-8401-eecfa2dda6e5","year":2018},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.017385Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:3a7d8052d15e75a220c464212005ff040634022567a88d25682e7eeffb6e8fd2","observation_id":"253af9c5-fe3b-4c00-b6b5-2a186ed3dfad","resolution":{"observed_at":"2026-08-07T22:12:51.379350Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.362576Z","title":"D., Soudry, D., and Srebro, N","venue":null,"work_id":"e8434117-37f6-45bd-b134-f9acc51aa547","year":2018},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.021248Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:e2a30e8617b8a5079bef7c91332e18cf426c4a3bf8fb96379b56e35da3d41344","observation_id":"5a00f6a6-e471-4af4-8dd1-0946aa292b97","resolution":{"observed_at":"2026-08-07T22:12:51.366609Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.351886Z","title":"and Tegmark, M","venue":null,"work_id":"c642baa7-acce-47f6-b409-c22baaf763b9","year":2024},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.025070Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:bedbe214a2992806f0b1b53bca65c2d49e0d391dbef0ffe5b611c604d6f3052c","observation_id":"af77a08a-f0ed-447d-b72f-71283a79f1b6","resolution":{"observed_at":"2026-08-07T22:12:51.355575Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1803.10122","last_updated":"2018-05-09T09:06:27Z","snapshot_observed_at":"2026-07-31T21:36:45.596575Z","submitted_at":"2018-03-27T15:08:55Z","title":"World Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1803.10122","snapshot_observed_at":"2026-08-07T22:12:50.028734Z","title":"and Schmidhuber, J","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.028734Z"},"links":{"cited_paper":"/paper/1803.10122","citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:1e6fb2e7753b81f44e2e11e431cab3f31146065f7d418f519bfa87eddcbe3eda","observation_id":"56212e73-2742-45ad-a217-78db692550fe","resolution":{"observed_at":"2026-08-07T22:12:50.028734Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.339375Z","title":"Masked autoencoders are scalable vision learners","venue":null,"work_id":"a6c5ab1b-d6ac-422b-a90e-817eb91eeb9e","year":2022},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.033145Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:7a9ebf6ce32c5a84b019ecbb950e7107882dde0daaf17c5b4119593a00a79b7a","observation_id":"07313a32-312c-42d1-909f-2945eb1d8d69","resolution":{"observed_at":"2026-08-07T22:12:51.343765Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.12001","last_updated":"2023-10-09T22:57:01Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-06-21T03:35:06Z","title":"An Overview of Catastrophic AI Risks","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.12001","snapshot_observed_at":"2026-08-07T22:12:50.037040Z","title":"An overview of catastrophic AI risks","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.037040Z"},"links":{"cited_paper":"/paper/2306.12001","citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:8d7bed24fb07c33b3ae18b8d77830c4097373618fd6ae186b43dc86083b3667b","observation_id":"d727f540-3a4d-4be9-a03a-22920e70273e","resolution":{"observed_at":"2026-08-07T22:12:50.037040Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.326553Z","title":"Multilayer feedforward networks are universal approximators","venue":null,"work_id":"0290cbdb-4e73-485f-bb11-46ee7e16153a","year":1989},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.041310Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:98aaea3327f0a33c6796ab7bff57098f5eb86b480318a7e7c7e5c025c18fac8c","observation_id":"8d07a6a0-8b8c-4416-a619-b7d42f506ae1","resolution":{"observed_at":"2026-08-07T22:12:51.331038Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.045189Z","title":"The low-rank simplicity bias in deep networks","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.045189Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:38fa07af4e2bfda49b4f0eeef833a3ca0d48896317d8e242e6b6d94d3bc41280","observation_id":"b214df8c-250d-4ff7-be1e-f5db1959fe37","resolution":{"observed_at":"2026-08-07T22:12:50.045189Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.307806Z","title":"The Platonic representation hypothesis","venue":null,"work_id":"208d9fe5-e193-4ffa-8b46-843e80e1e329","year":2024},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.048900Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:ff282cf4cbef19024b600f36b91df031876a8ef0f3dcc657d216b9013e7d1a8e","observation_id":"524469f1-7e19-47e9-829d-851e5b847ef7","resolution":{"observed_at":"2026-08-07T22:12:51.311716Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.297004Z","title":"and Pajunen, P","venue":null,"work_id":"d0ae6ffc-5c7c-4f7f-ada2-ba46b270f6a4","year":1999},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.052538Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:e98ccb4df37f8ae75c2a8d844f71443e723c956f907d08f38bf394e090034fae","observation_id":"98fee8e2-6723-4cf6-a032-301ff4bfd111","resolution":{"observed_at":"2026-08-07T22:12:51.300868Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.285575Z","title":null,"venue":null,"work_id":"f1c554dd-df3d-4732-aa15-83cb38b7bfd0","year":2019},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.056584Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:1f89366ed7bd9cd10668176cf1e393e45e4ae402e4dcd6783a138c2b0d873a5b","observation_id":"e5238713-1f81-46a9-bf45-42f5aa6cb62a","resolution":{"observed_at":"2026-08-07T22:12:51.289481Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.272861Z","title":"Neural tangent kernel: Convergence and generalization in neural networks","venue":null,"work_id":"127d1c12-f260-458a-a913-4465d5d7980e","year":2018},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.060426Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:31b8bec193dae9e8159f62f63e732d4217bb3578b1e9d0e5034870bbb20b83aa","observation_id":"a66ec036-12f5-4770-9bfb-02158e6379fe","resolution":{"observed_at":"2026-08-07T22:12:51.277345Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.260912Z","title":"Low -resource","venue":null,"work_id":"a50c28e5-1608-456c-82f2-86acd21a98ec","year":2023},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.065175Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:fc0a4dd17294e119ff99b8bbe2464bceb615582f92f0ada8dc6f3deb1ec95806","observation_id":"53f8925b-d809-4f7f-a6eb-9195f599ebd0","resolution":{"observed_at":"2026-08-07T22:12:51.264673Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.249641Z","title":"and Rinard, M","venue":null,"work_id":"f77910d6-7ae0-4937-b9a6-d8152c1b4052","year":2024},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.069104Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:1df9150912b195d6cf126e8904e4c05cff9bff1e833220e414ff0261127934c5","observation_id":"383d5e46-a0ab-47fa-96fc-e8d47b749c2e","resolution":{"observed_at":"2026-08-07T22:12:51.253389Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.237444Z","title":"SGD on neural networks learns functions of increasing complexity","venue":null,"work_id":"157ce29b-5cfd-430f-a2f9-28ecce37e3b1","year":2019},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.072865Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:b8d03275139e4b62c2c5a81aac3fa70e1d914780dd5f51116aad8a038285ef1d","observation_id":"71e72b80-285b-48d8-9f25-fe45fcf00e3a","resolution":{"observed_at":"2026-08-07T22:12:51.241544Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.225626Z","title":"How far is video generation from world model: A physical law perspective","venue":null,"work_id":"6873fe48-99e4-4a79-8390-3e8baf1f5478","year":2025},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.076745Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:0bed716297ba6ae04cda34252001776d7bf7098024f09977f1637559e3a9458c","observation_id":"4079ad2c-8bb3-4f4d-83e3-9582b82710e3","resolution":{"observed_at":"2026-08-07T22:12:51.229784Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.213410Z","title":"P., Monti, R","venue":null,"work_id":"d71a785d-12e0-47fe-95d3-8627b93ae1f4","year":2020},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.080610Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:4b73d9be7e72793fd52eb40abb0198eaf618a2dc20f35c68df3c8e74a02ef09a","observation_id":"7b6679a6-812b-4474-be8a-698245e27158","resolution":{"observed_at":"2026-08-07T22:12:51.217693Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.085095Z","title":"A path towards autonomous machine intelligence version 0.9","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.085095Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:a1c90767dfea541b31285a2ee3922309586103876de82a904d09d7b4c3f97d9a","observation_id":"350e01bf-faa8-4276-8a9e-6df604a566e5","resolution":{"observed_at":"2026-08-07T22:12:50.085095Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.194372Z","title":"D., Lei, Q., Saunshi, N., and Zhuo, J","venue":null,"work_id":"4a6beec0-7823-41a2-b9a9-88f4422f7fcb","year":2021},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.089203Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:4f344982128a845098025ea8ba5f674e51b6f9c6e800e01d3dc2260de7f8056d","observation_id":"a33d1aeb-66ea-49c1-bf63-3685a815c831","resolution":{"observed_at":"2026-08-07T22:12:51.198434Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2106.00737","last_updated":"2021-06-01T19:23:20Z","snapshot_observed_at":"2026-08-06T17:41:42.388577Z","submitted_at":"2021-06-01T19:23:20Z","title":"Implicit Representations of Meaning in Neural Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2106.00737","snapshot_observed_at":"2026-08-07T22:12:50.093046Z","title":"Z., Nye, M., and Andreas, J","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.093046Z"},"links":{"cited_paper":"/paper/2106.00737","citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:aeed6e64253010eb10c953c142fb4afdc2b979390af16f355233f23c9cdc96e8","observation_id":"e1eb8231-2132-4d55-a3f9-59afddfae5aa","resolution":{"observed_at":"2026-08-07T22:12:50.093046Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.183274Z","title":"K., and Bau, D","venue":null,"work_id":"a7bc689e-1193-4b1e-b821-37937b84ffbd","year":2023},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.097291Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:559838bea578d9d90270afd205b498c8a66b578e177492a86919af72ea046fba","observation_id":"2a29a239-b4a6-4f92-825a-98bf1b1045c2","resolution":{"observed_at":"2026-08-07T22:12:51.186932Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.171536Z","title":"An introduction to Kolmogorov complexity and its applications, volume 3","venue":null,"work_id":"72eb72bc-c3b0-4789-a969-f3f96f18696a","year":2008},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.101073Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:cafa140299593a171e62dd76d8031ff77583379c500c28ea271c445a3130d5fb","observation_id":"25e02eed-2f1d-4d43-82fc-1c66b6f12667","resolution":{"observed_at":"2026-08-07T22:12:51.175433Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.160373Z","title":"Self-supervised learning: Generative or contrastive","venue":null,"work_id":"be01c2a6-6ab0-40ad-89bd-3b3dd5844ac2","year":2021},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.104742Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:ce9b990d7b77fc3280f928672027855ba5febe744469501a7461f2ec1d9c8efc","observation_id":"1c07cacc-babe-4582-acac-031d88a283cd","resolution":{"observed_at":"2026-08-07T22:12:51.164172Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.05918","last_updated":"2023-10-09T17:59:18Z","snapshot_observed_at":"2026-08-02T03:11:30.434124Z","submitted_at":"2023-10-09T17:59:18Z","title":"Grokking as Compression: A Nonlinear Complexity Perspective","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.05918","snapshot_observed_at":"2026-08-07T22:12:50.108574Z","title":"Grokking as compression: A nonlinear complexity perspective","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.108574Z"},"links":{"cited_paper":"/paper/2310.05918","citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:5b6fafb5717e25ed144db421d83f2b8c0df89722b8b259ec1f36a3e3c458a638","observation_id":"68916361-4f92-48da-a5dc-0d6d628b0698","resolution":{"observed_at":"2026-08-07T22:12:50.108574Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.148374Z","title":"and Hutter, F","venue":null,"work_id":"cdc1a6ab-7238-421a-949d-7b2dfc4a86ea","year":2019},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.112683Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:df2554fe0edf2b02ec6daf33aed88adc466a3a2283ea60281092c28e77c9d0c3","observation_id":"9b908d7b-c5b9-4cf7-8d7e-02d8555a4616","resolution":{"observed_at":"2026-08-07T22:12:51.152233Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.135824Z","title":null,"venue":null,"work_id":"b669e0a4-88b4-477c-867d-d9a41b05a449","year":2022},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.116402Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:61f0c36d4071ac50c438c89ac16d02c8c63b50f6d99154c19c517460f354438a","observation_id":"81ca0171-85cf-40a0-93ee-46bcaa6aa881","resolution":{"observed_at":"2026-08-07T22:12:51.140387Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.124322Z","title":"Gradient descent on two-layer nets: Margin maximization and simplicity bias","venue":null,"work_id":"46780fd5-4f57-4f91-a20a-036a9144d82f","year":2021},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.120098Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:6b4be8e65dbccabc95da699cfe0ed299aced893cc29b3b326b64ecce3ffd301e","observation_id":"9fbee53d-5271-43b2-bfc5-d8305c9613d5","resolution":{"observed_at":"2026-08-07T22:12:51.128596Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.06824","last_updated":"2024-08-19T01:18:41Z","snapshot_observed_at":"2026-07-06T16:30:37.867641Z","submitted_at":"2023-10-10T17:54:39Z","title":"The Geometry of Truth: Emergent Linear Structure in Large Language Model Representations of True/False Datasets","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.06824","snapshot_observed_at":"2026-08-07T22:12:50.123935Z","title":"and Tegmark, M","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.123935Z"},"links":{"cited_paper":"/paper/2310.06824","citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:ef542d5074705de9e8e6a834dc6027cb03bbd51d173718906198892d0f3b01e2","observation_id":"c6d99339-6de8-46f0-aef0-6f5b6c9baa73","resolution":{"observed_at":"2026-08-07T22:12:50.123935Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.112689Z","title":"Linguistic regularities in continuous space word representations","venue":null,"work_id":"c14b22e3-037a-4267-895b-9eecdffd82b3","year":2013},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.128194Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:fd1409fce3563421764c05bdfe84c12501677d66f635305b3f7214ee1cb3adb7","observation_id":"390608a4-27da-4738-83d6-f55bc11999ac","resolution":{"observed_at":"2026-08-07T22:12:51.117106Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05229","last_updated":"2025-08-27T16:24:39Z","snapshot_observed_at":"2026-07-06T19:29:09.714725Z","submitted_at":"2024-10-07T17:36:37Z","title":"GSM-Symbolic: Understanding the Limitations of Mathematical Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05229","snapshot_observed_at":"2026-08-07T22:12:50.131856Z","title":"GSM -symbolic: Understanding the limitations of mathematical reasoning in large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.131856Z"},"links":{"cited_paper":"/paper/2410.05229","citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:620212e3d445e543fde01505aa421067d00beaacac0bf4abf622c6f666fa86b9","observation_id":"35dd6bf2-88c7-4a6c-a34c-0603ec54d8e3","resolution":{"observed_at":"2026-08-07T22:12:50.131856Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.101598Z","title":"Ai’s challenge of understanding the world","venue":null,"work_id":"b3044062-804a-43b4-a4d1-d54abfa7038a","year":2023},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.135699Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:71680a28aa28f4fa0f7c74aec14e054e5cc28818460e026a31f930047a93f916","observation_id":"3522f7f9-ae42-4232-87a7-3c56b59c3d90","resolution":{"observed_at":"2026-08-07T22:12:51.105905Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09038","last_updated":"2025-02-27T15:10:51Z","snapshot_observed_at":"2026-08-02T17:18:33.867969Z","submitted_at":"2025-01-14T20:59:37Z","title":"Do generative video models understand physical principles?","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.09038","snapshot_observed_at":"2026-08-07T22:12:50.139386Z","title":"Do generative video models learn physical principles from watching videos? arXiv preprint arXiv:2501.09038, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.139386Z"},"links":{"cited_paper":"/paper/2501.09038","citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:09dc8b7faa21b440155b957a9a4f3cdd74c611afaad32adcf28a20861595d253","observation_id":"421c1ee7-b03e-48f4-8b0c-3947c10e3a52","resolution":{"observed_at":"2026-08-07T22:12:50.139386Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2106.08295","last_updated":"2021-06-15T17:12:42Z","snapshot_observed_at":"2026-08-02T11:19:40.664702Z","submitted_at":"2021-06-15T17:12:42Z","title":"A White Paper on Neural Network Quantization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2106.08295","snapshot_observed_at":"2026-08-07T22:12:50.143399Z","title":"A., Bondarenko, Y., Baalen, M","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.143399Z"},"links":{"cited_paper":"/paper/2106.08295","citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:d9900846e87d3030e014c820917945ac48786fab8940298fc85b03dec7af65a1","observation_id":"c1861d6e-0c06-432d-97fd-4ee9eb30e478","resolution":{"observed_at":"2026-08-07T22:12:50.143399Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.00941","last_updated":"2023-09-07T20:36:48Z","snapshot_observed_at":"2026-07-06T16:13:34.788427Z","submitted_at":"2023-09-02T13:37:34Z","title":"Emergent Linear Representations in World Models of Self-Supervised Sequence Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.00941","snapshot_observed_at":"2026-08-07T22:12:50.147169Z","title":"Emergent linear representations in world models of self-supervised sequence models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.147169Z"},"links":{"cited_paper":"/paper/2309.00941","citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:60334db7177a1b04dc2f4cb15d44000bf215d43976786c4ec63678edfa7a7373","observation_id":"7b55d802-9926-4402-8e5f-7d467a5a5e89","resolution":{"observed_at":"2026-08-07T22:12:50.147169Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.090849Z","title":"Analysis of boolean functions","venue":null,"work_id":"c71a54d8-045b-4e27-ae2d-9a10a3356edd","year":2014},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.152096Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:f79485e9914e3b85b03f58f57db275301beb7e9f5d72c0bdfca08ef1518c01b6","observation_id":"a672be43-3bbd-4148-a1c9-0fec5e7a3367","resolution":{"observed_at":"2026-08-07T22:12:51.094410Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.079058Z","title":"PyTorch : An imperative style, high-performance deep learning library","venue":null,"work_id":"154e7f7f-b958-467d-b3fc-d8b0359957ff","year":2019},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.156251Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:54d00e5915e0d30bf9166d30f08e4acf132721f1eab4f2d395a1740b0310acd4","observation_id":"5b7d7b98-8faa-4903-8c45-034d6885121e","resolution":{"observed_at":"2026-08-07T22:12:51.083188Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.066384Z","title":"Train short, test long: Attention with linear biases enables input length extrapolation","venue":null,"work_id":"50295064-baa3-4534-895a-4b6a1fb8dfea","year":2022},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.160268Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:a723ae2585911cae24df64f9cee648c434321723c2bf1fffa491572bd146d702","observation_id":"38d51689-b6f8-46fe-9546-220d890d29ae","resolution":{"observed_at":"2026-08-07T22:12:51.071357Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.054202Z","title":"V., Louis, A","venue":null,"work_id":"9375b800-686a-41a4-ba31-8ec7b6df7943","year":2019},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.163987Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:385409eac5d8c2abe5a95ba94d71ce4985a316caa306138deb41aba104aed88d","observation_id":"e6ebed76-4596-4b0b-a685-5193fd053587","resolution":{"observed_at":"2026-08-07T22:12:51.057879Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.042777Z","title":"Language models are unsupervised multitask learners","venue":null,"work_id":"c965a4c6-289e-431c-90d7-0cfea0f18ded","year":2019},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.167942Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:18e066e0de40df5604eee482e62ac9f373da90186f3fdbd8d233276a9df3e926","observation_id":"8d940cb7-9d52-4aea-8b28-e32391915644","resolution":{"observed_at":"2026-08-07T22:12:51.046693Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.030412Z","title":"Do vision transformers see like convolutional neural networks? Advances in Neural Information Processing Systems , 34: 0 12116--12128, 2021","venue":null,"work_id":"4d937650-d0a8-4afe-8667-2e13d4547157","year":2021},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.172371Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:cdb33eb3766c65b952544aa90d6e052286e1d3f93bdd20671f4087d271e20991","observation_id":"020f9ca2-5dc2-4799-9415-57c3a86f9c50","resolution":{"observed_at":"2026-08-07T22:12:51.035317Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:51.016762Z","title":"Position: Understanding LLMs requires more than statistical generalization","venue":null,"work_id":"29ede3ec-1f28-443d-8549-abea713017f3","year":2024},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.177662Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:9e90db2552a3641447c6dbf19fd321a9a8f48d9248335e181286052f75823d82","observation_id":"2bb49730-e3a3-495f-b19b-f878c5c06800","resolution":{"observed_at":"2026-08-07T22:12:51.022043Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.181766Z","title":"and Everitt, T","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":71,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.181766Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:3a9b7c0cee6048a1284c85b013405489b12cdad66c013eef899c1ef04c1b1ad3","observation_id":"bee7d0ec-dfe8-40aa-aa6e-c00095a9ffca","resolution":{"observed_at":"2026-08-07T22:12:50.181766Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.998349Z","title":"R., Kalchbrenner, N., Goyal, A., and Bengio, Y","venue":null,"work_id":"2d88540b-656c-421b-bac2-d1a9fb5f53f3","year":2021},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":72,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.186148Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:00c9dcfbd2f8c7c9c25217e5d7c603fc78a20c0c9980290e7a30cb835bdf580f","observation_id":"5c7895b0-8219-4ba1-b032-6ca424be538a","resolution":{"observed_at":"2026-08-07T22:12:51.002176Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.190104Z","title":"S., Gunasekar, S., and Srebro, N","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":73,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.190104Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:42b52271ad4fcf38efa86fcef607c4c2a759bcaf8d1c52715990f450fa976a10","observation_id":"f9762ea6-0fd8-4e2b-98e1-cf1671905d9e","resolution":{"observed_at":"2026-08-07T22:12:50.190104Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.02241","last_updated":"2025-04-29T19:25:29Z","snapshot_observed_at":"2026-07-06T17:39:20.748463Z","submitted_at":"2024-03-04T17:33:20Z","title":"Neural Redshift: Random Networks are not Random Functions","version":3},"cited_work":{"arxiv_id":"2403.02241","doi":null,"metadata_source":"pith","pith_arxiv_id":"2403.02241","snapshot_observed_at":"2026-08-07T22:12:50.361579Z","title":"Neural Redshift: Random Networks are not Random Functions","venue":"cs.LG","work_id":"522bda73-9a1a-49eb-b321-4be18011f00d","year":2024},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":74,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.194047Z"},"links":{"cited_paper":"/paper/2403.02241","citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:4c11a1f56309773f9d5eb996c009c7370e6cfae7764403f65a262a758a02d15a","observation_id":"03b92ed4-d460-49c7-a385-07b87f148495","resolution":{"observed_at":"2026-08-07T22:12:50.368226Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.980236Z","title":"Contrastive learning, multi-view redundancy, and linear models","venue":null,"work_id":"d38a618f-00ff-4366-a68b-61c1aa9c1dc4","year":2021},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":75,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.198467Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:c4064e89d6b2ea58ca5e6fd71258dbe2f3cfcfc547f2819dc8bffc923987a564","observation_id":"6493b6b1-f067-4419-bbd1-fc4291583ef9","resolution":{"observed_at":"2026-08-07T22:12:50.984141Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.968897Z","title":"The nature of statistical learning theory","venue":null,"work_id":"14e235f0-8216-4cd7-a9b0-55ea241272c9","year":1999},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":76,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.202869Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:6ca4d8c766c1af71fa4a8ca2d63df14ab6e51c8ec1c9118206d460bdba128cec","observation_id":"3c6ad341-62ba-4c4a-9805-49a9dfcfa7e5","resolution":{"observed_at":"2026-08-07T22:12:50.973002Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.957880Z","title":"N., Kaiser, L., and Polosukhin, I","venue":null,"work_id":"95d66acf-b8ce-46a7-99a4-daef42fce90e","year":2017},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.207001Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:a87c16255465c283fa0b9880c4a9870b19f9076b33cb2ece5c7495fcafbaff27","observation_id":"47e3215b-5b8f-4479-bfe3-356a2c036b38","resolution":{"observed_at":"2026-08-07T22:12:50.961569Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.945280Z","title":"Self-supervised learning with data augmentations provably isolates content from style","venue":null,"work_id":"b4bc442a-2de6-4107-a7fe-0820f070f001","year":2021},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.210865Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:c196bdeb72dd769a0099a4eafa428f57fc99044f83fc8795c92b5ee9945ab317","observation_id":"e37d0edc-05d1-4eca-bcb0-9bb1422a9011","resolution":{"observed_at":"2026-08-07T22:12:50.949418Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.934037Z","title":"Nonparametric identifiability of causal representations from unknown interventions","venue":null,"work_id":"cadb7a60-09a1-4cd5-a5be-349c74776680","year":2023},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":79,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.214570Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:a314b9e42a2bc1eee0d2446ac3584721ca580216c989bc5553f7f1f7be00040d","observation_id":"178533ce-0f40-4ae0-a51d-c638b15c0bfb","resolution":{"observed_at":"2026-08-07T22:12:50.937893Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.922618Z","title":"M., and Ma, T","venue":null,"work_id":"a061e780-1b24-4de6-8bc9-96f357535c01","year":2021},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":80,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.218788Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:614a722f6e898baaea9902c7184f58c2ced48f7d1b8ee59db4396a41299d95d6","observation_id":"086d8c97-ea7b-4806-a50a-1fe1cf4fe9c9","resolution":{"observed_at":"2026-08-07T22:12:50.926578Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.910395Z","title":"and Watson, J","venue":null,"work_id":"d6cab5ad-e488-46b9-9418-f3b0bc70c3de","year":2005},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":81,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.223010Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:482249ba56b01e659b1b3094d78f4487247d986a9aeb1801b4a5f67ed4b499fd","observation_id":"58566301-f9be-45cc-93cd-bf7213443ef5","resolution":{"observed_at":"2026-08-07T22:12:50.914569Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.896477Z","title":null,"venue":null,"work_id":"38f42e0e-11ea-4803-8431-646ff617a5e4","year":1996},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":82,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.226871Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:b80298e30de91aa4bdfbaf31d90ddca50f5f4678e98a904d1a1f3e9f261b0696","observation_id":"bfba607c-070b-42e1-8765-f4dda9082cb8","resolution":{"observed_at":"2026-08-07T22:12:50.901554Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.12672","last_updated":"2023-06-23T06:05:31Z","snapshot_observed_at":"2026-07-06T15:45:24.674515Z","submitted_at":"2023-06-22T05:14:00Z","title":"From Word Models to World Models: Translating from Natural Language to the Probabilistic Language of Thought","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.12672","snapshot_observed_at":"2026-08-07T22:12:50.230706Z","title":"K., Goodman, N","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":83,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.230706Z"},"links":{"cited_paper":"/paper/2306.12672","citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:45d15ec176eb04824018cf554e345982d3414c0cffc732a7f3805a011b26d4bc","observation_id":"cbc82b42-5c0c-41b8-97f9-df9927e62127","resolution":{"observed_at":"2026-08-07T22:12:50.230706Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.02477","last_updated":"2024-03-28T23:37:24Z","snapshot_observed_at":"2026-07-06T15:50:41.385379Z","submitted_at":"2023-07-05T17:50:42Z","title":"Reasoning or Reciting? Exploring the Capabilities and Limitations of Language Models Through Counterfactual Tasks","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.02477","snapshot_observed_at":"2026-08-07T22:12:50.234654Z","title":"Reasoning or reciting? Exploring the capabilities and limitations of language models through counterfactual tasks","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":84,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.234654Z"},"links":{"cited_paper":"/paper/2307.02477","citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:bfc64a930823d5f969872ed9a0d920f14cf2cb4c3a0a00f781a53100f2c7712b","observation_id":"7c070b59-3f4c-41e0-b02b-0124c0af69c5","resolution":{"observed_at":"2026-08-07T22:12:50.234654Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.12278","last_updated":"2024-10-02T23:37:50Z","snapshot_observed_at":"2026-07-06T19:17:51.711221Z","submitted_at":"2024-09-18T19:28:04Z","title":"Making Large Language Models into World Models with Precondition and Effect Knowledge","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.12278","snapshot_observed_at":"2026-08-07T22:12:50.238938Z","title":"Making large language models into world models with precondition and effect knowledge","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":85,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.238938Z"},"links":{"cited_paper":"/paper/2409.12278","citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:518917d6392b5572882ffea080d1d750c728a5db25cc906387a5fddf9ecfe45d","observation_id":"be03045d-3c11-40a2-8474-d83bb899ba1f","resolution":{"observed_at":"2026-08-07T22:12:50.238938Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.884363Z","title":"S., Kawarabayashi, K.-i., and Jegelka, S","venue":null,"work_id":"f9f8eaa5-4ae7-4d36-8a81-a64ee14b844b","year":2020},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":86,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.242996Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:9d1ecbba6b2a65c82493db25ea323512b2a264f55e987de76be9999d6d1769fb","observation_id":"da98412d-ceaf-4cb7-9b5b-11ca6de19073","resolution":{"observed_at":"2026-08-07T22:12:50.888552Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.871486Z","title":"S., Kawarabayashi, K.-i., and Jegelka, S","venue":null,"work_id":"c2399e80-672e-44c2-8c06-638fe3bfda19","year":2021},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":87,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.246636Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:090698d30e42beed3393afb221740ef2f698a00e004af0ff56f850bd8c5b2290","observation_id":"c59a29d8-6a60-468d-97a2-20c58d040297","resolution":{"observed_at":"2026-08-07T22:12:50.876137Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1901.06523","last_updated":"2024-05-17T02:22:47Z","snapshot_observed_at":"2026-08-02T09:28:33.784961Z","submitted_at":"2019-01-19T13:37:39Z","title":"Frequency Principle: Fourier Analysis Sheds Light on Deep Neural Networks","version":7},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1901.06523","snapshot_observed_at":"2026-08-07T22:12:50.250323Z","title":"J., Zhang, Y., Luo, T., Xiao, Y., and Ma, Z","venue":null,"work_id":null,"year":1901},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":88,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.250323Z"},"links":{"cited_paper":"/paper/1901.06523","citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:dbc700b56fff51f9d3785f0c726155f714428f56b981865fe3192af12869bb50","observation_id":"2ab0add8-681f-4078-b1fb-a6369f5eaa73","resolution":{"observed_at":"2026-08-07T22:12:50.250323Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.860337Z","title":"and Paul, L","venue":null,"work_id":"f5214b6b-1500-41a8-8c2b-7e1fb838af7b","year":2024},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":89,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.254517Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:046e4ba4225fa55c72534167b2fd20b4099480efb3fbd638c1ede83ada304e42","observation_id":"e3fb2fc2-51a9-439c-8ba8-3e21315440e1","resolution":{"observed_at":"2026-08-07T22:12:50.864271Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.847653Z","title":"Understanding deep learning requires rethinking generalization","venue":null,"work_id":"8e2cf55b-0e13-48d3-b875-58dd4fe3bd16","year":2017},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":90,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.258678Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:a88bcfade0dcd553524cf6ff1246217a6245903b239662a84ebb61ca46fe7e20","observation_id":"68a47a00-642e-4916-8b3d-2828699f6722","resolution":{"observed_at":"2026-08-07T22:12:50.852125Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.833776Z","title":"Feature contamination: Neural networks learn uncorrelated features and fail to generalize","venue":null,"work_id":"84bf7929-3cf3-4ab5-9933-9059653f05ef","year":2024},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":91,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.262692Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:20de138bfaa420176d1b0acc25f1728ef1b1ebf61f18700cf0ad2fe7697578d5","observation_id":"218c12be-49ec-44bf-afdb-c4b9012798e8","resolution":{"observed_at":"2026-08-07T22:12:50.838112Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.821686Z","title":"M ^3 PL : Identifying and exploiting view bias of prompt learning","venue":null,"work_id":"b16f851f-b165-49d7-8e62-3ac4d4a515aa","year":2024},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":92,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.267116Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:a629d221549a4df3b168e22e034f28a4844e87dca2e0a07b680ec21749fb1dcd","observation_id":"b46c0945-daa1-4e5b-aa86-fc2f4e17a12a","resolution":{"observed_at":"2026-08-07T22:12:50.825682Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.807457Z","title":"P., and Orbanz, P","venue":null,"work_id":"084ef4c1-31ed-4ee4-91ec-26979675fc80","year":2019},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":93,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.270904Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:f1be867497365bf15c13352df5579ef77ced1ecacc22ff6d6d28dead1c0f809a","observation_id":"19e111a0-3695-4b22-b212-12ced909c0d6","resolution":{"observed_at":"2026-08-07T22:12:50.812099Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.794553Z","title":"S., Sharma, Y., Schneider, S., Bethge, M., and Brendel, W","venue":null,"work_id":"4a84b367-c118-4fc8-b578-3d54375a7f83","year":2021},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":94,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.274433Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:8e278d98cd2f3943428c3f5c9c268786e71d8b2f1a8a56f93834c642dfdd12d7","observation_id":"cf5c1cf8-9636-4845-82b3-6c64265e3d2b","resolution":{"observed_at":"2026-08-07T22:12:50.798854Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T22:12:50.781590Z","title":"Neural networks fail to learn periodic functions and how to fix it","venue":null,"work_id":"835b6700-2fc7-4098-982d-d17544b68647","year":2020},"citing_paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?","version":5},"reference_index":95,"source":"arxiv_source","source_observed_at":"2026-08-07T22:12:50.278134Z"},"links":{"citing_paper":"/paper/2502.09297"},"observation_digest":"sha256:8c670b1c139002088496c0777fe5d60c50c3565da7d4b5055f1081ad3db2af54","observation_id":"3a9d48a9-e311-42a9-a831-df8c758a4a16","resolution":{"observed_at":"2026-08-07T22:12:50.785901Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2502.09297","last_updated":"2025-09-09T07:21:48Z","latest_version":5,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-07T21:57:21.116942Z","submitted_at":"2025-02-13T13:11:54Z","title":"When Do Neural Networks Learn World Models?"},"reference_resolution":{"displayed":95,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":42,"verified_exact":2,"verified_fuzzy":51},"total_outbound_references":95},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 95 of 95 outbound references and 1 inbound Pith citation observation for arXiv:2502.09297."}