{"as_of":"2026-08-18T18:45:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:c7371ce8fbb41553245eb0a739134966acddd4a8aba43ddc3fe097370b326cad","coverage":[{"denominator":38,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":38,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T20:13:29.997803Z","state":"measured"},{"denominator":38,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":38,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-18T06:34:40.430872+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2505.13900/citation-record","integrity":"/paper/2505.13900/integrity","json":"/paper/2505.13900/citation-record.json","paper":"/paper/2505.13900"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"1711.08856","last_updated":"2019-02-25T11:08:56Z","snapshot_observed_at":"2026-08-17T03:04:27.775014Z","submitted_at":"2017-11-24T01:58:54Z","title":"Critical Learning Periods in Deep Neural Networks","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1711.08856","snapshot_observed_at":"2026-08-15T20:13:29.679072Z","title":"Critical learning periods in deep neural networks","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.679072Z"},"links":{"cited_paper":"/paper/1711.08856","citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:240f277a93767d58372448f48a263712e7b8cf4d0cf7b9ece588942dfbeb66e2","observation_id":"7de84ed3-6723-4b3e-8b1d-43f7b343714d","resolution":{"observed_at":"2026-08-15T20:13:29.679072Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:31.921276Z","title":"Ainsworth, Jonathan Hayase, and Siddhartha S","venue":null,"work_id":"c4a6e9a4-c6c7-4d4e-9871-e96ceeb3495e","year":2023},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.706925Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:9907afbda995bb1dc717c72361a4db48247106a9802a021fcb6344f4acaf5e4e","observation_id":"48e33b7d-6855-4f01-83d1-1a0ff7438316","resolution":{"observed_at":"2026-08-15T20:13:31.927424Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:31.759523Z","title":"A convergence theory for deep learning via over- parameterization","venue":null,"work_id":"1bb3b90b-7d38-495e-abec-e99383404eed","year":2019},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.712005Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:20a2930362d3d84fabd548c2f3d54e25a7be6cd0bed76761d218994f3a8a2d90","observation_id":"7e6b8784-264c-4951-988e-0a961ded2d71","resolution":{"observed_at":"2026-08-15T20:13:31.846345Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:31.740583Z","title":"On exact computation with an infinitely wide neural net","venue":null,"work_id":"24e46560-a10f-42f2-a929-a8721c7d6ff4","year":2019},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.717644Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:3743bb58b464e37acbeb0498438e64c3c1dd658144d1cecac8b0914860c4bcbd","observation_id":"b21adc22-f271-4819-9ea7-c3893b9e1465","resolution":{"observed_at":"2026-08-15T20:13:31.746850Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:31.723362Z","title":"A survey on deep learning applied to medical images: from simple artificial neural networks to generative models.Neural Computing and Applications, 35(3):2291–2323, 2023","venue":null,"work_id":"3fd462be-dd17-4964-ad10-81b57c828a62","year":2023},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.722341Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:8b8a83a8e45371c516398d81c352b0bc13d1b3d8e8f0e79b23e67c5ad7618aec","observation_id":"9d3ef832-a61a-43fb-b6b6-0bc56548dd49","resolution":{"observed_at":"2026-08-15T20:13:31.729169Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:31.707720Z","title":"On lazy training in differentiable programming","venue":null,"work_id":"aa31b8e1-6e9e-41f1-9d59-fc0f099e71da","year":2019},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.727088Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:7f1aecdb09fd45a35378b95e3347a47b672b0f4e91352619a349637d0b67f549","observation_id":"c42b6945-b3a4-43e1-a93a-46d7d09e3c3a","resolution":{"observed_at":"2026-08-15T20:13:31.712861Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2103.00065","last_updated":"2022-11-23T18:09:57Z","snapshot_observed_at":"2026-08-16T18:42:37.696014Z","submitted_at":"2021-02-26T22:08:19Z","title":"Gradient Descent on Neural Networks Typically Occurs at the Edge of Stability","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2103.00065","snapshot_observed_at":"2026-08-15T20:13:29.732360Z","title":"Gradient descent on neural networks typically occurs at the edge of stability.arXiv preprint arXiv:2103.00065, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.732360Z"},"links":{"cited_paper":"/paper/2103.00065","citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:37291f8e1d51870ad37b23516649551a466f1837b07897331ac6406dd06bb4d5","observation_id":"38b682b7-7f87-4130-9190-42710c9def53","resolution":{"observed_at":"2026-08-15T20:13:29.732360Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:31.591868Z","title":null,"venue":null,"work_id":"7a678c9e-c2de-4c22-84de-46e2b18f2dc1","year":2023},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.738029Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:5dfa3d1d4c7feb460263fb7f3e9c3003a71700a836dad1aef6d5091bdf092f86","observation_id":"8953a0cf-5c37-4f25-9f2a-5e6ff50212d9","resolution":{"observed_at":"2026-08-15T20:13:31.696308Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:31.361234Z","title":"Gradient descent finds global minima of deep neural networks","venue":null,"work_id":"9876c803-4588-4494-ab2f-5ac6b5e3e234","year":2019},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.742645Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:ede345e320f8cef17f3f1959303ece891072c0524db83452b3bc01b6aa1b682c","observation_id":"34de2252-102f-465d-85bc-96155e2fba25","resolution":{"observed_at":"2026-08-15T20:13:31.494528Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:29.747693Z","title":"Du, Xiyu Zhai, Barnabas Poczos, and Aarti Singh","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.747693Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:56625f724f8f121c2593c702a45a5104c05fee103203f4602290ee58f924816e","observation_id":"6eb35913-093a-4500-b6ac-63e937934eb5","resolution":{"observed_at":"2026-08-15T20:13:29.747693Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:29.752369Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.752369Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:60de11425cad5386fd1c179133938ffb812b2161062d93eda9120d96948ed9cf","observation_id":"8c0e2b99-c62b-46aa-8b2c-b8a018eacb16","resolution":{"observed_at":"2026-08-15T20:13:29.752369Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:29.757180Z","title":"Linear mode connectivity and the lottery ticket hypothesis","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.757180Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:84ad194760ca53e372879e717041f7ab49f92ed269259475ed1f38ac241a9fb4","observation_id":"34b1e7c5-c93c-4568-b4ca-bac2e5a1de86","resolution":{"observed_at":"2026-08-15T20:13:29.757180Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:31.159257Z","title":"Schwab, and Ari S","venue":null,"work_id":"71fad06e-950d-4dbc-a7e7-83d0e362cbc5","year":2020},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.761859Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:0d33dc2afb330bced6cb284319e18e604dd6de0c931550afd0b31dbd6874643b","observation_id":"3f56ddbc-07be-410b-b785-1b7af9124770","resolution":{"observed_at":"2026-08-15T20:13:31.229786Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:29.766480Z","title":"An investigation into neural net optimization via hessian eigenvalue density","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.766480Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:25c8a9e219ea8d14af50daf986c44d8a56842e4f2d2ee7b8b5f2d53cd1c7d124","observation_id":"b50c45ab-e12c-47e1-b6c9-02e4ceecb103","resolution":{"observed_at":"2026-08-15T20:13:29.766480Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1812.04754","last_updated":"2018-12-12T00:36:17Z","snapshot_observed_at":"2026-08-14T17:44:58.723571Z","submitted_at":"2018-12-12T00:36:17Z","title":"Gradient Descent Happens in a Tiny Subspace","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1812.04754","snapshot_observed_at":"2026-08-15T20:13:29.771210Z","title":"Gradient descent happens in a tiny subspace.arXiv preprint arXiv:1812.04754, 2018","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.771210Z"},"links":{"cited_paper":"/paper/1812.04754","citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:c512379a2c202173b2633d76aa8240f25f272bb0d2bba94d5c0ee9bb3fd058e2","observation_id":"61b9e24e-fe07-48a8-8973-65a4ab906024","resolution":{"observed_at":"2026-08-15T20:13:29.771210Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:29.776453Z","title":"Deep residual learning for image recognition","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.776453Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:959424e967e968e9e7c323442ee5085e02f65218d1280bfe7546026215dbf50a","observation_id":"3ac6bff8-0044-4deb-af4a-4b081f6eb9b5","resolution":{"observed_at":"2026-08-15T20:13:29.776453Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:30.998317Z","title":"Neural tangent kernel: Convergence and generalization in neural networks","venue":null,"work_id":"462d9bde-7df1-4db8-9fc5-ca4911d09b34","year":2018},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.781510Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:966357ace473bd137c3d34cf4440da9367c37aeddc3c2708e2b84e452398a37c","observation_id":"15be1226-57b9-48bb-a542-bb0b4f2bdb77","resolution":{"observed_at":"2026-08-15T20:13:31.044548Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2002.09572","last_updated":"2020-02-21T22:55:51Z","snapshot_observed_at":"2026-08-13T23:16:37.866586Z","submitted_at":"2020-02-21T22:55:51Z","title":"The Break-Even Point on Optimization Trajectories of Deep Neural Networks","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2002.09572","snapshot_observed_at":"2026-08-15T20:13:29.786110Z","title":"The break-even point on optimization trajectories of deep neural networks.arXiv preprint arXiv:2002.09572, 2020","venue":null,"work_id":null,"year":2002},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.786110Z"},"links":{"cited_paper":"/paper/2002.09572","citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:8763f6f38c80947edf718d3b2faccbdf74e6415f82e4e12407d7154bd4e4c2d6","observation_id":"5f37249c-2dc9-42d2-bbd2-97e11019ac2a","resolution":{"observed_at":"2026-08-15T20:13:29.786110Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:30.982818Z","title":"Assessing generalization of SGD via disagreement","venue":null,"work_id":"b5aa1611-7936-46da-ab5a-e5df0966b149","year":2022},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.791873Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:235f1ca6164954799426150bfaf2797988820ad620a34ba0127086b9c5673822","observation_id":"688b3afa-7afd-4dcb-9f3e-3a9a89728fba","resolution":{"observed_at":"2026-08-15T20:13:30.987294Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:30.965328Z","title":"Wide neural networks of any depth evolve as linear models under gradient descent","venue":null,"work_id":"b726d1a5-c097-44a7-949a-be91fb78b403","year":2019},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.798150Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:95f103954606f76573e146bf03a1206d4230c39136c24dd5b363e88a85ed616d","observation_id":"950d4d1b-84a3-4791-9c0d-388bd69531be","resolution":{"observed_at":"2026-08-15T20:13:30.971993Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:30.856896Z","title":"Learning overparameterized neural networks via stochastic gradient descent on structured data","venue":null,"work_id":"a4b43a5a-2473-48e1-b979-dcb8a0f88680","year":2018},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.802958Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:f93b98c6886f2beeea6a47bbd2171156de9633ef8f9569a8a77c4c6034aabdac","observation_id":"df13fba8-fcfc-417e-bc83-06aeabefc0c4","resolution":{"observed_at":"2026-08-15T20:13:30.898297Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:30.789442Z","title":"What happens after SGD reaches zero loss? –a mathematical framework","venue":null,"work_id":"6bafa4ff-738d-4939-863e-b9977aada15b","year":2022},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.807699Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:041f71de82222df9a5afe0a1cad674fac11622bbcf5a60d01c8cf36ac3ed5a30","observation_id":"57864035-d4bf-4107-b8aa-be124642e3a7","resolution":{"observed_at":"2026-08-15T20:13:30.806664Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:30.772649Z","title":"Task arithmetic in the tangent space: Improved editing of pre-trained models","venue":null,"work_id":"09da5e0a-226f-4250-89ae-ed3d5ed4b3f0","year":2023},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.812599Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:305e7377a3e1cf6c82b07907363e67bfd8729250e7dee6638a8249b5ee857800","observation_id":"4aaa7dca-ce47-40be-a908-c86e90f94368","resolution":{"observed_at":"2026-08-15T20:13:30.778035Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:30.756173Z","title":"Deep learning-based weather prediction: a survey.Big Data Research, 23:100178, 2021","venue":null,"work_id":"4695364e-a6f7-4401-8d5f-d4f07beeeb09","year":2021},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.817843Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:d97dd0efbcb25e70f531e41f5fcc5c476d298d156a589fc29d6a7ae4214b1a87","observation_id":"9a8f9344-aaf0-48bf-8437-bbd304c96c3e","resolution":{"observed_at":"2026-08-15T20:13:30.761658Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1409.1556","last_updated":"2015-04-10T16:25:04Z","snapshot_observed_at":"2026-08-17T19:17:06.411141Z","submitted_at":"2014-09-04T19:48:04Z","title":"Very Deep Convolutional Networks for Large-Scale Image Recognition","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1409.1556","snapshot_observed_at":"2026-08-15T20:13:29.823505Z","title":"Very deep convolutional networks for large-scale image recogni- tion","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.823505Z"},"links":{"cited_paper":"/paper/1409.1556","citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:74b56a17124b60e21bae24ecdcded129059a2e9b66f2d9ba47e37c640f37f97d","observation_id":"47e81bda-07a5-4567-bbd6-c98b13dfb8a4","resolution":{"observed_at":"2026-08-15T20:13:29.823505Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:30.741118Z","title":"The directionality of optimization trajectories in neural networks","venue":null,"work_id":"f34dbab0-c69b-47be-bfc3-7b4ffcb3556c","year":null},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.828561Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:3f8b1a60ee18dae1a2eb4c7f970744feb206cb8c032f8fcd5478289c905f06cd","observation_id":"aa1e766a-4168-4371-b31a-0d395e49246e","resolution":{"observed_at":"2026-08-15T20:13:30.745789Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:30.668935Z","title":"A survey on statistical theory of deep learning: Approximation, training dynamics, and generative models.Annual Review of Statistics and Its Application, 12, 2024","venue":null,"work_id":"0b77d202-489d-44a9-a1ba-31024105edb8","year":2024},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.892646Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:6f1d091c41925bd82a2a8a713418776392738f76ebd5da701ad8ab694b74db7d","observation_id":"e895e113-11a2-4826-8939-15d878945eca","resolution":{"observed_at":"2026-08-15T20:13:30.714657Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:30.536164Z","title":"Deep reinforcement learning for robotics: A survey of real-world successes","venue":null,"work_id":"95665243-4148-4734-ac0e-45ab625ccd62","year":2025},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.953165Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:6b59237b2bd8f91b1a5dd71ab2ba285e74faaf17a179fac351a6afd5caf110b9","observation_id":"b53df735-f7d5-4d19-a196-2d22720fbe4f","resolution":{"observed_at":"2026-08-15T20:13:30.598999Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:30.519303Z","title":"Analyzing sharpness along gd trajectory: Progressive sharpening and edge of stability.Advances in Neural Information Processing Systems, 35:9983–9994, 2022","venue":null,"work_id":"d99b2b41-19d2-480a-a252-50e99749395b","year":2022},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.958446Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:ea3181bd285acaa2dd6ca8462e9f749a46a1f0160fbe662d36fe0e6acedc0566","observation_id":"a4311636-0f69-4d2e-94ee-99c02db9abb2","resolution":{"observed_at":"2026-08-15T20:13:30.524046Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:30.501650Z","title":"Model soups: aver- aging weights of multiple fine-tuned models improves accuracy without increasing inference time","venue":null,"work_id":"4da76941-447d-4542-b4fb-4af0e220a78d","year":2022},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.963171Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:957ccf3f09eec5d8ffc7732f9b409ea135372c0340dd290542ebc6c9407c5e67","observation_id":"603f0ea1-1e6c-42be-84d0-f4a962cd7324","resolution":{"observed_at":"2026-08-15T20:13:30.507224Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:30.426518Z","title":"Robust fine-tuning of zero-shot models","venue":null,"work_id":"9ac5a0da-d543-4e15-b065-dbf0a7cc55f6","year":2022},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.967761Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:e833d6296d43e611f41387111756e01c98d46b21bb2c889b2c7ac1574232bacd","observation_id":"4670cb38-136d-44f5-997e-1f5fe9c7dcde","resolution":{"observed_at":"2026-08-15T20:13:30.488597Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:30.287314Z","title":"How sgd selects the global minima in over-parameterized learning: A dynamical stability perspective","venue":null,"work_id":"2a03e97f-d1ca-4296-97be-cbb48e07a150","year":2018},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.972392Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:d986645f8f9583d73b765189306367633352baec93a66fe9efb9985847e47609","observation_id":"b35d1f7b-bccf-4022-9799-d1e85cd941d7","resolution":{"observed_at":"2026-08-15T20:13:30.354567Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1902.04760","last_updated":"2020-04-04T22:53:19Z","snapshot_observed_at":"2026-08-18T09:45:12.030203Z","submitted_at":"2019-02-13T06:09:18Z","title":"Scaling Limits of Wide Neural Networks with Weight Sharing: Gaussian Process Behavior, Gradient Independence, and Neural Tangent Kernel Derivation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1902.04760","snapshot_observed_at":"2026-08-15T20:13:29.977799Z","title":null,"venue":null,"work_id":null,"year":1902},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.977799Z"},"links":{"cited_paper":"/paper/1902.04760","citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:31e61aeba38fdfcb7b95687c35884e0fe138c468fc932194c63b4eed56d0f95a","observation_id":"2c5c2bfd-66fd-47a1-bdc5-76a3ab71567e","resolution":{"observed_at":"2026-08-15T20:13:29.977799Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:30.271698Z","title":"Swing-by dynamics in concept learning and compositional generalization","venue":null,"work_id":"3a763b97-0a37-420e-b2ac-b3a969b18a15","year":2025},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.982834Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:c14ea863ee5e91e98a1da6fe52142294a3d18d3607289e168b418156193e4aa2","observation_id":"44e1d77c-0e50-45a0-b3d1-5ae2a9bf10c5","resolution":{"observed_at":"2026-08-15T20:13:30.276528Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:29.988384Z","title":"Going beyond linear mode connectivity: The layerwise linear feature connectivity.Advances in neural information processing systems, 36:60853–60877, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.988384Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:979d4e47f4c03197b35f93d79603d44b0073326264287750659e8535a0947d44","observation_id":"47a22bdb-fa18-4cb6-b48b-2ea7b7a69f22","resolution":{"observed_at":"2026-08-15T20:13:29.988384Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:30.243706Z","title":"On the emergence of cross-task linearity in pretraining-finetuning paradigm","venue":null,"work_id":"a289785e-8cf3-40ab-8833-d32816af5c1c","year":2024},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.993072Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:e5ea5d8dd95bd9d8074760233f317309fadd14eef2eb0d951e70acf5837c0363","observation_id":"f1e2713f-3f23-40e6-8833-f8390927f994","resolution":{"observed_at":"2026-08-15T20:13:30.248172Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:30.149291Z","title":"Gradient descent optimizes over-parameterized deep relu networks.Machine learning, 109:467–492, 2020","venue":null,"work_id":"b2922c25-f78d-48ca-97a1-4ac490dab185","year":2020},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.997803Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:16d71c6f5a3c8162f3744f988f647ccb69be8cbaa636932bd8e8da19a01ed481","observation_id":"408770bd-781c-4fb3-9838-f1d5b8573072","resolution":{"observed_at":"2026-08-15T20:13:30.207822Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:13:30.725884Z","title":null,"venue":null,"work_id":"878a504b-55d4-4c05-a628-18da66fbaa46","year":null},"citing_paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-15T20:13:29.833306Z"},"links":{"citing_paper":"/paper/2505.13900"},"observation_digest":"sha256:4c6d324ab9bbdca3f508f3ea873e001db0e9e374b162488e54a839a986f80c60","observation_id":"caf29d74-1275-450b-980c-7bb2cbd59895","resolution":{"observed_at":"2026-08-15T20:13:30.730930Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2505.13900","last_updated":"2025-05-20T04:03:52Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-18T09:46:25.765165Z","submitted_at":"2025-05-20T04:03:52Z","title":"New Evidence of the Two-Phase Learning Dynamics of Neural Networks"},"reference_resolution":{"displayed":38,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":14,"verified_exact":0,"verified_fuzzy":24},"total_outbound_references":38},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"thesis":"As of 18 August 2026, this Paper Citation Record lists 38 of 38 outbound references and 0 inbound Pith citation observations for arXiv:2505.13900."}