{"as_of":"2026-08-10T18:06:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:4af635a046b67ce576683a25a2b8b7f3c7258b506761548efa96b4a68a59c133","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":36,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":36,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":36,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":36,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-08T16:31:36.516166Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":60,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2010.14701","last_updated":"2020-11-06T04:16:36Z","snapshot_observed_at":"2026-07-06T10:09:17.078776Z","submitted_at":"2020-10-28T02:17:24Z","title":"Scaling Laws for Autoregressive Generative Modeling","version":2},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-05-13T07:49:43.711653Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2010.14701"},"observation_digest":"sha256:faaf915bf93d659b3bbc43deee1d067982f26ce63fc0a98a2823f51eaa4f0bda","observation_id":"aa995168-f597-4e21-9019-e9e4d2652a75","resolution":{"observed_at":"2026-05-13T07:49:43.819446Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2102.01293","last_updated":"2021-02-02T04:07:38Z","snapshot_observed_at":"2026-08-01T22:46:19.170916Z","submitted_at":"2021-02-02T04:07:38Z","title":"Scaling Laws for Transfer","version":1},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-05-18T00:58:13.116663Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2102.01293"},"observation_digest":"sha256:1481ef48ec67feaae4631381e23fee71ee1f3f237d7e9425708230c42a22c492","observation_id":"19ece7e6-6614-46d0-b7da-52a487f6eaa4","resolution":{"observed_at":"2026-05-18T00:58:13.693119Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2112.00861","last_updated":"2021-12-09T21:40:22Z","snapshot_observed_at":"2026-08-10T12:03:10.373653Z","submitted_at":"2021-12-01T22:24:34Z","title":"A General Language Assistant as a Laboratory for Alignment","version":3},"reference_index":94,"source":"arxiv_source","source_observed_at":"2026-05-11T14:22:57.925354Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2112.00861"},"observation_digest":"sha256:498f0e25dd22eda5f8a78592321ceb2df0eedf11dc2e8bc5e5eedbefe0b44951","observation_id":"617e79ba-f6d4-4378-84b8-de100a8cfc73","resolution":{"observed_at":"2026-05-11T14:22:58.981470Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2207.05221","last_updated":"2022-11-21T16:38:35Z","snapshot_observed_at":"2026-08-06T08:34:11.887259Z","submitted_at":"2022-07-11T22:59:39Z","title":"Language Models (Mostly) Know What They Know","version":4},"reference_index":152,"source":"arxiv_source","source_observed_at":"2026-05-10T15:42:47.274448Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2207.05221"},"observation_digest":"sha256:053e23ae2dfa3da89cb177f4751585b17758bf71778e01f063ec2709d6a0a11a","observation_id":"6261559f-2e0e-4b22-b45a-4d8df32b9d36","resolution":{"observed_at":"2026-05-10T15:42:47.810297Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2501.02378","last_updated":"2026-04-15T05:10:03Z","snapshot_observed_at":"2026-07-31T16:53:57.321287Z","submitted_at":"2025-01-04T20:49:20Z","title":"A ghost mechanism: An analytical model of abrupt learning in recurrent networks","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-05-23T06:31:21.006691Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2501.02378"},"observation_digest":"sha256:15f3a4ebb513eb0c2750dc5f2bc8b7b1cbf068b7f5bbafb3312a5610389a9015","observation_id":"a15c71d2-eff2-4c84-a603-85886ace9913","resolution":{"observed_at":"2026-05-23T06:32:38.851832Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-08T16:31:36.516166Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":null,"work_id":null,"year":2003},"citing_paper":{"arxiv_id":"2502.06192","last_updated":"2025-05-19T14:51:05Z","snapshot_observed_at":"2026-08-10T04:02:22.336385Z","submitted_at":"2025-02-10T06:48:04Z","title":"Right Time to Learn:Promoting Generalization via Bio-inspired Spacing Effect in Knowledge Distillation","version":2},"reference_index":2019,"source":"pdf_text","source_observed_at":"2026-08-08T16:31:36.516166Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2502.06192"},"observation_digest":"sha256:09aaf474f3b78ed9a90206baa656926f6837ad5d56b3558f24e52780c2af9f5e","observation_id":"2b72b369-736d-41cd-b45f-90aec27f8d62","resolution":{"observed_at":"2026-08-08T16:31:36.516166Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-07T10:42:54.707684Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":null,"work_id":null,"year":2003},"citing_paper":{"arxiv_id":"2506.04805","last_updated":"2026-05-25T11:37:32Z","snapshot_observed_at":"2026-08-07T10:30:20.168640Z","submitted_at":"2025-06-05T09:31:41Z","title":"Adaptive Preconditioners Trigger Loss Spikes in Adam","version":2},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-07T10:42:54.707684Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2506.04805"},"observation_digest":"sha256:3a7653cd3380a6e36452d72afe99667855aea2364f32db533a10b58d50c8b653","observation_id":"9f45d863-9706-4ebf-8ea7-1d696367aee8","resolution":{"observed_at":"2026-08-07T10:42:54.707684Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2603.10079","last_updated":"2026-04-13T13:15:13Z","snapshot_observed_at":"2026-08-02T23:21:16.952968Z","submitted_at":"2026-03-10T09:27:17Z","title":"Large Spikes in Stochastic Gradient Descent: A Large-Deviations View","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-15T13:30:04.366752Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2603.10079"},"observation_digest":"sha256:f54312940aab863db2ac9ffba9a7d34893ce35ca2f7cefcafc608ac99947be60","observation_id":"2d4679b0-9100-4b52-8d3a-6bd1049139c2","resolution":{"observed_at":"2026-05-15T13:30:51.043994Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2604.13627","last_updated":"2026-08-03T12:13:35Z","snapshot_observed_at":"2026-08-06T23:31:01.769852Z","submitted_at":"2026-04-15T08:53:42Z","title":"(How) Learning Rates Regulate Catastrophic Overtraining","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-10T13:40:19.845409Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2604.13627"},"observation_digest":"sha256:6d195a2b5dde0dc58ef31a5e511fac109605382cb595e17f6fa44fe3beb39b7e","observation_id":"6f90fd1c-8a12-40b0-bdc5-5f8197777bfc","resolution":{"observed_at":"2026-05-10T13:40:26.497872Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-04T05:30:35.676773Z","title":"The large learning rate phase of deep learning: the catapult mechanism.arXiv preprint arXiv:2003.02218,","venue":null,"work_id":null,"year":2003},"citing_paper":{"arxiv_id":"2604.13627","last_updated":"2026-08-03T12:13:35Z","snapshot_observed_at":"2026-08-06T23:31:01.769852Z","submitted_at":"2026-04-15T08:53:42Z","title":"(How) Learning Rates Regulate Catastrophic Overtraining","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-04T05:30:35.676773Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2604.13627"},"observation_digest":"sha256:445f7a18e1716c846e5d8793679174857b2c6f3aee819054f48a9e163c4531c1","observation_id":"758cdf53-4afe-46cd-94f4-7c0e3a9d543b","resolution":{"observed_at":"2026-08-04T05:30:35.676773Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2604.14669","last_updated":"2026-07-01T04:37:57Z","snapshot_observed_at":"2026-08-02T03:19:12.190552Z","submitted_at":"2026-04-16T06:23:18Z","title":"Zeroth-Order Optimization at the Edge of Stability","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-10T11:18:30.486247Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2604.14669"},"observation_digest":"sha256:ffde23983df532aef29864149bae927ca9d099eebb97142e93a3a108ef03743c","observation_id":"abbf4d66-e621-4a0d-b426-2cab6f18c99e","resolution":{"observed_at":"2026-05-10T11:20:10.144195Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-07-12T20:06:39.488359Z","title":"The large learning rate phase of deep learning: the catapult mechanism.arXiv preprint arXiv:2003.02218,","venue":null,"work_id":null,"year":2003},"citing_paper":{"arxiv_id":"2604.14669","last_updated":"2026-07-01T04:37:57Z","snapshot_observed_at":"2026-08-02T03:19:12.190552Z","submitted_at":"2026-04-16T06:23:18Z","title":"Zeroth-Order Optimization at the Edge of Stability","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-07-12T20:06:39.488359Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2604.14669"},"observation_digest":"sha256:d6166fa031864375f8f70abe88e094ceb7825d1962387fc657d44b5253f6b9d1","observation_id":"f1f63203-12e3-4bab-a3b4-d460b16298df","resolution":{"observed_at":"2026-07-12T20:06:39.488359Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2604.20446","last_updated":"2026-04-22T11:08:49Z","snapshot_observed_at":"2026-07-06T23:06:55.880484Z","submitted_at":"2026-04-22T11:08:49Z","title":"The Origin of Edge of Stability","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-05-10T01:24:55.474104Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2604.20446"},"observation_digest":"sha256:8ddfab4c00d7cc624312fe1faf291e01c82708c91dff3ca72ba3f5de4a537bb3","observation_id":"9f1484d4-1a70-4671-b587-2d69ce24461c","resolution":{"observed_at":"2026-05-11T13:36:07.948407Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2604.21016","last_updated":"2026-04-22T19:02:52Z","snapshot_observed_at":"2026-08-08T23:57:01.801387Z","submitted_at":"2026-04-22T19:02:52Z","title":"SGD at the Edge of Stability: The Stochastic Sharpness Gap","version":1},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-05-10T00:59:39.800662Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2604.21016"},"observation_digest":"sha256:4f827502976b148d0df5671acb537bb1bcc199851083fe04ed0b60320e6b86e0","observation_id":"9602fa3f-807d-4725-8944-0d29c099e53f","resolution":{"observed_at":"2026-05-10T00:59:49.391917Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2604.21691","last_updated":"2026-04-23T13:58:12Z","snapshot_observed_at":"2026-08-02T12:48:50.592713Z","submitted_at":"2026-04-23T13:58:12Z","title":"There Will Be a Scientific Theory of Deep Learning","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-05-09T20:11:17.616190Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2604.21691"},"observation_digest":"sha256:10eced8045c9a836e7139554e2eae86324abf9e5eba0358b33be17de4c0e4cd6","observation_id":"82f3b924-b8a0-420e-8775-d27a2173f4d6","resolution":{"observed_at":"2026-05-11T15:21:08.996891Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2605.02968","last_updated":"2026-05-03T12:21:14Z","snapshot_observed_at":"2026-07-06T23:15:55.848885Z","submitted_at":"2026-05-03T12:21:14Z","title":"Finite-Size Gradient Transport in Large Language Model Pretraining: From Cascade Size to Intensive Transport Efficiency","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-10T14:47:47.383941Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2605.02968"},"observation_digest":"sha256:4efb802dabf40304fe66a894a8e20a2014f3be255ba5fe72b2dc55200bb13197","observation_id":"fa3bea90-131e-48d7-94a2-4a0d1c1a8b4d","resolution":{"observed_at":"2026-05-10T15:00:31.040454Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2605.04054","last_updated":"2026-04-10T02:40:51Z","snapshot_observed_at":"2026-07-12T23:17:30.936793Z","submitted_at":"2026-04-10T02:40:51Z","title":"Endogenous Regime Switching Driven by Scalar-Irreducible Learning Dynamics","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-10T17:34:29.652388Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2605.04054"},"observation_digest":"sha256:73586664b06aef1e217c19087fdaec565d0d372046e92e8a84a9f7495ba1e9a7","observation_id":"f0c3f25a-a64b-4ad5-9fd7-6b6673039343","resolution":{"observed_at":"2026-05-11T06:36:01.856918Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2605.06821","last_updated":"2026-05-07T18:21:59Z","snapshot_observed_at":"2026-07-06T23:19:15.382283Z","submitted_at":"2026-05-07T18:21:59Z","title":"A Rod Flow Model for Adam at the Edge of Stability","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-05-11T00:54:24.615112Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2605.06821"},"observation_digest":"sha256:6f93317d1d3dbd08a5b43ed33fcc922db9420e6b01c8469e272f05ab193bdd4e","observation_id":"6bbaf906-a3d7-44ce-9fde-e2b624d0f59e","resolution":{"observed_at":"2026-05-11T05:00:56.647680Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2605.10468","last_updated":"2026-05-11T12:34:20Z","snapshot_observed_at":"2026-08-08T10:39:37.475240Z","submitted_at":"2026-05-11T12:34:20Z","title":"Can Muon Fine-tune Adam-Pretrained Models?","version":1},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-05-12T03:53:11.469583Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2605.10468"},"observation_digest":"sha256:1b4f18174dcf2ef82e6a97c8c0cf99f01057ed8e0e9780ef3d0ffd7b94af24e1","observation_id":"22c671cb-62d2-45e7-a830-c0bd1247e408","resolution":{"observed_at":"2026-05-12T06:51:28.677058Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2605.14200","last_updated":"2026-05-13T23:32:00Z","snapshot_observed_at":"2026-07-06T23:25:39.415637Z","submitted_at":"2026-05-13T23:32:00Z","title":"How to Scale Mixture-of-Experts: From muP to the Maximally Scale-Stable Parameterization","version":1},"reference_index":90,"source":"arxiv_source","source_observed_at":"2026-05-15T04:45:20.091598Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2605.14200"},"observation_digest":"sha256:281fe0c1c4a97b339f3b57c890f83181f60e080153fd4ca334ab765d10e06c2f","observation_id":"dcff6084-c81d-47e3-a253-5b8e87ae07eb","resolution":{"observed_at":"2026-05-15T04:49:44.826858Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2605.20005","last_updated":"2026-05-19T15:36:52Z","snapshot_observed_at":"2026-07-06T23:30:39.512029Z","submitted_at":"2026-05-19T15:36:52Z","title":"Fine-Tuning Without Forgetting via Loss-Adaptive Learning Rates","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-20T07:14:59.396900Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2605.20005"},"observation_digest":"sha256:100e1611269f6d10f56611f703776cc7cc82c86df37315a0a11207b643a33a1d","observation_id":"49dd278e-4c23-412f-8a62-8c5246844818","resolution":{"observed_at":"2026-05-20T07:18:07.041459Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2605.21292","last_updated":"2026-05-20T15:25:13Z","snapshot_observed_at":"2026-08-01T20:42:28.027915Z","submitted_at":"2026-05-20T15:25:13Z","title":"Large-Step Training Dynamics of a Two-Factor Linear Transformer Model","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-05-21T03:50:43.744189Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2605.21292"},"observation_digest":"sha256:7740d130043872e12bd1b7df9b38b877c7dc1ccf682ca3b96a43550b02df7791","observation_id":"d622889e-88bb-4c7b-a0d8-efc3a61dd255","resolution":{"observed_at":"2026-05-21T03:53:56.329810Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2606.00293","last_updated":"2026-05-29T19:24:38Z","snapshot_observed_at":"2026-08-02T12:57:03.676117Z","submitted_at":"2026-05-29T19:24:38Z","title":"Accurate Large-sample Uncertainty Quantification using Stochastic Gradient Markov Chain Monte Carlo","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-06-28T22:47:17.603350Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2606.00293"},"observation_digest":"sha256:bbf437c0a95a0811bab63be48a414ac1dbb7ea4de59b6c14ae396df7e62dba82","observation_id":"a38262b2-cb62-44b6-a9ba-85ffa08037ec","resolution":{"observed_at":"2026-07-01T19:16:01.182666Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2606.04212","last_updated":"2026-06-08T18:25:59Z","snapshot_observed_at":"2026-07-06T23:44:23.633099Z","submitted_at":"2026-06-02T20:58:40Z","title":"Edge of Stability Selectively Shapes Learning Across the Data Distribution","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-06-28T10:27:39.029895Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2606.04212"},"observation_digest":"sha256:cc775e3ffe9a0a19d1f94fa77925c9d24ca35d89cb35db8c2ce3351b72c899df","observation_id":"aea36370-95e3-4e8b-a880-45dff5b07730","resolution":{"observed_at":"2026-07-02T02:56:29.501282Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2606.05219","last_updated":"2026-05-29T13:32:18Z","snapshot_observed_at":"2026-08-08T23:43:07.652820Z","submitted_at":"2026-05-29T13:32:18Z","title":"Gradient Descent with Large Step Size Restores Symmetry in Deep Linear Networks with Multi-Pathway","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-06-28T23:15:59.394256Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2606.05219"},"observation_digest":"sha256:9c0e61573596b49a8a658e4dcd3f0a099b9ae94fe1e4016031adbe31f7d5a875","observation_id":"5c7b22c6-ee91-4a3a-b8ca-083ab82e857a","resolution":{"observed_at":"2026-06-28T23:22:46.406450Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2606.18080","last_updated":"2026-06-16T15:45:47Z","snapshot_observed_at":"2026-08-06T19:00:47.837914Z","submitted_at":"2026-06-16T15:45:47Z","title":"Edge Flow: A Tractable and Predictive Continuous-Time Model for Gradient Descent at the Edge of Stability","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-06-27T01:17:37.449377Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2606.18080"},"observation_digest":"sha256:f4b5e47853f8618ac4d0c54aa36d8c841c7b165ce172e9428175d5c2ad333595","observation_id":"3f38c97f-f7c8-48c7-bcb0-6ee12b2e4c95","resolution":{"observed_at":"2026-07-03T20:28:55.905813Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2606.22053","last_updated":"2026-06-20T14:10:49Z","snapshot_observed_at":"2026-08-08T00:11:02.573395Z","submitted_at":"2026-06-20T14:10:49Z","title":"Gradient-Descent Steps to Success over Mean Accuracy: A Paradigm Shift for ML","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-26T12:37:25.115951Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2606.22053"},"observation_digest":"sha256:f97bacf44c35fa9f76115ccc11c6a9a0f9049b85c020e833d55d40df01651170","observation_id":"af75fc6f-3063-471f-97e7-19aaa4a2d3d2","resolution":{"observed_at":"2026-07-04T07:49:39.677056Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":"2003.02218","doi":"10.48550/arxiv.2003.02218","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The large learning rate phase of deep learning: the catapult mechanism","venue":"arXiv (Cornell University)","work_id":"fdbcc7a5-2170-42dd-8424-9e228495e447","year":2003},"citing_paper":{"arxiv_id":"2606.23364","last_updated":"2026-06-22T14:00:26Z","snapshot_observed_at":"2026-08-02T17:37:15.368213Z","submitted_at":"2026-06-22T14:00:26Z","title":"Convergence of Gradient Descent for General Neural Network Architectures Beyond the NTK Regime","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-06-26T08:53:46.285233Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2606.23364"},"observation_digest":"sha256:3f1a87a6e10633aafa56b8d60fdf995fe6cf74b20b688fe5b8a4e4230d3786c5","observation_id":"147cf3b2-b3b7-46b3-a6fb-387931b0fe8e","resolution":{"observed_at":"2026-07-04T10:29:44.290800Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T15:49:36.486557+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-07-12T01:11:50.444218Z","title":"arXiv preprint arXiv:2003.02218 , year=","venue":null,"work_id":null,"year":2003},"citing_paper":{"arxiv_id":"2607.03613","last_updated":"2026-07-03T22:03:13Z","snapshot_observed_at":"2026-08-10T02:28:20.618211Z","submitted_at":"2026-07-03T22:03:13Z","title":"Implicit Bias of SGD in Multivariate ReLU Networks: Effective Width Collapse","version":1},"reference_index":189,"source":"arxiv_source","source_observed_at":"2026-07-12T01:11:50.444218Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2607.03613"},"observation_digest":"sha256:8ef906988155f1e5e38a479d83db2f1c8da5783588ead4e7dd1b885311fc5382","observation_id":"f5590f2d-b115-4234-8a37-eca57540006c","resolution":{"observed_at":"2026-07-12T01:11:50.444218Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-07-11T22:26:03.646701Z","title":"The Large Learning Rate Phase of Deep Learning: The Catapult Mechanism.arXiv preprint arXiv:2003.02218, 2020.https://doi.org/10.48550/arXiv.2003.02218","venue":null,"work_id":null,"year":2003},"citing_paper":{"arxiv_id":"2607.03998","last_updated":"2026-07-15T20:08:42Z","snapshot_observed_at":"2026-08-10T04:02:21.809335Z","submitted_at":"2026-07-04T19:49:01Z","title":"Directional Curvature from Armijo Backtracking: A Low-Cost Sharpness Probe and a Calibration-Free Learning-Rate Safeguard for Adam","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-07-11T22:26:03.646701Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2607.03998"},"observation_digest":"sha256:721e0c9a93f4a7417733c6f68d4d51a250bdf624e000aa190c9ca79fad1f2646","observation_id":"99af2c4f-03dd-481a-b5fc-2eeb432f6685","resolution":{"observed_at":"2026-07-11T22:26:03.646701Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-07-14T16:26:22.217462Z","title":"The Large Learning Rate Phase of Deep Learning: The Catapult Mechanism.arXiv preprint arXiv:2003.02218, 2020.https://doi.org/10.48550/arXiv.2003.02218","venue":null,"work_id":null,"year":2003},"citing_paper":{"arxiv_id":"2607.03998","last_updated":"2026-07-15T20:08:42Z","snapshot_observed_at":"2026-08-10T04:02:21.809335Z","submitted_at":"2026-07-04T19:49:01Z","title":"Directional Curvature from Armijo Backtracking: A Low-Cost Sharpness Probe and a Calibration-Free Learning-Rate Safeguard for Adam","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-07-14T16:26:22.217462Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2607.03998"},"observation_digest":"sha256:0f36d47d677986494ac9914f7dfa542285828036010987b508e3fb7b618c33b0","observation_id":"69446ef4-97dc-40ba-8097-aa60a77e02ff","resolution":{"observed_at":"2026-07-14T16:26:22.217462Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-02T08:48:43.532111Z","title":"The Large Learning Rate Phase of Deep Learning: The Catapult Mechanism.arXiv preprint arXiv:2003.02218, 2020.https://doi.org/10.48550/arXiv.2003.02218","venue":null,"work_id":null,"year":2003},"citing_paper":{"arxiv_id":"2607.03998","last_updated":"2026-07-15T20:08:42Z","snapshot_observed_at":"2026-08-10T04:02:21.809335Z","submitted_at":"2026-07-04T19:49:01Z","title":"Directional Curvature from Armijo Backtracking: A Low-Cost Sharpness Probe and a Calibration-Free Learning-Rate Safeguard for Adam","version":3},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-02T08:48:43.532111Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2607.03998"},"observation_digest":"sha256:a238d9629767c7e9062557a298879b708abf41d64acd1c750817e23fadf49a11","observation_id":"bcb3f7be-9112-4054-a632-51be920dd956","resolution":{"observed_at":"2026-08-02T08:48:43.532111Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-07-11T10:24:00.719150Z","title":"arXiv preprint arXiv:2003.02218 , year =","venue":null,"work_id":null,"year":2003},"citing_paper":{"arxiv_id":"2607.04993","last_updated":"2026-07-06T12:28:24Z","snapshot_observed_at":"2026-08-10T17:20:08.037438Z","submitted_at":"2026-07-06T12:28:24Z","title":"The Map Behind the Flow: Finite-Step Gradient Descent as a Dynamical System","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-07-11T10:24:00.719150Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2607.04993"},"observation_digest":"sha256:e89fb338fe4e2cad454fc011d9ba88e0dc71af56ef77a5a7d8297606fd13a789","observation_id":"8ed05444-e79c-417c-b07c-cc62055f153e","resolution":{"observed_at":"2026-07-11T10:24:00.719150Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-01T08:50:41.724748Z","title":"arXiv preprint arXiv:2003.02218 , year=","venue":null,"work_id":null,"year":2003},"citing_paper":{"arxiv_id":"2607.21005","last_updated":"2026-07-23T07:39:03Z","snapshot_observed_at":"2026-08-06T16:38:33.932762Z","submitted_at":"2026-07-23T07:39:03Z","title":"Weight-norm Criticality: A Mechanism for Loss Spikes Induced by the Normalization and Weight Decay","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-01T08:50:41.724748Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2607.21005"},"observation_digest":"sha256:235f01c33898613a752f7ae16ca5a287210d847911c66f773a1222759df3f4ae","observation_id":"71e7f561-54f2-4a81-ad47-b6e249b99043","resolution":{"observed_at":"2026-08-01T08:50:41.724748Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-01T06:58:09.772532Z","title":"The large learning rate phase of deep learning: the catapult mechanism.arXiv preprint arXiv:2003.02218,","venue":null,"work_id":null,"year":2003},"citing_paper":{"arxiv_id":"2607.21716","last_updated":"2026-07-23T18:02:03Z","snapshot_observed_at":"2026-08-09T14:14:32.562223Z","submitted_at":"2026-07-23T18:02:03Z","title":"A Defense of the Quadratic Model","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-01T06:58:09.772532Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2607.21716"},"observation_digest":"sha256:3d068dd041b0d3ffe25dbd6625f8d2ed3617080a03336e3ee3c7db80fe3662f6","observation_id":"ce33ad15-ee68-4472-b792-648ee5d9b15b","resolution":{"observed_at":"2026-08-01T06:58:09.772532Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2003.02218","snapshot_observed_at":"2026-08-06T00:43:14.981584Z","title":"The large learning rate phase of deep learning: The catapult mechanism, 2020","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2608.01032","last_updated":"2026-08-02T06:27:55Z","snapshot_observed_at":"2026-08-09T06:27:10.466682Z","submitted_at":"2026-08-02T06:27:55Z","title":"The Fourth Quadrant: A Stylized View of Benign Misfitting","version":1},"reference_index":118,"source":"pdf_text","source_observed_at":"2026-08-06T00:43:14.981584Z"},"links":{"cited_paper":"/paper/2003.02218","citing_paper":"/paper/2608.01032"},"observation_digest":"sha256:18bda9e86450ce0feb8b42656d663a43d4dce0e0d33dbf50056977ddd7a09d68","observation_id":"54c77455-4e96-4490-aced-f7b75d138663","resolution":{"observed_at":"2026-08-06T00:43:14.981584Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2003.02218/citation-record","integrity":"/paper/2003.02218/integrity","json":"/paper/2003.02218/citation-record.json","paper":"/paper/2003.02218"},"outbound":[],"paper":{"arxiv_id":"2003.02218","last_updated":"2020-03-04T17:52:48Z","latest_version":1,"primary_category":"stat.ML","snapshot_observed_at":"2026-08-08T21:43:40.840815Z","submitted_at":"2020-03-04T17:52:48Z","title":"The large learning rate phase of deep learning: the catapult mechanism"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 36 inbound Pith citation observations for arXiv:2003.02218."}