{"as_of":"2026-08-07T10:27:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:0a3f8baa6212c65915bdf1371f84072d4daa1602471226a971466647e854762b","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":7,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":7,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":7,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":7,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T19:31:20.239832Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-02T16:07:09.448556Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2206.05794","last_updated":"2024-10-18T21:32:39Z","snapshot_observed_at":"2026-08-06T01:24:27.374782Z","submitted_at":"2022-06-12T17:06:35Z","title":"SGD and Weight Decay Secretly Minimize the Rank of Your Neural Network","version":7},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2206.05794","snapshot_observed_at":"2026-08-06T19:31:20.239832Z","title":"Galanti, Z","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.05644","last_updated":"2025-09-05T01:58:57Z","snapshot_observed_at":"2026-08-06T19:18:52.304504Z","submitted_at":"2025-07-08T03:52:48Z","title":"The Features at Convergence Theorem: a first-principles alternative to the Neural Feature Ansatz for how networks learn representations","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T19:31:20.239832Z"},"links":{"cited_paper":"/paper/2206.05794","citing_paper":"/paper/2507.05644"},"observation_digest":"sha256:f5643f2f00b1306e376f6a98f40e36817ad68accbbc106e8419cc6f8231e79e7","observation_id":"706d7da5-4244-4ea5-aa3d-85df53b06cf3","resolution":{"observed_at":"2026-08-06T19:31:20.239832Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2206.05794","last_updated":"2024-10-18T21:32:39Z","snapshot_observed_at":"2026-08-06T01:24:27.374782Z","submitted_at":"2022-06-12T17:06:35Z","title":"SGD and Weight Decay Secretly Minimize the Rank of Your Neural Network","version":7},"cited_work":{"arxiv_id":"2206.05794","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2206.05794","snapshot_observed_at":"2026-07-02T16:07:09.448556Z","title":"arXiv preprint arXiv:2206.05794 , year=","venue":null,"work_id":"4e72f082-dab4-4555-8e12-2ce3304f0433","year":2022},"citing_paper":{"arxiv_id":"2604.03473","last_updated":"2026-04-03T21:41:41Z","snapshot_observed_at":"2026-08-07T02:03:37.863182Z","submitted_at":"2026-04-03T21:41:41Z","title":"Evolutionary Search for Automated Design of Uncertainty Quantification Methods","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-13T19:17:56.912708Z"},"links":{"cited_paper":"/paper/2206.05794","citing_paper":"/paper/2604.03473"},"observation_digest":"sha256:8b64b5e3ab19f0708b6907aad92276aa86f5688c5f5407f8ab3b941fe8a86f8e","observation_id":"1fdfa8de-876d-4140-be48-13626623bd21","resolution":{"observed_at":"2026-05-13T19:18:09.217340Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.05794","last_updated":"2024-10-18T21:32:39Z","snapshot_observed_at":"2026-08-06T01:24:27.374782Z","submitted_at":"2022-06-12T17:06:35Z","title":"SGD and Weight Decay Secretly Minimize the Rank of Your Neural Network","version":7},"cited_work":{"arxiv_id":"2206.05794","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2206.05794","snapshot_observed_at":"2026-07-02T16:07:09.448556Z","title":"arXiv preprint arXiv:2206.05794 , year=","venue":null,"work_id":"4e72f082-dab4-4555-8e12-2ce3304f0433","year":2022},"citing_paper":{"arxiv_id":"2605.16622","last_updated":"2026-05-15T20:43:26Z","snapshot_observed_at":"2026-08-01T15:07:39.814716Z","submitted_at":"2026-05-15T20:43:26Z","title":"Does Weight Decay Enhance Training Stability?","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-20T19:49:01.351717Z"},"links":{"cited_paper":"/paper/2206.05794","citing_paper":"/paper/2605.16622"},"observation_digest":"sha256:d0b3787bfc0c20dc2950bd01188b4fd5e5e20566dcd747ab40fce6242f9e69d0","observation_id":"a0e1c6e9-3caf-40c9-ad18-a3986cf9d05a","resolution":{"observed_at":"2026-05-20T19:53:42.998253Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.05794","last_updated":"2024-10-18T21:32:39Z","snapshot_observed_at":"2026-08-06T01:24:27.374782Z","submitted_at":"2022-06-12T17:06:35Z","title":"SGD and Weight Decay Secretly Minimize the Rank of Your Neural Network","version":7},"cited_work":{"arxiv_id":"2206.05794","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2206.05794","snapshot_observed_at":"2026-07-02T16:07:09.448556Z","title":"arXiv preprint arXiv:2206.05794 , year=","venue":null,"work_id":"4e72f082-dab4-4555-8e12-2ce3304f0433","year":2022},"citing_paper":{"arxiv_id":"2605.20441","last_updated":"2026-05-19T19:48:40Z","snapshot_observed_at":"2026-07-06T23:31:03.878152Z","submitted_at":"2026-05-19T19:48:40Z","title":"Weight Decay Regimes in Grokking Transformers: Cheap Online Diagnostics","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-21T07:30:07.287938Z"},"links":{"cited_paper":"/paper/2206.05794","citing_paper":"/paper/2605.20441"},"observation_digest":"sha256:3707fd12c16557fb52f6d471c7ec595d90c0506648bc22913d5fd7134850e326","observation_id":"99b448d0-bb6a-4de8-9beb-de8a7e182b88","resolution":{"observed_at":"2026-05-21T07:34:02.899180Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.05794","last_updated":"2024-10-18T21:32:39Z","snapshot_observed_at":"2026-08-06T01:24:27.374782Z","submitted_at":"2022-06-12T17:06:35Z","title":"SGD and Weight Decay Secretly Minimize the Rank of Your Neural Network","version":7},"cited_work":{"arxiv_id":"2206.05794","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2206.05794","snapshot_observed_at":"2026-07-02T16:07:09.448556Z","title":"arXiv preprint arXiv:2206.05794 , year=","venue":null,"work_id":"4e72f082-dab4-4555-8e12-2ce3304f0433","year":2022},"citing_paper":{"arxiv_id":"2605.23087","last_updated":"2026-05-21T22:37:25Z","snapshot_observed_at":"2026-07-06T23:33:19.375527Z","submitted_at":"2026-05-21T22:37:25Z","title":"The Implicit Bias of Depth: From Neural Collapse to Softmax Codes","version":1},"reference_index":137,"source":"arxiv_source","source_observed_at":"2026-05-25T05:26:15.556205Z"},"links":{"cited_paper":"/paper/2206.05794","citing_paper":"/paper/2605.23087"},"observation_digest":"sha256:3262d344fd8e75d91c51c1b8418b3c96a73c9c76d7f37211e7ac6671717eba42","observation_id":"ec9025ee-74ed-47e5-9d30-0ac4460eb6e1","resolution":{"observed_at":"2026-05-25T05:26:38.933941Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.05794","last_updated":"2024-10-18T21:32:39Z","snapshot_observed_at":"2026-08-06T01:24:27.374782Z","submitted_at":"2022-06-12T17:06:35Z","title":"SGD and Weight Decay Secretly Minimize the Rank of Your Neural Network","version":7},"cited_work":{"arxiv_id":"2206.05794","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2206.05794","snapshot_observed_at":"2026-07-02T16:07:09.448556Z","title":"arXiv preprint arXiv:2206.05794 , year=","venue":null,"work_id":"4e72f082-dab4-4555-8e12-2ce3304f0433","year":2022},"citing_paper":{"arxiv_id":"2606.05863","last_updated":"2026-06-04T08:39:04Z","snapshot_observed_at":"2026-08-07T06:24:53.080556Z","submitted_at":"2026-06-04T08:39:04Z","title":"Deciphering Two Training Clocks in Grokking via Deep Linear Network Theory with Conditional ReLU Reduction","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-06-28T02:26:16.631418Z"},"links":{"cited_paper":"/paper/2206.05794","citing_paper":"/paper/2606.05863"},"observation_digest":"sha256:ac4b1369623fbd49c148963b2a4ea0369456c4368445ba6b8ddcdd4de5985ffc","observation_id":"8a2839c8-1691-49ff-8d3e-24e1994385e0","resolution":{"observed_at":"2026-07-02T12:06:56.089444Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.05794","last_updated":"2024-10-18T21:32:39Z","snapshot_observed_at":"2026-08-06T01:24:27.374782Z","submitted_at":"2022-06-12T17:06:35Z","title":"SGD and Weight Decay Secretly Minimize the Rank of Your Neural Network","version":7},"cited_work":{"arxiv_id":"2206.05794","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2206.05794","snapshot_observed_at":"2026-07-02T16:07:09.448556Z","title":"arXiv preprint arXiv:2206.05794 , year=","venue":null,"work_id":"4e72f082-dab4-4555-8e12-2ce3304f0433","year":2022},"citing_paper":{"arxiv_id":"2606.07404","last_updated":"2026-06-05T15:48:42Z","snapshot_observed_at":"2026-07-06T23:47:05.104295Z","submitted_at":"2026-06-05T15:48:42Z","title":"Reversible Foundations: Training a 120B Sparse MoE through State-Preserving Scaling","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-06-27T22:55:09.477413Z"},"links":{"cited_paper":"/paper/2206.05794","citing_paper":"/paper/2606.07404"},"observation_digest":"sha256:590bdf1eb7113d648d1c20722ded947f3bda708ea99fef200e8e016c95ff58a5","observation_id":"f3c56663-8160-484e-8fb3-f7c1313bdbc8","resolution":{"observed_at":"2026-07-02T16:07:09.450255Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2206.05794/citation-record","integrity":"/paper/2206.05794/integrity","json":"/paper/2206.05794/citation-record.json","paper":"/paper/2206.05794"},"outbound":[],"paper":{"arxiv_id":"2206.05794","last_updated":"2024-10-18T21:32:39Z","latest_version":7,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-06T01:24:27.374782Z","submitted_at":"2022-06-12T17:06:35Z","title":"SGD and Weight Decay Secretly Minimize the Rank of Your Neural Network"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 7 inbound Pith citation observations for arXiv:2206.05794."}