{"as_of":"2026-08-10T11:44:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:b016570631216a7a6b35e6c5a78214bf1c911b331f128c341359dd5906e54356","coverage":[{"denominator":91,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":91,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-08T19:58:27.272665Z","state":"measured"},{"denominator":95,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":95,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":4,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":4,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-01T15:35:24.823836Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"cited_work":{"arxiv_id":"2502.05300","doi":"10.48550/arxiv.2502.05300","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.05300","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Ziyin, Y","venue":"ArXiv.org","work_id":"8e11bd7e-9d06-4c6c-b342-bbbdf465dd00","year":2025},"citing_paper":{"arxiv_id":"2605.21933","last_updated":"2026-05-21T03:04:44Z","snapshot_observed_at":"2026-08-02T20:32:58.137009Z","submitted_at":"2026-05-21T03:04:44Z","title":"Thermodynamic Irreversibility of Training Algorithms","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-22T04:38:51.773377Z"},"links":{"cited_paper":"/paper/2502.05300","citing_paper":"/paper/2605.21933"},"observation_digest":"sha256:19976bc7ea1bd7dffb032f4acbaf681c32a9f917e2c4dd73a791334e1d461b04","observation_id":"fb451b88-15e2-4900-b577-343e509179c3","resolution":{"observed_at":"2026-05-22T04:41:04.261521Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"cited_work":{"arxiv_id":"2502.05300","doi":"10.48550/arxiv.2502.05300","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.05300","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Ziyin, Y","venue":"ArXiv.org","work_id":"8e11bd7e-9d06-4c6c-b342-bbbdf465dd00","year":2025},"citing_paper":{"arxiv_id":"2606.04754","last_updated":"2026-08-07T10:31:15Z","snapshot_observed_at":"2026-08-10T11:09:41.310278Z","submitted_at":"2026-06-03T11:37:58Z","title":"Beyond Structural Symmetries: Linear Mode Connectivity via Neuron Identifiability","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-06-28T07:26:25.407267Z"},"links":{"cited_paper":"/paper/2502.05300","citing_paper":"/paper/2606.04754"},"observation_digest":"sha256:d024a3d5518e43f26a596f02e0f2213020f4243a86b133ce10467bfda3281f2a","observation_id":"304ba9d5-e647-474b-b3e4-323991dfb6b5","resolution":{"observed_at":"2026-06-28T07:31:44.820712Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.05300","snapshot_observed_at":"2026-07-12T04:51:08.903724Z","title":"Parameter symmetry potentially unifies deep learning theory.arXiv preprint arXiv:2502.05300, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.03108","last_updated":"2026-07-03T08:43:47Z","snapshot_observed_at":"2026-08-09T04:50:02.780514Z","submitted_at":"2026-07-03T08:43:47Z","title":"Observable- and Positional-Encoding-Dependent Symmetry Readout from Neural Network Weights","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-07-12T04:51:08.903724Z"},"links":{"cited_paper":"/paper/2502.05300","citing_paper":"/paper/2607.03108"},"observation_digest":"sha256:03b7cb6df30fd6a0c878faac817c26dc5ad69b06ba44ff61f71089101c44f435","observation_id":"c2521a33-80e7-40bc-b476-70eada18c351","resolution":{"observed_at":"2026-07-12T04:51:08.903724Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.05300","snapshot_observed_at":"2026-08-01T15:35:24.823836Z","title":"Parameter symmetry potentially unifies deep learning theory","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.18422","last_updated":"2026-07-20T18:12:00Z","snapshot_observed_at":"2026-08-09T18:11:12.381228Z","submitted_at":"2026-07-20T18:12:00Z","title":"PAC--Bayes Bounds on Quotient Parameter Spaces: Geometry-induced Implicit-Bias Priors","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-01T15:35:24.823836Z"},"links":{"cited_paper":"/paper/2502.05300","citing_paper":"/paper/2607.18422"},"observation_digest":"sha256:a7d3caf0eb3af4135997439cc14b64bfdbc21d9924d2178711072590bcbdfd9c","observation_id":"bba359f4-e947-4753-9606-1fc1a34c5940","resolution":{"observed_at":"2026-08-01T15:35:24.823836Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2502.05300/citation-record","integrity":"/paper/2502.05300/integrity","json":"/paper/2502.05300/citation-record.json","paper":"/paper/2502.05300"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:25.762761Z","title":"Sgd learning on neural networks: leap complexity and saddle-to-saddle dynamics","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.762761Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:4a911189d8e78dc47f5d99d3d23e97b78430af13906cdabb19948b62595347cc","observation_id":"fbb9e94e-4724-4b5b-a17a-df6072a4940b","resolution":{"observed_at":"2026-08-08T19:58:25.762761Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1610.01644","last_updated":"2018-11-22T23:40:00Z","snapshot_observed_at":"2026-08-09T22:50:03.459615Z","submitted_at":"2016-10-05T20:59:01Z","title":"Understanding intermediate layers using linear classifier probes","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1610.01644","snapshot_observed_at":"2026-08-08T19:58:25.840740Z","title":"Understanding intermediate layers using linear classifier probes","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.840740Z"},"links":{"cited_paper":"/paper/1610.01644","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:c82d6120afce6c210ee5ebb4c00f946c7838578c24ebda2dee14363628048392","observation_id":"0fa5db58-a585-4d2d-8338-1a896c4d37d4","resolution":{"observed_at":"2026-08-08T19:58:25.840740Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1902.02366","last_updated":"2019-02-06T19:18:13Z","snapshot_observed_at":"2026-08-10T02:50:43.334043Z","submitted_at":"2019-02-06T19:18:13Z","title":"Negative eigenvalues of the Hessian in deep neural networks","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1902.02366","snapshot_observed_at":"2026-08-08T19:58:25.888569Z","title":"Negative eigenvalues of the hessian in deep neural networks","venue":null,"work_id":null,"year":1902},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.888569Z"},"links":{"cited_paper":"/paper/1902.02366","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:cbc7a984984b3e4ccd9d50d322443fde789a50cd6effdf2dc4bc5c35631c1621","observation_id":"cdf6a915-c7de-4b25-880b-9a5bcdc5b5ef","resolution":{"observed_at":"2026-08-08T19:58:25.888569Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:25.897570Z","title":"More is different: Broken symmetry and the nature of the hierarchical structure of science","venue":null,"work_id":null,"year":1972},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.897570Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:60c0772be4ed5dbc0a58bd256036f8051e03c9f98cc4514fe06fcff97dcc6266","observation_id":"e2b513c0-a0e9-4cac-bfff-1554ad1eabe4","resolution":{"observed_at":"2026-08-08T19:58:25.897570Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:25.903734Z","title":"The prevalence of neural collapse in neural multivariate regression","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.903734Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:db3e1fe150ca4fa996a861c71b74735f424b6da526cc237c9c03e5f53dc6fbe2","observation_id":"e84e0daf-8005-4e49-82d7-113abc672980","resolution":{"observed_at":"2026-08-08T19:58:25.903734Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:25.911184Z","title":"Thermal forces from a micro- scopic perspective","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.911184Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:cb67bdb5dd4a433bf3bf6ca3f352db6a6e78148bccd3cca55b3dc8f92f1f708a","observation_id":"abd5a8bc-4def-48d5-a4e1-b72d8947100f","resolution":{"observed_at":"2026-08-08T19:58:25.911184Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2111.00034","last_updated":"2021-12-02T21:06:20Z","snapshot_observed_at":"2026-08-02T07:28:20.038757Z","submitted_at":"2021-10-29T18:22:46Z","title":"Neural Networks as Kernel Learners: The Silent Alignment Effect","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2111.00034","snapshot_observed_at":"2026-08-08T19:58:25.917872Z","title":"Neural networks as kernel learners: The silent alignment effect","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.917872Z"},"links":{"cited_paper":"/paper/2111.00034","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:399ad17e69d9929dc9a82db325c22a148ae362dc44ef8844e4328c1486334b53","observation_id":"b0430263-04f5-4855-ae3f-fd4448a80077","resolution":{"observed_at":"2026-08-08T19:58:25.917872Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:25.927928Z","title":"Revisiting model stitching to compare neural representations","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.927928Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:bc466d96ec1d201e06cc963567b248e0ce67de68ef7b21bece87e9f4181eee44","observation_id":"c9793350-deb5-4771-85fc-08d0cebcbacb","resolution":{"observed_at":"2026-08-08T19:58:25.927928Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:25.931858Z","title":"Nearly-tight vc-dimension and pseudodimension bounds for piecewise linear neural networks","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.931858Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:6bd27331d2bb49f99ebe0ff46a562c6b487ff85585852438db2fa6092791af3e","observation_id":"09de887b-d52e-4a54-8fa5-026717370d41","resolution":{"observed_at":"2026-08-08T19:58:25.931858Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:25.935606Z","title":"Representation learning: A review and new perspectives","venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.935606Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:c458c330809d1760f651f6fecee7dda417d030dbc764d6c1922ec9b57ca2e2c7","observation_id":"90b23e28-6e07-4800-a5d8-3b47a599ee1d","resolution":{"observed_at":"2026-08-08T19:58:25.935606Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2104.13478","last_updated":"2021-05-02T16:16:03Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2021-04-27T21:09:51Z","title":"Geometric Deep Learning: Grids, Groups, Graphs, Geodesics, and Gauges","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2104.13478","snapshot_observed_at":"2026-08-08T19:58:25.939981Z","title":"Geometric deep learning: Grids, groups, graphs, geodesics, and gauges","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.939981Z"},"links":{"cited_paper":"/paper/2104.13478","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:bf07705a25ff9c72f5b997db26cf0cf9197c08ba631a88f89b9989ba500fe231","observation_id":"0fbe2ccd-20b6-4fc0-ab03-14cdeb041108","resolution":{"observed_at":"2026-08-08T19:58:25.939981Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04251","last_updated":"2024-05-29T01:03:31Z","snapshot_observed_at":"2026-07-06T15:39:37.586787Z","submitted_at":"2023-06-07T08:44:51Z","title":"Stochastic Collapse: How Gradient Noise Attracts SGD Dynamics Towards Simpler Subnetworks","version":3},"cited_work":{"arxiv_id":"2306.04251","doi":null,"metadata_source":"pith","pith_arxiv_id":"2306.04251","snapshot_observed_at":"2026-08-08T19:58:27.789679Z","title":"Stochastic Collapse: How Gradient Noise Attracts SGD Dynamics Towards Simpler Subnetworks","venue":"cs.LG","work_id":"95f5b8ec-9420-4864-8709-7ef8e52fcbd8","year":2023},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.944385Z"},"links":{"cited_paper":"/paper/2306.04251","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:b05f80b416fc5ed73150a4f0069a11fe1e07952f4942255b504f8ad39b513a69","observation_id":"e4b0e55b-3a10-42cd-98ae-e70c325ebaf6","resolution":{"observed_at":"2026-08-08T19:58:27.795012Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1812.07956","last_updated":"2020-01-07T16:11:56Z","snapshot_observed_at":"2026-07-06T07:22:15.380277Z","submitted_at":"2018-12-19T14:11:20Z","title":"On Lazy Training in Differentiable Programming","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1812.07956","snapshot_observed_at":"2026-08-08T19:58:25.948535Z","title":"On lazy training in differentiable programming","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.948535Z"},"links":{"cited_paper":"/paper/1812.07956","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:0233d8b1ced850e80f4be7aa35052ff79570dbef2e754ce697a3b7280097a87a","observation_id":"3c788795-92da-4ab6-a283-1c306e49f781","resolution":{"observed_at":"2026-08-08T19:58:25.948535Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2103.00065","last_updated":"2022-11-23T18:09:57Z","snapshot_observed_at":"2026-08-06T02:48:08.712124Z","submitted_at":"2021-02-26T22:08:19Z","title":"Gradient Descent on Neural Networks Typically Occurs at the Edge of Stability","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2103.00065","snapshot_observed_at":"2026-08-08T19:58:25.952588Z","title":"Gradient descent on neural networks typically occurs at the edge of stability","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.952588Z"},"links":{"cited_paper":"/paper/2103.00065","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:216108d26ede6ae41177ba5d0f5e62bf9a1ee426c1017efb202da97935d2c75b","observation_id":"39320332-5ae0-423f-86a8-b2440f25ad47","resolution":{"observed_at":"2026-08-08T19:58:25.952588Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:25.956622Z","title":"A kernel theory of modern data augmentation","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.956622Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:8606274a5abfcebe93871c27ebb7b6013a9e36967c7591822ca99b0f9ee818a0","observation_id":"19961cba-2e61-434a-872a-31d01e79c2c4","resolution":{"observed_at":"2026-08-08T19:58:25.956622Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:25.960204Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.960204Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:e717e6467699d77428b4ea66a08ee4299d929c04b5de9c6768d67da9734cde3e","observation_id":"6f1525b9-7840-4905-96f6-892d1d58783b","resolution":{"observed_at":"2026-08-08T19:58:25.960204Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2010.11929","last_updated":"2021-06-03T13:08:56Z","snapshot_observed_at":"2026-08-10T01:12:16.468283Z","submitted_at":"2020-10-22T17:55:59Z","title":"An Image is Worth 16x16 Words: Transformers for Image Recognition at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.11929","snapshot_observed_at":"2026-08-08T19:58:25.963539Z","title":"An image is worth 16x16 words: Transformers for image recognition at scale","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.963539Z"},"links":{"cited_paper":"/paper/2010.11929","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:833d99256366114ab737be08e637364883dd2df84fd7d32ba00a0165c23504e8","observation_id":"11c9f8cf-a270-42f7-9359-3763d30833a2","resolution":{"observed_at":"2026-08-08T19:58:25.963539Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:25.967069Z","title":"Algorithmic regularization in learning deep homogeneous models: Layers are automatically balanced","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.967069Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:2190ba7877154729e5a95ddabb3f1510328e3ec11b768176b65e87ba5410d0f0","observation_id":"7ff311ab-441b-4304-9d1a-bbe0788ab887","resolution":{"observed_at":"2026-08-08T19:58:25.967069Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:25.970771Z","title":"Why molecules move along a temperature gradient","venue":null,"work_id":null,"year":2006},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.970771Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:5dd1298b2140a57328600bd9e09ab4b29c50f5da39dade28effed921b3492876","observation_id":"fdbeb665-d801-4bcc-8a41-376587748315","resolution":{"observed_at":"2026-08-08T19:58:25.970771Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2110.06296","last_updated":"2022-07-05T11:40:49Z","snapshot_observed_at":"2026-08-10T01:11:52.514832Z","submitted_at":"2021-10-12T19:28:48Z","title":"The Role of Permutation Invariance in Linear Mode Connectivity of Neural Networks","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2110.06296","snapshot_observed_at":"2026-08-08T19:58:25.973999Z","title":"The role of permutation invariance in linear mode connectivity of neural networks","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:25.973999Z"},"links":{"cited_paper":"/paper/2110.06296","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:c1ad8a814428a74111b2f94d9047896358c1331ec25a77526a2902dd4e9f8eef","observation_id":"cdb943ee-f0b3-468f-92d9-7f72fc7730b2","resolution":{"observed_at":"2026-08-08T19:58:25.973999Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:26.006533Z","title":"The representation theory of finite groups","venue":null,"work_id":null,"year":1982},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.006533Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:1e0d9f34562043c8502d5e6436712dffc56aa9769d249ec7c921500acdecb7f4","observation_id":"a376c2a7-05d1-46ca-ba39-d57fae89a905","resolution":{"observed_at":"2026-08-08T19:58:26.006533Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1710.06096","last_updated":"2017-10-17T04:55:14Z","snapshot_observed_at":"2026-08-01T22:56:36.109739Z","submitted_at":"2017-10-17T04:55:14Z","title":"Spontaneous Symmetry Breaking in Neural Networks","version":1},"cited_work":{"arxiv_id":"1710.06096","doi":null,"metadata_source":"pith","pith_arxiv_id":"1710.06096","snapshot_observed_at":"2026-08-08T19:58:27.731738Z","title":"Spontaneous Symmetry Breaking in Neural Networks","venue":"stat.CO","work_id":"e6eff227-2cd3-4328-ba1d-d9e5f5f07f00","year":2017},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.116820Z"},"links":{"cited_paper":"/paper/1710.06096","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:8cb61e9f472ef884d4d2bd75295dad15f9d12fca9175cc49cb4a496b70a50218","observation_id":"4169e64c-e8f4-4592-84b0-13241ca2e1a5","resolution":{"observed_at":"2026-08-08T19:58:27.736581Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.964828Z","title":"A regularity condition of the information matrix of a multilayer perceptron network","venue":null,"work_id":"9a0ecf54-e348-43f3-8c59-bcd6d9dd3976","year":1996},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.207039Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:24a4ed5370d19a33c8b2248283334157294d87b00c14079188bcffab5781b134","observation_id":"6fa6e918-8115-48b7-929b-d35752a271fc","resolution":{"observed_at":"2026-08-08T19:58:28.968538Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:26.248751Z","title":"Local minima and plateaus in hierarchical structures of multilayer perceptrons","venue":null,"work_id":null,"year":2000},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.248751Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:439b056ad08a2cbb396dad2cda85062eaeaf08af8c4bb4381bc60bf2e4f7c93d","observation_id":"b1dc4198-fa78-4955-899b-159ef9168f97","resolution":{"observed_at":"2026-08-08T19:58:26.248751Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2112.15121","last_updated":"2022-01-04T04:06:26Z","snapshot_observed_at":"2026-07-06T12:23:35.429909Z","submitted_at":"2021-12-30T16:36:26Z","title":"On the Role of Neural Collapse in Transfer Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2112.15121","snapshot_observed_at":"2026-08-08T19:58:26.281082Z","title":"On the role of neural collapse in transfer learning","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.281082Z"},"links":{"cited_paper":"/paper/2112.15121","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:5ed91e090d026e24314578b39358364b6e82c3093b0dcb1a95e82570afccb380","observation_id":"ffae2493-54b9-48df-b30a-9c8b1ad6c152","resolution":{"observed_at":"2026-08-08T19:58:26.281082Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.12033","last_updated":"2023-01-28T00:06:22Z","snapshot_observed_at":"2026-08-04T16:52:28.285866Z","submitted_at":"2023-01-28T00:06:22Z","title":"Norm-based Generalization Bounds for Compositionally Sparse Neural Networks","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.12033","snapshot_observed_at":"2026-08-08T19:58:26.287608Z","title":"Norm-based generalization bounds for compositionally sparse neural networks","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.287608Z"},"links":{"cited_paper":"/paper/2301.12033","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:72b68ecca7df172887da1f3cb2376fe3106f9865a9c6fab77bad3c1043361356","observation_id":"2da32e80-2fc6-4087-b0a0-4e39681af113","resolution":{"observed_at":"2026-08-08T19:58:26.287608Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2109.14119","last_updated":"2022-04-19T22:01:13Z","snapshot_observed_at":"2026-08-06T07:19:24.500872Z","submitted_at":"2021-09-29T00:50:00Z","title":"Stochastic Training is Not Necessary for Generalization","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2109.14119","snapshot_observed_at":"2026-08-08T19:58:26.292899Z","title":"Stochastic training is not necessary for generalization","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.292899Z"},"links":{"cited_paper":"/paper/2109.14119","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:c51a620b2c482620426a9dab188ff18103138153ffda59689438c6af211236fd","observation_id":"b4df2c74-8030-4e9d-8d83-3dfa0771e781","resolution":{"observed_at":"2026-08-08T19:58:26.292899Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1909.12051","last_updated":"2019-12-28T10:44:16Z","snapshot_observed_at":"2026-08-10T05:41:04.159293Z","submitted_at":"2019-09-26T12:38:41Z","title":"The Implicit Bias of Depth: How Incremental Learning Drives Generalization","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1909.12051","snapshot_observed_at":"2026-08-08T19:58:26.299922Z","title":"The implicit bias of depth: How incremental learning drives generalization","venue":null,"work_id":null,"year":1909},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.299922Z"},"links":{"cited_paper":"/paper/1909.12051","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:0885525fed91f4e7168591a7604ec4efb9acb97b5e450fb4b623fa200bbdf02f","observation_id":"ef9fc4a5-d2b6-44de-8c06-d9a67e185918","resolution":{"observed_at":"2026-08-08T19:58:26.299922Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.948114Z","title":"On the symmetries of deep learn- ing models and their internal representations","venue":null,"work_id":"fc0ae19c-ab69-4016-97dd-4391db4cc64c","year":2022},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.304284Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:1bd2f2174a1724e01ffb9404fe06e60a5afdadca12edb27028c91ee8075e5d96","observation_id":"d720ecd6-0036-4238-b7cb-07aa0aa7361b","resolution":{"observed_at":"2026-08-08T19:58:28.952586Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.07987","last_updated":"2024-07-25T09:33:50Z","snapshot_observed_at":"2026-08-06T11:02:06.607024Z","submitted_at":"2024-05-13T17:58:30Z","title":"The Platonic Representation Hypothesis","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.07987","snapshot_observed_at":"2026-08-08T19:58:26.309237Z","title":"The platonic representation hypoth- esis","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.309237Z"},"links":{"cited_paper":"/paper/2405.07987","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:d75d940e3475224aeb43576336a054dd2b3c4ebcbd9f193619bf8d7079a77d29","observation_id":"110c634b-c307-48a8-b6ee-8de3875bbd1e","resolution":{"observed_at":"2026-08-08T19:58:26.309237Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1502.03167","last_updated":"2015-03-02T20:44:12Z","snapshot_observed_at":"2026-07-06T04:08:54.419941Z","submitted_at":"2015-02-11T01:44:18Z","title":"Batch Normalization: Accelerating Deep Network Training by Reducing Internal Covariate Shift","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1502.03167","snapshot_observed_at":"2026-08-08T19:58:26.313348Z","title":"Batch normalization: Accelerating deep network training by reduc- ing internal covariate shift","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.313348Z"},"links":{"cited_paper":"/paper/1502.03167","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:cf735f6cc9cd1beccc1b1dcaf09579a138203bab26f93c9369a498c13687ef80","observation_id":"1df1b910-eea1-4f7e-b90a-91cc327d765f","resolution":{"observed_at":"2026-08-08T19:58:26.313348Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1806.07572","last_updated":"2020-02-10T08:39:09Z","snapshot_observed_at":"2026-08-09T19:02:12.754548Z","submitted_at":"2018-06-20T06:35:46Z","title":"Neural Tangent Kernel: Convergence and Generalization in Neural Networks","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1806.07572","snapshot_observed_at":"2026-08-08T19:58:26.366078Z","title":"Neural tangent kernel: Convergence and general- ization in neural networks","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.366078Z"},"links":{"cited_paper":"/paper/1806.07572","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:72314a1d44b3d70929318b46bb84e9e1e66b6849079f2c2d5bc99eac4b9546ee","observation_id":"60506ad6-3c9f-420e-a2a0-29a8d19b2ab5","resolution":{"observed_at":"2026-08-08T19:58:26.366078Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2106.15933","last_updated":"2022-01-31T12:58:39Z","snapshot_observed_at":"2026-08-08T08:13:53.762544Z","submitted_at":"2021-06-30T09:34:05Z","title":"Saddle-to-Saddle Dynamics in Deep Linear Networks: Small Initialization Training, Symmetry, and Sparsity","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2106.15933","snapshot_observed_at":"2026-08-08T19:58:26.454460Z","title":"Saddle-to-saddle dynamics in deep linear networks: Small initialization training, symmetry, and sparsity","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.454460Z"},"links":{"cited_paper":"/paper/2106.15933","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:202bad5d0bcab0b5c52a63ea04e21b0c635195590be8878a58df3ed37403fe92","observation_id":"f0835da8-74a9-4d74-b960-dd28cc4b84d8","resolution":{"observed_at":"2026-08-08T19:58:26.454460Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.937442Z","title":"Sgd on neural networks learns functions of increasing complexity","venue":null,"work_id":"bfae45cd-0be5-4b75-ac54-471d843ce701","year":2019},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.519897Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:ecdf82daded4626e65da34ab6e247c8ad2078ed505f431266470d75d593f483e","observation_id":"f6e49f4c-baf7-4153-8aad-6de5f82cba47","resolution":{"observed_at":"2026-08-08T19:58:28.941148Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.23819","last_updated":"2024-10-31T11:04:07Z","snapshot_observed_at":"2026-08-04T19:15:20.641971Z","submitted_at":"2024-10-31T11:04:07Z","title":"Weight decay induces low-rank attention layers","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.23819","snapshot_observed_at":"2026-08-08T19:58:26.557007Z","title":"Weight decay induces low-rank attention layers","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.557007Z"},"links":{"cited_paper":"/paper/2410.23819","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:e2748b400514ccaace9fbd9993fc3abfcb192d040978735e9a42492133e0d15a","observation_id":"e6cd0852-1bdd-4edd-8e26-659ecad58608","resolution":{"observed_at":"2026-08-08T19:58:26.557007Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:26.562736Z","title":"Similarity of neural network representations revisited","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.562736Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:d1450cafdd7eda1c0d240eaacceca0de6925f44adf66a14fc2c99cb9d00ccf59","observation_id":"544fad2f-adca-4135-bd67-0e9a7ef93002","resolution":{"observed_at":"2026-08-08T19:58:26.562736Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.919506Z","title":"Get rich quick: exact solutions reveal how unbalanced initializations promote rapid feature learning","venue":null,"work_id":"eef06671-aa71-4fd3-b9f7-a3fff221045e","year":2024},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.566785Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:f853cde6dd5c3e254484d64da53165df03e822fb50ef4d8d165444e38d7dffe8","observation_id":"6250652a-1005-48bb-b201-ccfe001bf837","resolution":{"observed_at":"2026-08-08T19:58:28.923284Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.908475Z","title":"Rethinking the limiting dynamics of sgd: modified loss, phase space oscillations, and anomalous diffusion","venue":null,"work_id":"2c8b47e9-74fd-4ce0-82a2-a3f1df9b1b23","year":2021},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.570590Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:867d1b665adc8344962e0da85edacf787e782b9e5821bb6ff17ad8bb0c7b35e5","observation_id":"4f519f3b-3803-4bf3-9ca3-2de54a38fd3e","resolution":{"observed_at":"2026-08-08T19:58:28.912293Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:26.574614Z","title":"Statistical Physics: Volume 5 , volume 5","venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.574614Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:4a9cf78f8211966bf646022614f729a6adce36dff5ec4350cfda741693e94f33","observation_id":"456c805d-9502-40e4-96f5-88fb1a94a332","resolution":{"observed_at":"2026-08-08T19:58:26.574614Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.891762Z","title":"Neural network renormalization group","venue":null,"work_id":"a33059f8-da31-42db-ba72-0c4f2b08fab1","year":2018},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.577845Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:938ac9b673668f93caf6eb8eae3e585e2013cfb544399c691826b8c9b34e1d7c","observation_id":"5f5290bd-32cd-4182-bb75-fda59fc30080","resolution":{"observed_at":"2026-08-08T19:58:28.895037Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1612.09296","last_updated":"2018-01-20T02:45:55Z","snapshot_observed_at":"2026-07-06T05:24:30.687281Z","submitted_at":"2016-12-29T20:57:19Z","title":"Symmetry, Saddle Points, and Global Optimization Landscape of Nonconvex Matrix Factorization","version":3},"cited_work":{"arxiv_id":"1612.09296","doi":null,"metadata_source":"pith","pith_arxiv_id":"1612.09296","snapshot_observed_at":"2026-08-08T19:58:27.625273Z","title":"Symmetry, Saddle Points, and Global Optimization Landscape of Nonconvex Matrix Factorization","venue":"cs.LG","work_id":"4db576b4-9eda-4601-8754-f84aa2fcf87e","year":2016},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.581343Z"},"links":{"cited_paper":"/paper/1612.09296","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:658ad82c387dac6c8ff2c46fc9511649fc7b55f8689a76bb184456d67b2cdf46","observation_id":"fb754d69-c8b8-42e4-a628-e39313cb3ef2","resolution":{"observed_at":"2026-08-08T19:58:27.629445Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.881125Z","title":"Reconciling modern deep learning with traditional op- timization analyses: The intrinsic learning rate","venue":null,"work_id":"5c46e1b8-aea8-4005-ae15-f0e2cd7c03b5","year":2020},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.584859Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:71805cce979102bd4782649587e0368b25324b26c1c59b5c7232417973531265","observation_id":"c6ddf403-17cc-4e02-b6ab-32186839424e","resolution":{"observed_at":"2026-08-08T19:58:28.884853Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.870564Z","title":"What happens after sgd reaches zero loss?–a mathe- matical framework","venue":null,"work_id":"5872b3ed-dcde-478e-8a5c-752e7c92eada","year":2021},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.588978Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:d9b78c6cd705f4865bd1d38a2b83ce1f5ecad6735aebe0d42214ac9455844318","observation_id":"2c4dd624-9b6d-4444-b53c-1455317097e6","resolution":{"observed_at":"2026-08-08T19:58:28.874676Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20231","last_updated":"2024-10-15T12:53:48Z","snapshot_observed_at":"2026-07-06T18:22:48.671656Z","submitted_at":"2024-05-30T16:32:31Z","title":"The Empirical Impact of Neural Parameter Symmetries, or Lack Thereof","version":3},"cited_work":{"arxiv_id":"2405.20231","doi":null,"metadata_source":"pith","pith_arxiv_id":"2405.20231","snapshot_observed_at":"2026-08-08T19:58:27.610734Z","title":"The Empirical Impact of Neural Parameter Symmetries, or Lack Thereof","venue":"cs.LG","work_id":"6381f226-b372-43e8-8c83-90e9f24e5f5e","year":2024},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.592870Z"},"links":{"cited_paper":"/paper/2405.20231","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:db6f2815e7c14496916b3b9162dec06d521ed59b70a82754c0d5f512fbf688c6","observation_id":"afd92de6-9b39-4aba-860c-ad26593f4493","resolution":{"observed_at":"2026-08-08T19:58:27.614501Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.776364Z","title":"Abide by the law and follow the flow: Conserva- tion laws for gradient flows","venue":null,"work_id":"e0ca797a-8ea4-4fe2-98c7-2fe6b5725161","year":2023},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.597134Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:3df2d9d324590e261e254c9ff6d0020515b01ec08baff32fa14f75aace904504","observation_id":"f3e0bf8a-e24f-4158-a8fb-d551e69b26b6","resolution":{"observed_at":"2026-08-08T19:58:28.828220Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1812.09902","last_updated":"2019-04-30T06:01:53Z","snapshot_observed_at":"2026-07-06T07:23:15.236496Z","submitted_at":"2018-12-24T11:52:27Z","title":"Invariant and Equivariant Graph Networks","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1812.09902","snapshot_observed_at":"2026-08-08T19:58:26.601242Z","title":"Invariant and equivariant graph networks","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.601242Z"},"links":{"cited_paper":"/paper/1812.09902","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:db710c3d490716735c7cd7cf95f37d860338285934867d7e27174cc25018b357","observation_id":"898d4dc3-f6a6-426f-9513-e5db7676fbcf","resolution":{"observed_at":"2026-08-08T19:58:26.601242Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.646220Z","title":"The tunnel effect: Building data representations in deep neural networks","venue":null,"work_id":"75a3b2db-bbb9-4fa8-8459-33a7c6fe3c10","year":2024},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.605328Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:138176060ac229eb10a908c532e98bcd11d717ba816aa4792c1d9c3c8ab870ca","observation_id":"cfda142d-1c66-4908-9c0c-dda9b81177a0","resolution":{"observed_at":"2026-08-08T19:58:28.694619Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.602331Z","title":"Charge symmetry, quarks and mesons","venue":null,"work_id":"e1401c25-a308-454c-a139-814c905f49d7","year":1990},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.609171Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:aa15b791fa3c2b7b22c36edcb11f604674bdd974754cdec8ceda6648dc97ab2a","observation_id":"0db7458a-be17-4815-bfa6-e78cff579db5","resolution":{"observed_at":"2026-08-08T19:58:28.605574Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.589990Z","title":"Deep neural networks have an inbuilt occam’s razor","venue":null,"work_id":"575280b4-82f6-4480-bd65-37c311cf75bb","year":2025},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.658635Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:3a2b968797ca005f5d3046889fb46df627a6b389dd9b75b1ba43e9bc37b1a909","observation_id":"baf69777-51bf-421a-b154-3fdc7476cd58","resolution":{"observed_at":"2026-08-08T19:58:28.593785Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1805.12076","last_updated":"2018-05-30T16:50:28Z","snapshot_observed_at":"2026-07-06T06:42:05.870167Z","submitted_at":"2018-05-30T16:50:28Z","title":"Towards Understanding the Role of Over-Parametrization in Generalization of Neural Networks","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1805.12076","snapshot_observed_at":"2026-08-08T19:58:26.724182Z","title":"Towards understanding the role of over-parametrization in generalization of neural networks","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.724182Z"},"links":{"cited_paper":"/paper/1805.12076","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:ef39f13093a6b547af48a049c8396bd2e83c43ef7d028cf723c8900dbd7a4f23","observation_id":"18a02589-0534-4f8f-b85a-64cbb8e78447","resolution":{"observed_at":"2026-08-08T19:58:26.724182Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1412.6614","last_updated":"2015-04-16T18:48:31Z","snapshot_observed_at":"2026-08-04T18:00:43.680629Z","submitted_at":"2014-12-20T06:52:25Z","title":"In Search of the Real Inductive Bias: On the Role of Implicit Regularization in Deep Learning","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1412.6614","snapshot_observed_at":"2026-08-08T19:58:26.794022Z","title":"In search of the real inductive bias: On the role of implicit regularization in deep learning","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.794022Z"},"links":{"cited_paper":"/paper/1412.6614","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:8b16f6d5efd3ac207b2ade29df2b5678bce3b099b3e03cf3a0c93c129b0459c3","observation_id":"b2aebb79-c594-4e82-9bf2-3c550279583c","resolution":{"observed_at":"2026-08-08T19:58:26.794022Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.579065Z","title":"On connected sublevel sets in deep learning","venue":null,"work_id":"c4c14511-0490-40bc-bd9f-eb494c7f900b","year":2019},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.806755Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:9a4dbb186ff6342b0d390ce2a7826e3b8d7cc4a4d9554008b8c37c5f38481278","observation_id":"67a462be-6771-4795-b213-328be8f3b68d","resolution":{"observed_at":"2026-08-08T19:58:28.582580Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.568049Z","title":"Neural networks should be wide enough to learn disconnected decision regions","venue":null,"work_id":"040a6283-8e2a-4532-b040-cbf596ff2c87","year":2018},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.811246Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:f263c7502432f04507d28a28b30b0e81fabd4c264200cc5ac372797443ed732c","observation_id":"dbf40799-3e8b-4468-a31b-500aa07e46c2","resolution":{"observed_at":"2026-08-08T19:58:28.572413Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.556676Z","title":"Prevalence of neural collapse during the terminal phase of deep learning training","venue":null,"work_id":"98f777f5-5433-4d6f-8fc3-7422009d7fb1","year":2020},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.814175Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:53640b58b03bf0f0f94f1224221f0b76364f2869d14ec53f3dbabcecf393a24b","observation_id":"0b336871-9644-47a4-906e-8f30d510ecc8","resolution":{"observed_at":"2026-08-08T19:58:28.560677Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.544270Z","title":"An introduction to quantum field theory","venue":null,"work_id":"4b570610-a092-4393-99d4-5096245faa8f","year":2018},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.817646Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:8d28c603ddead58bf83d494f5c5e9ce65374f81a03c70fcb5cc6dc9556296166","observation_id":"60496c21-2dc2-40c2-a278-1da106efc2c7","resolution":{"observed_at":"2026-08-08T19:58:28.549195Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.13733","last_updated":"2024-11-20T22:20:47Z","snapshot_observed_at":"2026-07-06T19:53:30.020280Z","submitted_at":"2024-11-20T22:20:47Z","title":"On Generalization Bounds for Neural Networks with Low Rank Layers","version":1},"cited_work":{"arxiv_id":"2411.13733","doi":null,"metadata_source":"pith","pith_arxiv_id":"2411.13733","snapshot_observed_at":"2026-08-08T19:58:27.566398Z","title":"On Generalization Bounds for Neural Networks with Low Rank Layers","venue":"cs.LG","work_id":"1c802066-24ef-4d5c-a158-9a9fa75fe202","year":2024},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.820993Z"},"links":{"cited_paper":"/paper/2411.13733","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:a2332e8774216780adffb6db73e12ca326e3cbdad8a055fe4f608dbc2a26f839","observation_id":"32e4a70b-92df-4506-b716-770f59b22ddf","resolution":{"observed_at":"2026-08-08T19:58:27.570879Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.534384Z","title":"Cosmology and broken discrete symmetry","venue":null,"work_id":"859c78cd-c021-4722-9300-ace50234a3cc","year":1991},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.825259Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:1b0f5e51f433c77abd89364cdddf8ff408066252b22f2f3c485693e0eaac2ecb","observation_id":"618f192e-4796-40b4-aa21-f9bb670883da","resolution":{"observed_at":"2026-08-08T19:58:28.537985Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.524268Z","title":"Neural collapse in deep homogeneous classifiers and the role of weight decay","venue":null,"work_id":"ce20c85b-1c9e-4952-8459-e76338358fcc","year":2022},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.828724Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:227921770fee7bf9b31bee7d908ba97c0339d81367707ffe5b71e5ad484887bd","observation_id":"703716d4-3c5a-4dd8-8cd9-1c9da2b03824","resolution":{"observed_at":"2026-08-08T19:58:28.527888Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.513555Z","title":"Feature learning in deep classifiers through intermediate neural collapse","venue":null,"work_id":"7c0c3ca7-67d9-4321-8a80-ff9a6bd15bbe","year":2023},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.832362Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:1dbde7f53cd0003938dee04b1727931a398a43682413d777d585fb16021534b6","observation_id":"73154d2c-26e7-4dd7-87d0-85267098e1ab","resolution":{"observed_at":"2026-08-08T19:58:28.517575Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1312.6120","last_updated":"2014-02-19T17:26:57Z","snapshot_observed_at":"2026-08-08T04:46:43.735461Z","submitted_at":"2013-12-20T20:24:00Z","title":"Exact solutions to the nonlinear dynamics of learning in deep linear neural networks","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1312.6120","snapshot_observed_at":"2026-08-08T19:58:26.835795Z","title":"Exact solutions to the nonlinear dynamics of learning in deep linear neural networks","venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.835795Z"},"links":{"cited_paper":"/paper/1312.6120","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:be975727e344504aa250252297d4ce308b31ffc311eed28e653886a00d143e4a","observation_id":"e379ef95-079f-4e57-ab90-9bc29f88a4a2","resolution":{"observed_at":"2026-08-08T19:58:26.835795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.15438","last_updated":"2023-05-30T17:25:42Z","snapshot_observed_at":"2026-07-06T15:08:34.018146Z","submitted_at":"2023-03-27T17:59:20Z","title":"On the Stepwise Nature of Self-Supervised Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.15438","snapshot_observed_at":"2026-08-08T19:58:26.839502Z","title":"On the stepwise nature of self-supervised learning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.839502Z"},"links":{"cited_paper":"/paper/2303.15438","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:b278af77b029960b5990014d7df90ea4521d43bcef0be29c3399e09b3b3ac763","observation_id":"8b10588b-ee63-4d89-aa7b-ea3ed7c11c06","resolution":{"observed_at":"2026-08-08T19:58:26.839502Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.502925Z","title":"Geometry of the loss landscape in overparameterized neural networks: Symmetries and invariances","venue":null,"work_id":"f6ed4f56-96a1-4aee-8f10-08393f8a3730","year":2021},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.843502Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:4cdb234c7d68d1c093f2d6fa9bf6dfc79bc3228123c12be35bfb763f18bdac84","observation_id":"10a9c39c-b4f6-44ff-bcda-4d942355da75","resolution":{"observed_at":"2026-08-08T19:58:28.506504Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2101.12176","last_updated":"2021-01-28T18:32:14Z","snapshot_observed_at":"2026-08-04T15:33:31.839842Z","submitted_at":"2021-01-28T18:32:14Z","title":"On the Origin of Implicit Regularization in Stochastic Gradient Descent","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2101.12176","snapshot_observed_at":"2026-08-08T19:58:26.846533Z","title":"On the origin of implicit regular- ization in stochastic gradient descent","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.846533Z"},"links":{"cited_paper":"/paper/2101.12176","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:8bb3a9df3281d99457d9e1922811093f3ca68c222e301a458fe66d75c7210b87","observation_id":"3bc75599-3baf-4d5f-b66d-7cd7fac037ac","resolution":{"observed_at":"2026-08-08T19:58:26.846533Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.491677Z","title":"Noether’s learning dynamics: Role of symmetry breaking in neural networks","venue":null,"work_id":"37a39812-9d6a-44bc-ba0d-ddd12499f481","year":2021},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.850721Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:6373deaf729519e13d1d2fb4db8815b513ad75fe005fe8211f276ece3de63319","observation_id":"a5d911d4-e216-452e-98c7-84ce07ac8792","resolution":{"observed_at":"2026-08-08T19:58:28.496013Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.413135Z","title":"Optimizing mode connectivity via neuron alignment.Advances in Neural Information Processing Systems, 33:15300– 15311, 2020","venue":null,"work_id":"6371079d-b946-4c24-b776-6235ada6bc0c","year":2020},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.919279Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:2da438a9c9fd60bd1c6841f69aa37427986c770fab22b7d40aadac77e20c18e0","observation_id":"f675996b-90c8-4c06-b8fd-aa2866ec6b66","resolution":{"observed_at":"2026-08-08T19:58:28.458441Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.312273Z","title":"Equivalences between sparse models and neural networks","venue":null,"work_id":"60ab42fa-0b7b-413c-a17a-ef24e7c05ab7","year":2021},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:26.964311Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:a52078d95d8afb08d2ba84cfdd3155277c63e436a96e86a299bf0f6aacccfcc3","observation_id":"1c54e815-9fbc-4ddf-8101-e43ade21467d","resolution":{"observed_at":"2026-08-08T19:58:28.378434Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"physics/0004057","last_updated":"2000-04-24T15:22:30Z","snapshot_observed_at":"2026-08-07T10:46:08.274157Z","submitted_at":"2000-04-24T15:22:30Z","title":"The information bottleneck method","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"physics/0004057","snapshot_observed_at":"2026-08-08T19:58:27.030727Z","title":"The information bottleneck method","venue":null,"work_id":null,"year":2000},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.030727Z"},"links":{"cited_paper":"/paper/physics/0004057","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:76a93291faf8a7149542f11fdb6dcc4ca34f96f17f5dfb174abe51b0442c7b19","observation_id":"d514f935-1d64-400e-9fbb-14c1375cb4af","resolution":{"observed_at":"2026-08-08T19:58:27.030727Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.262184Z","title":"Deep learning and the information bottleneck principle","venue":null,"work_id":"95b82163-9323-4231-be0d-d7967530a492","year":2015},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.046652Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:7eb620bba3a24b39f1674bb2ad75df487f416373104197efb41463704869293f","observation_id":"62d34c9b-fee2-402a-b445-9d74a027b254","resolution":{"observed_at":"2026-08-08T19:58:28.278050Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:27.051405Z","title":"Attention is all you need","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.051405Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:f8c6047274aea046ed21dee5b9137adfa495b013c02581c23c21c4ad0ddc2c88","observation_id":"bc1831b4-c7d1-473b-b4f4-2df82f419471","resolution":{"observed_at":"2026-08-08T19:58:27.051405Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.244599Z","title":"Asymptotic equivalence of bayes cross validation and widely applicable information criterion in singular learning theory","venue":null,"work_id":"4dadba32-635f-4c24-b847-bf3339a2b0dc","year":2010},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.055243Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:be2720b1e082b7ce48fa840df586bb77a6e253d1a4384885a150f5977f98edb1","observation_id":"a8b11102-9b4e-4e8c-ac4f-2ed442b955ce","resolution":{"observed_at":"2026-08-08T19:58:28.248119Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.17767","last_updated":"2024-11-26T03:54:09Z","snapshot_observed_at":"2026-08-08T09:36:21.431403Z","submitted_at":"2024-05-28T02:46:11Z","title":"Linguistic Collapse: Neural Collapse in (Large) Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.17767","snapshot_observed_at":"2026-08-08T19:58:27.058975Z","title":"Linguistic collapse: Neural collapse in (large) language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.058975Z"},"links":{"cited_paper":"/paper/2405.17767","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:a2325be4a45c9a508d5b1ec249213c458ef6ec2a24756f7840b4f5180c4f4ae1","observation_id":"e1210920-b7e8-4909-8da9-7114c1221a0b","resolution":{"observed_at":"2026-08-08T19:58:27.058975Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.232882Z","title":"The janus effects of sgd vs gd: high noise and low rank","venue":null,"work_id":"8edbbfa0-34cf-4361-935b-1bd7907efda5","year":2023},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.063004Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:bdbcf2eaafe116c20ad9b882c6297047953ef2e193193c1aae603f1a1b14d106","observation_id":"0b57a276-0e39-441c-97ff-e589130f5877","resolution":{"observed_at":"2026-08-08T19:58:28.236617Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.221991Z","title":"Removed for anonymity","venue":null,"work_id":"c8f922f5-2deb-4621-a8bb-d0b66b4b4128","year":null},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.066730Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:5c03a365fc04dc4c1878c0112d8d19d55be57b619acab410a105c7fe6b0161ce","observation_id":"67aeef13-5422-4f85-8dce-df450fb2d042","resolution":{"observed_at":"2026-08-08T19:58:28.225989Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.07085","last_updated":"2025-02-21T11:50:09Z","snapshot_observed_at":"2026-08-04T06:52:50.677485Z","submitted_at":"2024-01-13T14:21:46Z","title":"Three Mechanisms of Feature Learning in a Linear Network","version":3},"cited_work":{"arxiv_id":"2401.07085","doi":null,"metadata_source":"pith","pith_arxiv_id":"2401.07085","snapshot_observed_at":"2026-08-08T19:58:27.474021Z","title":"Three Mechanisms of Feature Learning in a Linear Network","venue":"cs.LG","work_id":"95ba17e5-9286-40ee-a8ce-a26ed2faae35","year":2024},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.070903Z"},"links":{"cited_paper":"/paper/2401.07085","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:4c06dc59d7d61265e89995d75ff0808a73fb77d6124f78471d318fd6447b7de9","observation_id":"97d1b892-f228-43fb-9d0b-1162e6f74a0b","resolution":{"observed_at":"2026-08-08T19:58:27.504143Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.210201Z","title":"Performance-optimized hierarchical models predict neural responses in higher visual cortex","venue":null,"work_id":"2af7cfe0-1a7e-4b05-9e1e-14009920071f","year":2014},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.075225Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:b7f15c6f378a9ea9099b21538f188a9089145b6f51de4fa192f122676afd900c","observation_id":"e1193183-e1e0-4a58-92e4-fb33b9683e8f","resolution":{"observed_at":"2026-08-08T19:58:28.214127Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2011.14522","last_updated":"2022-07-15T17:04:24Z","snapshot_observed_at":"2026-08-09T04:48:30.109419Z","submitted_at":"2020-11-30T03:21:05Z","title":"Feature Learning in Infinite-Width Neural Networks","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2011.14522","snapshot_observed_at":"2026-08-08T19:58:27.079130Z","title":"Feature learning in infinite-width neural networks","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.079130Z"},"links":{"cited_paper":"/paper/2011.14522","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:a828303f8d41f606830d60419dbc96ed349375404569622484b336f48f4d2939","observation_id":"c537dd00-9eef-43a2-9640-cbb631fcc321","resolution":{"observed_at":"2026-08-08T19:58:27.079130Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.199111Z","title":"Visualizing and understanding convolutional networks","venue":null,"work_id":"525129ff-4096-4ed2-a989-546268e87e30","year":2014},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.082756Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:427f6aae11595269fde7c5424a0c89f4ff3e021bb31faf021bc6ffd8c3592b68","observation_id":"1ad77583-9e71-477a-aeff-8cca6dce8487","resolution":{"observed_at":"2026-08-08T19:58:28.202908Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:27.086235Z","title":"Understanding deep learning requires rethinking generalization","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.086235Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:97fecae2a35aa208a95f06f66cce0ce68b5024f198951efca4d39f58dd755d42","observation_id":"ce01cb2b-f913-43a3-be7a-43c305ea4658","resolution":{"observed_at":"2026-08-08T19:58:27.086235Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.180811Z","title":"Symmetry teleportation for accelerated opti- mization","venue":null,"work_id":"868601e8-f9ea-445f-aa09-12fa1eddbbbb","year":2022},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.089776Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:12a283c40c407df1c4837694d11298cc9329e91afca20e9e46a4edc2650fc242","observation_id":"afd32a63-f25f-46c0-a847-6ae1a945278b","resolution":{"observed_at":"2026-08-08T19:58:28.184718Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.170448Z","title":"Understanding mode connectivity via param- eter space symmetry","venue":null,"work_id":"5ffc16be-64b0-4c79-b94f-11eaffa50254","year":2023},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.093114Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:b96314436475cfc805f6c461d7f15be7a749fbf1e60af72a0d2254ce51dc7f19","observation_id":"6c11d115-8c1b-4c79-9f52-197843f2b0f3","resolution":{"observed_at":"2026-08-08T19:58:28.174227Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.13404","last_updated":"2024-04-13T18:28:52Z","snapshot_observed_at":"2026-08-10T06:23:41.230990Z","submitted_at":"2023-05-22T18:35:42Z","title":"Improving Convergence and Generalization Using Parameter Symmetries","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.13404","snapshot_observed_at":"2026-08-08T19:58:27.096742Z","title":"Improving convergence and generalization using parameter symmetries","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.096742Z"},"links":{"cited_paper":"/paper/2305.13404","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:0bf70352efb0be58a9ffd37d6a28662c30bd1333a7ef41e09a9dbf7f89c996e2","observation_id":"0e132255-ffe7-4f2a-ae40-fefbecd25cd5","resolution":{"observed_at":"2026-08-08T19:58:27.096742Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.11787","last_updated":"2024-05-01T23:15:37Z","snapshot_observed_at":"2026-08-10T08:24:33.772076Z","submitted_at":"2022-05-24T05:03:06Z","title":"Quadratic models for understanding catapult dynamics of neural networks","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.11787","snapshot_observed_at":"2026-08-08T19:58:27.100942Z","title":"Quadratic models for understanding neural network dynamics","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.100942Z"},"links":{"cited_paper":"/paper/2205.11787","citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:2504a0498b014293385c4c2e694cfa3d2362191369e688bb566ade82919b64fc","observation_id":"31cb38e9-0646-4d01-a54a-715e3bad237f","resolution":{"observed_at":"2026-08-08T19:58:27.100942Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.160347Z","title":"Symmetry induces structure and constraint of learning","venue":null,"work_id":"88968304-3e18-44dd-adb5-3dd5e0ad8008","year":2024},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.148814Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:49316ebf12d2c353d8a0e915cc4adde01a038fdb29c12254c82ac28d57dd10eb","observation_id":"cfcde326-caec-4264-a254-6f58443deeb0","resolution":{"observed_at":"2026-08-08T19:58:28.163906Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.150158Z","title":"Formation of representations in neural networks","venue":null,"work_id":"12175a8a-ca76-42ae-b89b-2e7e5d49025d","year":2025},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.184656Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:de96fb2d2f7c843bf1b540b397122c0fbe579d91278a3b2b81193dbbdb276b5f","observation_id":"c4a11f64-bc11-4224-ab6d-2b3d5c51fbc2","resolution":{"observed_at":"2026-08-08T19:58:28.153538Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.139122Z","title":"The probabilistic stability of stochastic gradient descent, 2023","venue":null,"work_id":"879d6070-5174-482f-8db0-523f0bba4c6a","year":2023},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.228432Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:1c237229e43a2a4f973333f77d766a823c5e5195030b2fdafd8df6489aa2456d","observation_id":"1fc5c95f-ee4a-4f0f-a201-fd92e0e64528","resolution":{"observed_at":"2026-08-08T19:58:28.142672Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.108213Z","title":"Parameter symmetry and noise equilibrium of stochastic gradient descent","venue":null,"work_id":"eb364295-c5d5-4396-8523-13506e5a7d05","year":2024},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.252739Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:ba25478f2b1f49a0cc36168f2b2ac89e9c14ac357a2d68406a1556fbdb371df1","observation_id":"64140146-7ba7-4918-934b-3957c4774f2d","resolution":{"observed_at":"2026-08-08T19:58:28.131090Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2505.12387","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:27.361665Z","title":"Neural thermodynamics i: Entropic forces in deep and universal representation learning","venue":null,"work_id":"4eb01824-e125-4317-96c2-971db2e9af7f","year":2025},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":87,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.256398Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:1088561f284f6188395979625056cfa48b912b27e5822412a58218e5f2a3a547","observation_id":"aba9f7ed-3c4f-4259-be47-2b4a08e10013","resolution":{"observed_at":"2026-08-08T19:58:27.389789Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:28.061407Z","title":"Remove symmetries to control model expressivity","venue":null,"work_id":"7594fb8e-343d-484a-b1fa-d8b05338ea92","year":2025},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":88,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.260375Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:fdc4f49ef4e076727b49dee9cf34d4a1653d71201e649a5ea1d6d9a1a072ad4e","observation_id":"1f65beb9-de34-45fb-9379-d3b809dae6a2","resolution":{"observed_at":"2026-08-08T19:58:28.082683Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:27.998062Z","title":null,"venue":null,"work_id":"25e3d5fe-6f2a-4b80-a287-addbacb404b7","year":null},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":89,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.264524Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:107890000c21998e05a7e0af9173fcd21ff5313eb5e10b17c77b983f08dd9ced","observation_id":"afeaa83c-6586-4a35-a8ec-e4d68a06e506","resolution":{"observed_at":"2026-08-08T19:58:28.029857Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:27.943751Z","title":"Here, A ∶= PG ∑x∇θf(x, θ0)T∇θf(x, θ0)PG and A+ denotes the Moore–Penrose inverse of A","venue":null,"work_id":"cccbc574-f263-47d0-8a1b-89e1f8533787","year":null},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":90,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.268644Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:d7a84624b022246c408033823b2ea8189267a31d24eb76de15af05552f622cb2","observation_id":"d153e544-46d9-4131-b1c8-ddbf21638d23","resolution":{"observed_at":"2026-08-08T19:58:27.963863Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T19:58:27.915023Z","title":"(21) Therefore, g(x, θ) simplifies to a kernel model g(x, θ)=∇θ0 f(x, θ0)PGθ","venue":null,"work_id":"3ee25f63-d28c-4575-947e-e14d7ab4b89c","year":null},"citing_paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory","version":2},"reference_index":91,"source":"pdf_text","source_observed_at":"2026-08-08T19:58:27.272665Z"},"links":{"citing_paper":"/paper/2502.05300"},"observation_digest":"sha256:444514f7d9f39943579edcc664e41a9a61cba2af01a22e8a97fe99183754e3f1","observation_id":"7c9dc0bf-cb19-4a5b-ad97-1d34d7d79b35","resolution":{"observed_at":"2026-08-08T19:58:27.918485Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2502.05300","last_updated":"2025-05-23T17:22:54Z","latest_version":2,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-09T04:49:29.861223Z","submitted_at":"2025-02-07T20:10:05Z","title":"Parameter Symmetry Potentially Unifies Deep Learning Theory"},"reference_resolution":{"displayed":91,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":46,"verified_exact":7,"verified_fuzzy":38},"total_outbound_references":91},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 91 of 91 outbound references and 4 inbound Pith citation observations for arXiv:2502.05300."}