{"as_of":"2026-08-08T23:29:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:7280fa2ff9299ea96bf47eb9457568cc70ca46ccd5390b3ed331b75d187dd7b9","coverage":[{"denominator":88,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":88,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T19:04:45.756667Z","state":"measured"},{"denominator":88,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":88,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2502.10216/citation-record","integrity":"/paper/2502.10216/integrity","json":"/paper/2502.10216/citation-record.json","paper":"/paper/2502.10216"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2209.04836","last_updated":"2023-03-01T22:17:16Z","snapshot_observed_at":"2026-07-06T13:50:57.502239Z","submitted_at":"2022-09-11T10:44:27Z","title":"Git Re-Basin: Merging Models modulo Permutation Symmetries","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2209.04836","snapshot_observed_at":"2026-08-07T19:04:45.453751Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.453751Z"},"links":{"cited_paper":"/paper/2209.04836","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:16f81f75220dec2147a9ff55646c2fabbf5e609174a20ccf47d5280d754c6b67","observation_id":"5830d426-4e57-4ddc-b6b9-f4903a8a5731","resolution":{"observed_at":"2026-08-07T19:04:45.453751Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.11983","last_updated":"2023-12-19T09:23:48Z","snapshot_observed_at":"2026-07-06T17:05:10.877197Z","submitted_at":"2023-12-19T09:23:48Z","title":"Fluctuation-based Adaptive Structured Pruning for Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.11983","snapshot_observed_at":"2026-08-07T19:04:45.458177Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.458177Z"},"links":{"cited_paper":"/paper/2312.11983","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:fbd03f2c29f0b3661a99ed476b689844b8c29b09ed7a3dfb49942e6ba79fff32","observation_id":"f813ff86-a2f7-4978-9cd4-021c77e2156f","resolution":{"observed_at":"2026-08-07T19:04:45.458177Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:45.462942Z","title":"Arduino nano 33 ble documentation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.462942Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:03bd82e1218635087fd58cf60cedd509873019d111fb2b0869c2067950bcb19f","observation_id":"c511df0f-c67d-49a5-9dfd-3d00ab79bef8","resolution":{"observed_at":"2026-08-07T19:04:45.462942Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.15024","last_updated":"2024-02-09T17:59:40Z","snapshot_observed_at":"2026-07-06T17:21:01.918787Z","submitted_at":"2024-01-26T17:35:45Z","title":"SliceGPT: Compress Large Language Models by Deleting Rows and Columns","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.15024","snapshot_observed_at":"2026-08-07T19:04:45.466680Z","title":"Ashkboos, M","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.466680Z"},"links":{"cited_paper":"/paper/2401.15024","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:929701ec8febba792f979a224faf943b9f7893b80c9cd8aa77ab4a706e911f51","observation_id":"8fe60b2f-ca4b-4a5f-a851-67bf1dcf51be","resolution":{"observed_at":"2026-08-07T19:04:45.466680Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1512.07548","last_updated":"2015-12-23T17:12:06Z","snapshot_observed_at":"2026-08-06T09:53:26.929066Z","submitted_at":"2015-12-23T17:12:06Z","title":"k-Means Clustering Is Matrix Factorization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1512.07548","snapshot_observed_at":"2026-08-07T19:04:45.470798Z","title":"Bauckhage","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.470798Z"},"links":{"cited_paper":"/paper/1512.07548","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:a933dca8cde55c3c061dc054a1b834b63d419a8ed661aeefd4dfcad744ddb28f","observation_id":"75832b77-9ecf-49d8-9024-59f8e4f2936b","resolution":{"observed_at":"2026-08-07T19:04:45.470798Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:45.474797Z","title":null,"venue":null,"work_id":null,"year":2006},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.474797Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:776f7656c6d3722466a7bdfcf252a29a19053e49bfb8a511f1a696d2e363ba6b","observation_id":"4f258fda-1227-4630-908a-630709ab3d8b","resolution":{"observed_at":"2026-08-07T19:04:45.474797Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2108.07258","last_updated":"2022-07-12T23:45:14Z","snapshot_observed_at":"2026-08-02T09:20:40.804790Z","submitted_at":"2021-08-16T17:50:08Z","title":"On the Opportunities and Risks of Foundation Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2108.07258","snapshot_observed_at":"2026-08-07T19:04:45.478998Z","title":"Bommasani, D","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.478998Z"},"links":{"cited_paper":"/paper/2108.07258","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:fe9348d253f9af6cf8a9eb030d39569b2c9aaa7ec38f997286f66b000f2ab743","observation_id":"307315ee-30e3-4696-bf92-59af69f1bcdd","resolution":{"observed_at":"2026-08-07T19:04:45.478998Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:45.483515Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.483515Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:5c39bf1d35b5d1a3b8c4849da828597efe584686d1f67a980e52d4ebec7f98b5","observation_id":"9502dd06-943d-4001-8c9c-468139f86b5a","resolution":{"observed_at":"2026-08-07T19:04:45.483515Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:45.486933Z","title":"Chang, X","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.486933Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:9b717b79d441166fa51cef974c646462544d2ed1625be6bb12495de1cc902d9a","observation_id":"396fc62e-d69f-4e10-82d2-31decb9c1025","resolution":{"observed_at":"2026-08-07T19:04:45.486933Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1904.01186","last_updated":"2019-12-31T06:58:35Z","snapshot_observed_at":"2026-07-06T07:43:14.473925Z","submitted_at":"2019-04-02T03:00:06Z","title":"Data-Free Learning of Student Networks","version":4},"cited_work":{"arxiv_id":"1904.01186","doi":null,"metadata_source":"pith","pith_arxiv_id":"1904.01186","snapshot_observed_at":"2026-08-07T19:04:46.429914Z","title":"Data-Free Learning of Student Networks","venue":"cs.LG","work_id":"a337b59b-1f12-4da7-b51c-862a12a41731","year":2019},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.490058Z"},"links":{"cited_paper":"/paper/1904.01186","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:8be80639fc5d55d139453fca2ac2c810b8f93fbedfbe7043dfb35bfd35d8499e","observation_id":"929ce3a0-c442-469a-b2fa-47ea3b379ef1","resolution":{"observed_at":"2026-08-07T19:04:46.434786Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:45.493285Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.493285Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:5b2bc5d5122f93d9fdac082e0cd77d5e44256566bd4472a6c59bc6458d4c6950","observation_id":"bfcd368e-77a0-4551-a918-312721890863","resolution":{"observed_at":"2026-08-07T19:04:45.493285Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.06756","last_updated":"2023-11-26T17:24:31Z","snapshot_observed_at":"2026-07-06T16:30:37.867641Z","submitted_at":"2023-10-10T16:27:12Z","title":"Going Beyond Neural Network Feature Similarity: The Network Feature Complexity and Its Interpretation Using Category Theory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.06756","snapshot_observed_at":"2026-08-07T19:04:45.496273Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.496273Z"},"links":{"cited_paper":"/paper/2310.06756","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:8076d160a44cd08f80a3de1dab5823875087b0b424f509ad44a38fe867977edc","observation_id":"53347b84-add1-4584-a5ff-ee18b7aac40d","resolution":{"observed_at":"2026-08-07T19:04:45.496273Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.06767","last_updated":"2024-08-09T03:44:50Z","snapshot_observed_at":"2026-07-06T16:05:42.952649Z","submitted_at":"2023-08-13T13:34:04Z","title":"A Survey on Deep Neural Network Pruning-Taxonomy, Comparison, Analysis, and Recommendations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.06767","snapshot_observed_at":"2026-08-07T19:04:45.499512Z","title":"Cheng, M","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.499512Z"},"links":{"cited_paper":"/paper/2308.06767","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:dfc421ebfc8d1fb3899b494c16247a9874061d3893245e6aa4156258861960c8","observation_id":"9dc4b57c-54b9-48e4-a385-7d720e75e0b8","resolution":{"observed_at":"2026-08-07T19:04:45.499512Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:45.502937Z","title":"Corti, B","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.502937Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:2ceef326d337e83801b2ccdf4a0076f4e1f6d2ab02681d6c45adc3ff3d1706e5","observation_id":"c4c9976d-f03a-445f-bc9c-8caad7729cc9","resolution":{"observed_at":"2026-08-07T19:04:45.502937Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.13349","last_updated":"2025-07-28T14:11:53Z","snapshot_observed_at":"2026-07-31T05:34:54.063652Z","submitted_at":"2023-11-22T12:34:51Z","title":"REDS: Resource-Efficient Deep Subnetworks for Dynamic Resource Constraints","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.13349","snapshot_observed_at":"2026-08-07T19:04:45.506062Z","title":"Corti, B","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.506062Z"},"links":{"cited_paper":"/paper/2311.13349","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:94c939df21fd8ef4ef015d4219c98e90d8edfa4d648c325befea73a19a9ab003","observation_id":"671eabdc-8d31-475c-8557-afd95809d041","resolution":{"observed_at":"2026-08-07T19:04:45.506062Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:45.508925Z","title":null,"venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.508925Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:663edc868a815ecbfd21acb51be5bb85a4afaad03416a357c72f2c38f20ad6fd","observation_id":"d69a7d1c-e0f0-4f0d-842d-e8706d45e406","resolution":{"observed_at":"2026-08-07T19:04:45.508925Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1909.10364","last_updated":"2020-04-19T15:47:42Z","snapshot_observed_at":"2026-07-06T08:23:53.657345Z","submitted_at":"2019-09-23T13:47:51Z","title":"Class-dependent Compression of Deep Neural Networks","version":3},"cited_work":{"arxiv_id":"1909.10364","doi":null,"metadata_source":"pith","pith_arxiv_id":"1909.10364","snapshot_observed_at":"2026-08-07T19:04:46.384053Z","title":"Class-dependent Compression of Deep Neural Networks","venue":"cs.LG","work_id":"b184bc87-4150-46fc-a735-bdb5d3d78824","year":2019},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.511545Z"},"links":{"cited_paper":"/paper/1909.10364","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:d1dc5435f4ecb5bbcef3640dadecd4d0f7ef99426265fa0d7bbab5c36873bbe1","observation_id":"f3e40d45-48f6-41f8-84f6-40d6b389c303","resolution":{"observed_at":"2026-08-07T19:04:46.388333Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2110.06296","last_updated":"2022-07-05T11:40:49Z","snapshot_observed_at":"2026-08-06T06:25:19.456607Z","submitted_at":"2021-10-12T19:28:48Z","title":"The Role of Permutation Invariance in Linear Mode Connectivity of Neural Networks","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2110.06296","snapshot_observed_at":"2026-08-07T19:04:45.514364Z","title":"Entezari, H","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.514364Z"},"links":{"cited_paper":"/paper/2110.06296","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:38aecee6d01a8a275243a4610f87c854c2de19b81ebae1b281c74eda9a133110","observation_id":"4d2fc30f-e89d-4b80-9214-25dd7d9f9e18","resolution":{"observed_at":"2026-08-07T19:04:45.514364Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:45.517367Z","title":"Esp-eye development board - espressif systems","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.517367Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:8207d6d31cbb0c482c9b40901a2bc44d5118d14081ef90b4dc3d966e015a331c","observation_id":"233aa405-3997-4989-92aa-1f595a90414f","resolution":{"observed_at":"2026-08-07T19:04:45.517367Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1912.11006","last_updated":"2020-03-02T12:12:43Z","snapshot_observed_at":"2026-08-04T04:04:57.045893Z","submitted_at":"2019-12-23T18:08:33Z","title":"Data-Free Adversarial Distillation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1912.11006","snapshot_observed_at":"2026-08-07T19:04:45.520302Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.520302Z"},"links":{"cited_paper":"/paper/1912.11006","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:b82a220b9ce9b530c01781b8e24f0aef45d3290c69692a972ae0a030faad0c26","observation_id":"d765240e-b12b-4d90-93f7-d0c44ab407bb","resolution":{"observed_at":"2026-08-07T19:04:45.520302Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1803.03635","last_updated":"2019-03-04T15:51:11Z","snapshot_observed_at":"2026-08-05T23:54:27.386622Z","submitted_at":"2018-03-09T18:51:28Z","title":"The Lottery Ticket Hypothesis: Finding Sparse, Trainable Neural Networks","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1803.03635","snapshot_observed_at":"2026-08-07T19:04:45.523882Z","title":"Frankle and M","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.523882Z"},"links":{"cited_paper":"/paper/1803.03635","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:d5149e87ea8191b206344f2f42fddbcea0606f3d28d26cb80497d2d4e25782f4","observation_id":"4fb39c4d-74c3-422d-b226-173d3a044b51","resolution":{"observed_at":"2026-08-07T19:04:45.523882Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:47.144393Z","title":"Frantar and D","venue":null,"work_id":"60806ad9-7aad-480b-8220-8b387a5b6a02","year":2022},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.527517Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:11cffd57f6a8c37312561f0e4e6989838c8fd0ccd03605cb0de4f8aa707ad8dd","observation_id":"adb2c203-bb89-43cc-ba77-75d0dc6d03b8","resolution":{"observed_at":"2026-08-07T19:04:47.148450Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:45.534347Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.534347Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:b859c0091b0175b49fa851f7d8656a3dcc9e23667f7d156849d88292a67dcfea","observation_id":"2584b9bd-59f7-417c-b509-094aea0096fb","resolution":{"observed_at":"2026-08-07T19:04:45.534347Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2103.13630","last_updated":"2021-06-21T21:01:12Z","snapshot_observed_at":"2026-07-06T10:53:11.986981Z","submitted_at":"2021-03-25T06:57:11Z","title":"A Survey of Quantization Methods for Efficient Neural Network Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2103.13630","snapshot_observed_at":"2026-08-07T19:04:45.538462Z","title":"Gholami, S","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.538462Z"},"links":{"cited_paper":"/paper/2103.13630","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:1dab72778c8b4a9916d213805cbec8e39f7ed0932aa9815f22e4fffebf40f970","observation_id":"99284c22-70a6-4e24-aaa0-5b6928c39c63","resolution":{"observed_at":"2026-08-07T19:04:45.538462Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:45.542033Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.542033Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:a8c00364919bb73c8a001bc1340d77369de0a1a87cf3284deff07dbccb093994","observation_id":"572be022-d713-4bfb-a5ad-3a948ad2d211","resolution":{"observed_at":"2026-08-07T19:04:45.542033Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1705.09280","last_updated":"2017-05-25T17:55:24Z","snapshot_observed_at":"2026-07-06T05:44:18.348582Z","submitted_at":"2017-05-25T17:55:24Z","title":"Implicit Regularization in Matrix Factorization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1705.09280","snapshot_observed_at":"2026-08-07T19:04:45.545536Z","title":"Gunasekar, B","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.545536Z"},"links":{"cited_paper":"/paper/1705.09280","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:b164cd271b3c4b2d1f3af75ac8db756c4afa7b8ed02e31b017f821e5bcf479aa","observation_id":"81489cd6-2d37-4ea2-a5f9-e91e3cdb26d8","resolution":{"observed_at":"2026-08-07T19:04:45.545536Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:47.133863Z","title":"Gupta, A","venue":null,"work_id":"7e8b54e3-14e5-47a2-a11c-b9acc67855d8","year":2015},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.549305Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:01c1f64d5b0b61a1c76f981d23e3c1b038ebe1ce8f5a6fad873767f139286f5b","observation_id":"182ca6f7-9e6a-429a-8227-767a275e80ee","resolution":{"observed_at":"2026-08-07T19:04:47.137545Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:45.552814Z","title":null,"venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.552814Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:bba7c55823a0e7a80d071fd5137a894aaa0cd06115a72bb57393f7047a22db15","observation_id":"212a428d-fce3-4225-82e4-f6443efd0bd9","resolution":{"observed_at":"2026-08-07T19:04:45.552814Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:45.556512Z","title":"Hassibi, D","venue":null,"work_id":null,"year":1993},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.556512Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:8034169e89d1718340b0b3f98b7e559784118ab91e2f5ff2d9b12b741257e21b","observation_id":"0353aa0c-e38e-4708-937b-267fd727e84a","resolution":{"observed_at":"2026-08-07T19:04:45.556512Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:45.560071Z","title":null,"venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.560071Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:dc25552713f1293cbe776e40a4c2c939e71195ba55c473ab3ec336816c9341e2","observation_id":"a1bc26a3-c372-498b-a23c-094b91e463d4","resolution":{"observed_at":"2026-08-07T19:04:45.560071Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:47.104454Z","title":null,"venue":null,"work_id":"d437a9d2-7f2f-40b3-a4e0-bb6552db62dc","year":2018},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.563673Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:707a99d8af2a23b5ac006e078dd7d9669a4a6bb5ce483acf2878851eb1dc6c89","observation_id":"1eea47c7-6314-4142-bf6d-e481f5a1ebfb","resolution":{"observed_at":"2026-08-07T19:04:47.108275Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:47.094208Z","title":null,"venue":null,"work_id":"82d2520a-6cc9-49ca-99b1-2ed6d7f25485","year":2017},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.567147Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:ccc4cec9a4e3222b329846f36036dead5f763a68f2b52e2f65fc33cc607eed79","observation_id":"53f6f026-911c-4425-b4de-2ee7404af7f3","resolution":{"observed_at":"2026-08-07T19:04:47.097674Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1503.02531","last_updated":"2015-03-09T15:44:49Z","snapshot_observed_at":"2026-07-06T04:11:24.157003Z","submitted_at":"2015-03-09T15:44:49Z","title":"Distilling the Knowledge in a Neural Network","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1503.02531","snapshot_observed_at":"2026-08-07T19:04:45.570627Z","title":"Hinton, O","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.570627Z"},"links":{"cited_paper":"/paper/1503.02531","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:522aa4a4f06277087f63ea64e6880ac5db047df15097c9b987e76e8e66d7ef1d","observation_id":"5576a2d9-66d2-4181-8490-2d5753b56cf1","resolution":{"observed_at":"2026-08-07T19:04:45.570627Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.14929","last_updated":"2024-06-14T17:40:29Z","snapshot_observed_at":"2026-08-01T20:58:46.699588Z","submitted_at":"2023-08-28T23:08:15Z","title":"Maestro: Uncovering Low-Rank Structures via Trainable Decomposition","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.14929","snapshot_observed_at":"2026-08-07T19:04:45.574257Z","title":"Horvath, S","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.574257Z"},"links":{"cited_paper":"/paper/2308.14929","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:3c1a6ae1f05f89c248edaf6f7f913a53fab511e37b2f188cb4c9961a4c53c485","observation_id":"74d17b3a-fa00-451d-a0f1-0e6c992e4ce2","resolution":{"observed_at":"2026-08-07T19:04:45.574257Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1607.03250","last_updated":"2016-07-12T07:43:01Z","snapshot_observed_at":"2026-07-06T05:03:13.994105Z","submitted_at":"2016-07-12T07:43:01Z","title":"Network Trimming: A Data-Driven Neuron Pruning Approach towards Efficient Deep Architectures","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1607.03250","snapshot_observed_at":"2026-08-07T19:04:45.577906Z","title":null,"venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.577906Z"},"links":{"cited_paper":"/paper/1607.03250","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:9b51e74c55b9061d8fc25dcd24419e6904f26337c94a05ab050771542a786317","observation_id":"83d9a645-873a-4e78-be6c-345fd995869d","resolution":{"observed_at":"2026-08-07T19:04:45.577906Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1502.03167","last_updated":"2015-03-02T20:44:12Z","snapshot_observed_at":"2026-07-06T04:08:54.419941Z","submitted_at":"2015-02-11T01:44:18Z","title":"Batch Normalization: Accelerating Deep Network Training by Reducing Internal Covariate Shift","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1502.03167","snapshot_observed_at":"2026-08-07T19:04:45.581476Z","title":"Ioffe and C","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.581476Z"},"links":{"cited_paper":"/paper/1502.03167","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:9dfdfea6042b2ecdaa8c8c0cc33ba509da4db07d83911f09c08befa9d736b6e6","observation_id":"d6ed70c4-b511-47fb-9be2-b46c0b90a25a","resolution":{"observed_at":"2026-08-07T19:04:45.581476Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:47.083915Z","title":null,"venue":null,"work_id":"90ea4335-0ca2-494b-97e6-d8d7ced4a67d","year":1991},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.584611Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:82ce7e83203d5e2d07c86e34317ae50560e15e08e6f31a3f74a3d64a9c41853b","observation_id":"6f540b23-3760-4dbd-b7c6-130d1589649e","resolution":{"observed_at":"2026-08-07T19:04:47.087536Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2212.09849","last_updated":"2025-05-21T23:53:00Z","snapshot_observed_at":"2026-08-07T00:34:34.195906Z","submitted_at":"2022-12-19T20:46:43Z","title":"Dataless Knowledge Fusion by Merging Weights of Language Models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2212.09849","snapshot_observed_at":"2026-08-07T19:04:45.588015Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.588015Z"},"links":{"cited_paper":"/paper/2212.09849","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:2d64d78af0a031f5306a4c6aa9a8fb6117de00b6bb4dc31b144ed8fa3d08a1ef","observation_id":"f528b48a-4664-43c4-874e-7f8310ccddea","resolution":{"observed_at":"2026-08-07T19:04:45.588015Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.03094","last_updated":"2024-05-06T14:32:35Z","snapshot_observed_at":"2026-08-06T05:06:28.877535Z","submitted_at":"2023-04-06T14:22:02Z","title":"PopulAtion Parameter Averaging (PAPA)","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.03094","snapshot_observed_at":"2026-08-07T19:04:45.591286Z","title":"Jolicoeur-Martineau, E","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.591286Z"},"links":{"cited_paper":"/paper/2304.03094","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:ce921126a767bcd52cbf4ef9caa4858ecde101ac159a1edc6a9edf48dd97c1e2","observation_id":"e26b8038-8904-4c10-a039-620ef76a52c8","resolution":{"observed_at":"2026-08-07T19:04:45.591286Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.08403","last_updated":"2023-09-25T17:21:00Z","snapshot_observed_at":"2026-08-04T23:44:38.227519Z","submitted_at":"2022-11-15T18:45:26Z","title":"REPAIR: REnormalizing Permuted Activations for Interpolation Repair","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.08403","snapshot_observed_at":"2026-08-07T19:04:45.594385Z","title":"Jordan, H","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.594385Z"},"links":{"cited_paper":"/paper/2211.08403","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:aecf2254db51c479b9b965f5941bdf7d8be426ecf1bc382563de62fd185c13f5","observation_id":"a4d96094-20b1-4b9f-8df0-e02f58942b13","resolution":{"observed_at":"2026-08-07T19:04:45.594385Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:47.072898Z","title":null,"venue":null,"work_id":"3a6e05e8-d7c1-47f4-8e2a-378342208a88","year":2006},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.597261Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:6775f1ffe1ae7db5b333259386d2415cc231dd798e938897f6549c5937241792","observation_id":"d02ab137-d734-4eff-aa19-a3cfd221fc00","resolution":{"observed_at":"2026-08-07T19:04:47.076370Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1511.06530","last_updated":"2016-02-24T11:52:12Z","snapshot_observed_at":"2026-07-06T04:37:12.487978Z","submitted_at":"2015-11-20T09:20:08Z","title":"Compression of Deep Convolutional Neural Networks for Fast and Low Power Mobile Applications","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1511.06530","snapshot_observed_at":"2026-08-07T19:04:45.599861Z","title":null,"venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.599861Z"},"links":{"cited_paper":"/paper/1511.06530","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:60ba1cac7d43724b16a4aba2edb7e1863bcb78452c1f889257968ceb83c8ffc5","observation_id":"663a6b4e-9036-4b4b-9637-5433e2720b7d","resolution":{"observed_at":"2026-08-07T19:04:45.599861Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:47.064071Z","title":"Krizhevsky, G","venue":null,"work_id":"4ff0c777-a414-4552-837d-60b38d276e04","year":2009},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.602730Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:159925c00c64c7c1e1cdd135633713e18a0e7f4368bdd336f8e7b41a081253e9","observation_id":"d8603e1d-c11c-4593-a210-ad5f54296713","resolution":{"observed_at":"2026-08-07T19:04:47.067002Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:47.054043Z","title":"Krizhevsky, V","venue":null,"work_id":"6c150071-fd14-4ed2-9395-c1da769f178c","year":2009},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.605217Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:067a4ed5cc72badc0e5ac7ddd57e0ae0d07fbc792d162fca20c33a392bc11dd2","observation_id":"22f021f2-c5c6-43d0-be70-b20f541f601c","resolution":{"observed_at":"2026-08-07T19:04:47.057371Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:47.044146Z","title":null,"venue":null,"work_id":"913f02b0-5277-40fe-a22f-8d5ad6a95248","year":1955},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.607616Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:9d93b40b8b887807a295458c097aa8d7cd36517d93b5ad6691815e7731e4615c","observation_id":"b9ed4fd1-0987-4729-b0bd-7e2b7fe681da","resolution":{"observed_at":"2026-08-07T19:04:47.047539Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:47.034141Z","title":"Kumar, S","venue":null,"work_id":"c5c7180e-8ee3-4836-9753-8ce7b3ff1867","year":1935},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.610342Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:f2fdbe4b207bda684fe30bea9dec0dcda02ebcc7f8ab92214e901c20ef0edeac","observation_id":"8025760c-8a37-4d33-a241-5bc9d1a1451b","resolution":{"observed_at":"2026-08-07T19:04:47.037980Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1412.6553","last_updated":"2015-04-24T11:40:54Z","snapshot_observed_at":"2026-07-06T04:04:16.777653Z","submitted_at":"2014-12-19T23:02:43Z","title":"Speeding-up Convolutional Neural Networks Using Fine-tuned CP-Decomposition","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1412.6553","snapshot_observed_at":"2026-08-07T19:04:45.613798Z","title":"Lebedev, Y","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.613798Z"},"links":{"cited_paper":"/paper/1412.6553","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:d127bcd060d8c06b31cd2204866fa358887a2d2fd2e165e7e87e0214cdf27485","observation_id":"9149a52a-4836-431b-9e4f-d93c28154a1a","resolution":{"observed_at":"2026-08-07T19:04:45.613798Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:47.022999Z","title":"LeCun, J","venue":null,"work_id":"1d81d1f7-9f87-4a1b-9de8-a42894a6effd","year":1989},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.617221Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:8a9ef9f3e6e921d311bc6711f517b0ea24f4b2cf0c5c7c388c23777cf287b2e3","observation_id":"87c7eb06-8168-496a-a4f5-9ab5a9da93e5","resolution":{"observed_at":"2026-08-07T19:04:47.027251Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.18953","last_updated":"2023-05-30T11:37:41Z","snapshot_observed_at":"2026-08-07T06:35:08.191974Z","submitted_at":"2023-05-30T11:37:41Z","title":"Sit Back and Relax: Learning to Drive Incrementally in All Weather Conditions","version":1},"cited_work":{"arxiv_id":"2305.18953","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.18953","snapshot_observed_at":"2026-08-07T19:04:46.154128Z","title":"Sit Back and Relax: Learning to Drive Incrementally in All Weather Conditions","venue":"cs.CV","work_id":"daa1cb15-7820-4dc5-98dd-8714d9fa94c6","year":2023},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.620555Z"},"links":{"cited_paper":"/paper/2305.18953","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:a3c60e67901393d847476628e6262d065c5bae8c44d88ba2854ece17dff504bd","observation_id":"55e69420-7767-4866-ab86-87c03fa4fc83","resolution":{"observed_at":"2026-08-07T19:04:46.158249Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1605.04711","last_updated":"2022-11-20T14:21:55Z","snapshot_observed_at":"2026-07-06T04:56:23.954139Z","submitted_at":"2016-05-16T10:21:25Z","title":"Ternary Weight Networks","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1605.04711","snapshot_observed_at":"2026-08-07T19:04:45.624168Z","title":null,"venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.624168Z"},"links":{"cited_paper":"/paper/1605.04711","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:142d1a39b4428ccd3587a83ddf59185cc5f3043255c68864fd6cb56f0297c1a2","observation_id":"da936086-d311-4845-81c0-cedf217746c0","resolution":{"observed_at":"2026-08-07T19:04:45.624168Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1608.08710","last_updated":"2017-03-10T17:57:56Z","snapshot_observed_at":"2026-07-06T05:08:42.861486Z","submitted_at":"2016-08-31T02:29:59Z","title":"Pruning Filters for Efficient ConvNets","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1608.08710","snapshot_observed_at":"2026-08-07T19:04:45.631433Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.631433Z"},"links":{"cited_paper":"/paper/1608.08710","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:fffb57cb8fd3086d084e6e93b1efc86d4316d8cb7815053f8e42cf5e287a5c73","observation_id":"d1b96a94-277c-487a-911c-0b93f409b906","resolution":{"observed_at":"2026-08-07T19:04:45.631433Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1511.07543","last_updated":"2016-02-28T22:04:54Z","snapshot_observed_at":"2026-08-03T16:55:19.885866Z","submitted_at":"2015-11-24T02:31:46Z","title":"Convergent Learning: Do different neural networks learn the same representations?","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1511.07543","snapshot_observed_at":"2026-08-07T19:04:45.634667Z","title":null,"venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.634667Z"},"links":{"cited_paper":"/paper/1511.07543","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:b97d105a21bebb43dfcd22818e4cda1d01af0892480c716f889429758bfd0a93","observation_id":"4cb9403f-ada6-41ae-b712-85d76b9b4329","resolution":{"observed_at":"2026-08-07T19:04:45.634667Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07236","last_updated":"2024-04-12T09:34:38Z","snapshot_observed_at":"2026-07-06T17:58:25.894552Z","submitted_at":"2024-04-08T08:50:09Z","title":"Lightweight Deep Learning for Resource-Constrained Environments: A Survey","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07236","snapshot_observed_at":"2026-08-07T19:04:45.638088Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.638088Z"},"links":{"cited_paper":"/paper/2404.07236","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:eaa2c2f47632362f3c6bee52bf93b2b95c8a9dd3c2d039700c2c0fc6259ea959","observation_id":"bb16dbef-3cbb-4a46-9369-6dd9548b7ea0","resolution":{"observed_at":"2026-08-07T19:04:45.638088Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:47.011644Z","title":null,"venue":null,"work_id":"aca69d58-3dbf-4731-987c-b90b79eb3749","year":2017},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.642839Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:6ca73f7619b38eca9036f595985f1a7d5405db14bf111eae3d4743cb96b8f396","observation_id":"b9266278-7316-4c8c-a3c4-6d639892e1b3","resolution":{"observed_at":"2026-08-07T19:04:47.015289Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:47.000447Z","title":null,"venue":null,"work_id":"782fa468-1425-490d-b8f1-e71dbe2a5352","year":2017},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.646146Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:b6d0163dfc7d3cf50a45664f8e6f7a0440aadcbc4c47e3aa827a1c237bc6648d","observation_id":"92fb97a4-81cf-45f5-bb48-f96d99f26464","resolution":{"observed_at":"2026-08-07T19:04:47.004340Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:45.649331Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.649331Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:2f47bb8a68b82c1d0f15274d48a159bf5497ec75765494e5d23ddece34387bdf","observation_id":"a8e475ab-d08b-4203-a58c-b8be45b6ae81","resolution":{"observed_at":"2026-08-07T19:04:45.649331Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2111.09832","last_updated":"2022-08-26T16:43:05Z","snapshot_observed_at":"2026-07-06T12:09:59.122020Z","submitted_at":"2021-11-18T17:59:35Z","title":"Merging Models with Fisher-Weighted Averaging","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2111.09832","snapshot_observed_at":"2026-08-07T19:04:45.652964Z","title":"Matena and C","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.652964Z"},"links":{"cited_paper":"/paper/2111.09832","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:30bbbb299f476d2d07c97bf5237d19b8e7cdce0c966f3917809a4d8f9f6f3c17","observation_id":"aa524705-3cc4-4f7f-9028-713b5264d042","resolution":{"observed_at":"2026-08-07T19:04:45.652964Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.03853","last_updated":"2024-10-11T09:43:32Z","snapshot_observed_at":"2026-08-03T01:48:18.654934Z","submitted_at":"2024-03-06T17:04:18Z","title":"ShortGPT: Layers in Large Language Models are More Redundant Than You Expect","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.03853","snapshot_observed_at":"2026-08-07T19:04:45.656509Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.656509Z"},"links":{"cited_paper":"/paper/2403.03853","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:1c525ad00fa9a996ef906a54513bf83764c6fdd7b60be54ec8ef0e0bceb20080","observation_id":"e734bb89-f805-4837-b4ab-f79626d263bd","resolution":{"observed_at":"2026-08-07T19:04:45.656509Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1609.07843","last_updated":"2016-09-26T04:06:13Z","snapshot_observed_at":"2026-07-06T05:12:10.387914Z","submitted_at":"2016-09-26T04:06:13Z","title":"Pointer Sentinel Mixture Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1609.07843","snapshot_observed_at":"2026-08-07T19:04:45.659885Z","title":"Merity, C","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.659885Z"},"links":{"cited_paper":"/paper/1609.07843","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:c8e5ee50dda23403b227b6d8c55b5c0afeb0ddfba5d9c0ae36b1643d676a9bef","observation_id":"e2f5b13e-3097-4204-9b3c-bf428547052c","resolution":{"observed_at":"2026-08-07T19:04:45.659885Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1905.09768","last_updated":"2019-11-25T20:50:29Z","snapshot_observed_at":"2026-08-05T23:49:08.239249Z","submitted_at":"2019-05-23T16:44:08Z","title":"Zero-shot Knowledge Transfer via Adversarial Belief Matching","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1905.09768","snapshot_observed_at":"2026-08-07T19:04:45.662932Z","title":"Micaelli and A","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.662932Z"},"links":{"cited_paper":"/paper/1905.09768","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:594fa64856cd454f0aa63b26b88219f1df2ed557b457b29e84ec9c19a7162170","observation_id":"c1e9fa88-a3b3-444b-a0a0-b0219e8e561d","resolution":{"observed_at":"2026-08-07T19:04:45.662932Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:46.983402Z","title":null,"venue":null,"work_id":"68ee743e-4fa8-4416-b89b-2e38c7441190","year":null},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.665957Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:c0ced5bc0a0c7ed427f31ff5405ca68436aeecf788b69b18b9e8f8c71face8e0","observation_id":"33bf2f67-7a60-481d-9b48-42c92b2bdaea","resolution":{"observed_at":"2026-08-07T19:04:46.987071Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:46.973410Z","title":"Mordvintsev, C","venue":null,"work_id":"2120d35d-0005-415c-acb1-6b691a119a78","year":2015},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.669140Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:9b16fd0fe459c941bb1af8ed73c0a301abe0fa8d253548112c3b6f2cad03b16d","observation_id":"cba92202-d2cf-4495-8e7c-22e7f5cac1ab","resolution":{"observed_at":"2026-08-07T19:04:46.976932Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:46.963186Z","title":"Jetson nano - nvidia developer","venue":null,"work_id":"1577f833-40eb-49f3-9d54-89c3e42e473b","year":2024},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.672217Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:76d2c918c3cc618125597e5fe1ce83b7b2d947b613be502180e87c428071645e","observation_id":"3969f1f1-120f-40e7-916e-6179c8bfc661","resolution":{"observed_at":"2026-08-07T19:04:46.967091Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:46.952918Z","title":"Papst, D","venue":null,"work_id":"a155e5b8-9c4a-4476-bf84-1171e80611ec","year":2024},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.675682Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:41e22263b204abed94732f03c2ccfc91db9c52a79b99ac37ba0ceb0f28a96da5","observation_id":"2f43983a-0863-4107-87b9-c6769ff38000","resolution":{"observed_at":"2026-08-07T19:04:46.956496Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.14152","last_updated":"2023-06-25T07:38:43Z","snapshot_observed_at":"2026-07-06T15:46:24.940417Z","submitted_at":"2023-06-25T07:38:43Z","title":"Low-Rank Prune-And-Factorize for Language Model Compression","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.14152","snapshot_observed_at":"2026-08-07T19:04:45.679663Z","title":"Ren and K","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.679663Z"},"links":{"cited_paper":"/paper/2306.14152","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:addc71f283d2370ab27326eb1d79c2b206f9f235d4b2ed799983170c03993a5d","observation_id":"d5bde621-4b57-4d03-a51d-687dd0beeaa8","resolution":{"observed_at":"2026-08-07T19:04:45.679663Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:46.942370Z","title":"Rombach, A","venue":null,"work_id":"c9a941b1-464d-4b8e-a077-f1ec7407d0b1","year":2022},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.682917Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:d9bdd833011b982cdab8c66804519e6df8503b7a1d4102132ba2308ec8ab367e","observation_id":"1eb4dad8-1428-40bd-b7b6-8e3fc56c7b8b","resolution":{"observed_at":"2026-08-07T19:04:46.946200Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1701.06538","last_updated":"2017-01-23T18:10:00Z","snapshot_observed_at":"2026-07-06T05:27:13.416519Z","submitted_at":"2017-01-23T18:10:00Z","title":"Outrageously Large Neural Networks: The Sparsely-Gated Mixture-of-Experts Layer","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1701.06538","snapshot_observed_at":"2026-08-07T19:04:45.685636Z","title":"Shazeer, A","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.685636Z"},"links":{"cited_paper":"/paper/1701.06538","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:32fd04a17a13ee8d873fef1e4699752c0ec6a19ccd397fac311e2052a0595caf","observation_id":"714406db-dab9-4dd7-8f33-369ea27abb49","resolution":{"observed_at":"2026-08-07T19:04:45.685636Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1409.1556","last_updated":"2015-04-10T16:25:04Z","snapshot_observed_at":"2026-07-06T03:53:32.549552Z","submitted_at":"2014-09-04T19:48:04Z","title":"Very Deep Convolutional Networks for Large-Scale Image Recognition","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1409.1556","snapshot_observed_at":"2026-08-07T19:04:45.688656Z","title":"Simonyan and A","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.688656Z"},"links":{"cited_paper":"/paper/1409.1556","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:aac2fe53cf105931226931e38c645607fce6e551c30db454d74e6725e1784279","observation_id":"3f00142e-91e5-44ed-a726-e049cc938e2b","resolution":{"observed_at":"2026-08-07T19:04:45.688656Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:46.930587Z","title":null,"venue":null,"work_id":"4ddf52be-cc7e-4307-b6f5-fd518199aae1","year":2020},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.691645Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:303b642e11c3e17c2ce77b0a4c3d6c508a0f3ceeedb92b40f1c1abf63f3ffcd8","observation_id":"09936e55-8cfe-43a8-9340-938f7fb9d73a","resolution":{"observed_at":"2026-08-07T19:04:46.934702Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:46.918810Z","title":"Solodskikh, A","venue":null,"work_id":"61426b9e-cf7d-42f4-a87a-7e25620ee63a","year":2023},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":71,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.694440Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:3211212e68799433417fdafb72bfc42a6efc450c829eff1b66dd084918a74320","observation_id":"3caefbef-7299-4375-96e8-7b0f72ab0653","resolution":{"observed_at":"2026-08-07T19:04:46.922771Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.03053","last_updated":"2024-03-13T02:04:06Z","snapshot_observed_at":"2026-08-04T16:00:39.222919Z","submitted_at":"2023-05-04T17:59:58Z","title":"ZipIt! Merging Models from Different Tasks without Training","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.03053","snapshot_observed_at":"2026-08-07T19:04:45.697164Z","title":"Stoica, D","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":72,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.697164Z"},"links":{"cited_paper":"/paper/2305.03053","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:471b5a1c7dd86e3061661d49d90c043dfbc0d03e36f325d08485fdaa94ce4591","observation_id":"652daf9d-6396-473f-975b-92ee447a2c0e","resolution":{"observed_at":"2026-08-07T19:04:45.697164Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.11695","last_updated":"2024-05-06T17:47:01Z","snapshot_observed_at":"2026-07-06T15:44:42.776459Z","submitted_at":"2023-06-20T17:18:20Z","title":"A Simple and Effective Pruning Approach for Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.11695","snapshot_observed_at":"2026-08-07T19:04:45.700618Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":73,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.700618Z"},"links":{"cited_paper":"/paper/2306.11695","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:a87ff0442f0dfdb8f9dc09bd10ad16d4eafad899284802df6a8a247773799c9e","observation_id":"1a078906-b549-44ac-8573-c7579dc0eca4","resolution":{"observed_at":"2026-08-07T19:04:45.700618Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.07839","last_updated":"2024-02-13T13:19:54Z","snapshot_observed_at":"2026-08-03T02:45:42.249287Z","submitted_at":"2024-02-12T17:50:56Z","title":"Towards Meta-Pruning via Optimal Transport","version":2},"cited_work":{"arxiv_id":"2402.07839","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.07839","snapshot_observed_at":"2026-08-07T19:04:46.017090Z","title":"Towards Meta-Pruning via Optimal Transport","venue":"cs.CV","work_id":"503fbd36-4c7c-4c15-bfd5-7db0df4af96e","year":2024},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":74,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.704396Z"},"links":{"cited_paper":"/paper/2402.07839","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:85604009647b400203e3fa3baa416fb0b07ce8fd64e1886b06850484c74dc191","observation_id":"a39a95a1-c867-4e98-9b46-8fdb9c56bad5","resolution":{"observed_at":"2026-08-07T19:04:46.020543Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2302.13971","last_updated":"2023-02-27T17:11:15Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-02-27T17:11:15Z","title":"LLaMA: Open and Efficient Foundation Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.13971","snapshot_observed_at":"2026-08-07T19:04:45.708042Z","title":"Touvron, T","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":75,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.708042Z"},"links":{"cited_paper":"/paper/2302.13971","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:8f880f79af8b640aae04506e8108d302d3855eb27d841f94d803c9ffdddb1712","observation_id":"530161e0-0af9-4d53-b9fa-da872f5a5f1e","resolution":{"observed_at":"2026-08-07T19:04:45.708042Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.09288","last_updated":"2023-07-19T17:08:59Z","snapshot_observed_at":"2026-08-07T12:56:43.323460Z","submitted_at":"2023-07-18T14:31:57Z","title":"Llama 2: Open Foundation and Fine-Tuned Chat Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.09288","snapshot_observed_at":"2026-08-07T19:04:45.711675Z","title":"Touvron, L","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":76,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.711675Z"},"links":{"cited_paper":"/paper/2307.09288","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:7d2e70603a206e03a428b8764c0b0e75e5bf456a41a5618ad181e3ae7d356268","observation_id":"0eef904b-9dc7-4eb2-999c-43ac8dad55dd","resolution":{"observed_at":"2026-08-07T19:04:45.711675Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:46.861113Z","title":null,"venue":null,"work_id":"858ca718-c89b-48e9-ac4e-8cefe6bd9fd3","year":2020},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.714941Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:4ef705ac46edcff927650fe9e16fc570ef61021fddc14f491b74ccfa9ed7076b","observation_id":"ed7f5154-c43c-46cc-b6c0-1d65dad5b737","resolution":{"observed_at":"2026-08-07T19:04:46.910990Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.13536","last_updated":"2024-05-28T09:28:53Z","snapshot_observed_at":"2026-07-06T15:31:01.106910Z","submitted_at":"2023-05-22T23:19:45Z","title":"Subspace-Configurable Networks","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.13536","snapshot_observed_at":"2026-08-07T19:04:45.718178Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.718178Z"},"links":{"cited_paper":"/paper/2305.13536","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:ff6a45cc106912d29578bb63c42fca164ba916af916c0212ea433720284854d2","observation_id":"c02198c1-6c88-4255-84f8-16f680da1c08","resolution":{"observed_at":"2026-08-07T19:04:45.718178Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:46.790584Z","title":null,"venue":null,"work_id":"2ea36255-ecf5-418a-a32d-817d5e8b3ac9","year":2020},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":79,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.721883Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:03e76eb9be81a0868e48275006bf4638c4988283b298284ab296088024492e64","observation_id":"ce98c7bc-e6f4-4b1d-af6d-2512df5a7fbe","resolution":{"observed_at":"2026-08-07T19:04:46.824182Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:46.750041Z","title":null,"venue":null,"work_id":"026ea7ef-2e26-43f3-a5ea-048aca0fc323","year":2016},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":80,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.725328Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:5334d80b77150378394c2267c215f66e4418641aef046893e3f4e019a4c0fff0","observation_id":"d7a66195-8fea-4d7c-8211-b583ea8b3499","resolution":{"observed_at":"2026-08-07T19:04:46.768334Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2203.05482","last_updated":"2022-07-01T23:48:19Z","snapshot_observed_at":"2026-07-06T12:46:30.292996Z","submitted_at":"2022-03-10T17:03:49Z","title":"Model soups: averaging weights of multiple fine-tuned models improves accuracy without increasing inference time","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2203.05482","snapshot_observed_at":"2026-08-07T19:04:45.728714Z","title":"Wortsman, G","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":81,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.728714Z"},"links":{"cited_paper":"/paper/2203.05482","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:a2bdf30aa27ee67bd27646c7db0251e4ef657d22a3ae7706f9fdcef76fe5151c","observation_id":"a18edc2d-7061-488d-a538-1c4728a36896","resolution":{"observed_at":"2026-08-07T19:04:45.728714Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.05641","last_updated":"2024-09-20T08:27:13Z","snapshot_observed_at":"2026-08-06T00:07:24.781214Z","submitted_at":"2023-06-09T03:00:34Z","title":"Toward Data Efficient Model Merging between Different Datasets without Performance Degradation","version":2},"cited_work":{"arxiv_id":"2306.05641","doi":null,"metadata_source":"pith","pith_arxiv_id":"2306.05641","snapshot_observed_at":"2026-08-07T19:04:45.965326Z","title":"Toward Data Efficient Model Merging between Different Datasets without Performance Degradation","venue":"cs.LG","work_id":"cf0a8c2c-6909-4914-94a2-57b577c81ea1","year":2023},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":82,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.732518Z"},"links":{"cited_paper":"/paper/2306.05641","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:8a2680a793b718e84ea2fb3c5f4d1c49fb6e3db866cecd6b51cae8d5152157ce","observation_id":"4c7a72b7-92ed-4287-9219-13ce0e8c5078","resolution":{"observed_at":"2026-08-07T19:04:45.969081Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1912.08795","last_updated":"2020-06-16T03:30:21Z","snapshot_observed_at":"2026-07-31T09:13:01.272655Z","submitted_at":"2019-12-18T18:50:10Z","title":"Dreaming to Distill: Data-free Knowledge Transfer via DeepInversion","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1912.08795","snapshot_observed_at":"2026-08-07T19:04:45.736217Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":83,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.736217Z"},"links":{"cited_paper":"/paper/1912.08795","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:73bb4af79f9d066e4789c24818a0810df4dee320f607561b98a6c474deba2bf3","observation_id":"3448c147-45bd-4be3-b39d-6105e7028dd9","resolution":{"observed_at":"2026-08-07T19:04:45.736217Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2202.04595","last_updated":"2022-03-11T13:47:55Z","snapshot_observed_at":"2026-08-03T16:28:22.442483Z","submitted_at":"2022-02-09T17:46:49Z","title":"Exploring Structural Sparsity in Neural Image Compression","version":4},"cited_work":{"arxiv_id":"2202.04595","doi":null,"metadata_source":"pith","pith_arxiv_id":"2202.04595","snapshot_observed_at":"2026-08-07T19:04:45.939072Z","title":"Exploring Structural Sparsity in Neural Image Compression","venue":"eess.IV","work_id":"0f2c0f5a-6100-4968-ad90-0073a14e80a0","year":2022},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":84,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.739849Z"},"links":{"cited_paper":"/paper/2202.04595","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:c4266ff759f14cb242c498d00cba38c4f46350e61ecd2ffd54abd8df1d715144","observation_id":"636cfec5-8ddd-4cbd-b8f6-6af3f0b4cacb","resolution":{"observed_at":"2026-08-07T19:04:45.944771Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:46.569348Z","title":null,"venue":null,"work_id":"347269df-e848-4f4c-bfcc-f22f3429a37d","year":2023},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":85,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.743419Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:8c01c5a737ba1ebf499665488b64ebd83b1e34cfc2afc347498e276f90eec151","observation_id":"f83ec558-cad0-48f0-87c1-a76345036b95","resolution":{"observed_at":"2026-08-07T19:04:46.726563Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1702.03044","last_updated":"2017-08-25T13:21:18Z","snapshot_observed_at":"2026-07-06T05:29:31.198406Z","submitted_at":"2017-02-10T02:30:22Z","title":"Incremental Network Quantization: Towards Lossless CNNs with Low-Precision Weights","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1702.03044","snapshot_observed_at":"2026-08-07T19:04:45.746480Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":86,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.746480Z"},"links":{"cited_paper":"/paper/1702.03044","citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:6356d8d23d1f85ba23e44057ff6337e4ab19a715c5b308cb2e41dd3aeb21df38","observation_id":"118d98d6-06e0-4b5c-b659-543c9f8279e9","resolution":{"observed_at":"2026-08-07T19:04:45.746480Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:45.749827Z","title":"@esa (Ref","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":87,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.749827Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:b2653193754561c78ab8114788674973bab15e3d293d0532f891a831c75dcd3f","observation_id":"a010022c-d46f-44b7-9bc5-5dabbf73ec8e","resolution":{"observed_at":"2026-08-07T19:04:45.749827Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:45.753270Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":88,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.753270Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:03a1ab622af19c8e30c2f9cc131e3425e3e5df93e30d1e7ceffe898faef423ea","observation_id":"b304ae6a-4fbd-44cb-87ac-0f27f6108894","resolution":{"observed_at":"2026-08-07T19:04:45.753270Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T19:04:45.756667Z","title":"hs @ @ Ծ-GĀgz Z(nuxʙK;]lv9qQǔ1g#DΝ","venue":null,"work_id":null,"year":1999},"citing_paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress","version":2},"reference_index":89,"source":"arxiv_source","source_observed_at":"2026-08-07T19:04:45.756667Z"},"links":{"citing_paper":"/paper/2502.10216"},"observation_digest":"sha256:4e73f63ac8d367736f678f3a1ad1adbdf7c8e0d82abd96afdb72166e0e3eb687","observation_id":"93a28408-840b-4906-9675-03f0ed39f4b5","resolution":{"observed_at":"2026-08-07T19:04:45.756667Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2502.10216","last_updated":"2025-08-12T07:34:30Z","latest_version":2,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-08T07:46:40.733541Z","submitted_at":"2025-02-14T15:10:43Z","title":"Forget the Data and Fine-Tuning! Just Fold the Network to Compress"},"reference_resolution":{"displayed":88,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":71,"verified_exact":6,"verified_fuzzy":11},"total_outbound_references":88},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 88 of 88 outbound references and 0 inbound Pith citation observations for arXiv:2502.10216."}