{"as_of":"2026-08-15T12:17:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:ca33615ce1c8c20e458ff01bfabcd50a619250d52c0a89179fd4eb17af1daf37","coverage":[{"denominator":47,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":47,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-11T20:57:20.333112Z","state":"measured"},{"denominator":49,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":49,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-15T06:32:42.880941+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-28T02:35:39.845487Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-02T11:56:56.059724Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"cited_work":{"arxiv_id":"2412.05169","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2412.05169","snapshot_observed_at":"2026-07-02T11:56:56.059724Z","title":"Towards understand- ing the role of sharpness-aware minimization algorithms for out-of-distribution generalization","venue":null,"work_id":"8e2bab07-8806-4cc9-aee6-9f5e4d99ceec","year":2024},"citing_paper":{"arxiv_id":"2604.08192","last_updated":"2026-04-09T12:44:19Z","snapshot_observed_at":"2026-08-15T01:23:14.279718Z","submitted_at":"2026-04-09T12:44:19Z","title":"Inside-Out: Measuring Generalization in Vision Transformers Through Inner Workings","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-05-10T18:28:49.855280Z"},"links":{"cited_paper":"/paper/2412.05169","citing_paper":"/paper/2604.08192"},"observation_digest":"sha256:a0a4739b530475a813cb7d0e4256a69315aea2135f763575ae6e05620bc125ea","observation_id":"0f7c16eb-4415-46e7-a2c6-b625c9a701f1","resolution":{"observed_at":"2026-05-11T00:30:54.169150Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"cited_work":{"arxiv_id":"2412.05169","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2412.05169","snapshot_observed_at":"2026-07-02T11:56:56.059724Z","title":"Towards understand- ing the role of sharpness-aware minimization algorithms for out-of-distribution generalization","venue":null,"work_id":"8e2bab07-8806-4cc9-aee6-9f5e4d99ceec","year":2024},"citing_paper":{"arxiv_id":"2606.06418","last_updated":"2026-06-04T17:22:58Z","snapshot_observed_at":"2026-08-02T08:47:15.917924Z","submitted_at":"2026-06-04T17:22:58Z","title":"Double Preconditioning (DoPr): Optimization for Test-Time Performance, not Validation Loss","version":1},"reference_index":188,"source":"arxiv_source","source_observed_at":"2026-06-28T02:35:39.845487Z"},"links":{"cited_paper":"/paper/2412.05169","citing_paper":"/paper/2606.06418"},"observation_digest":"sha256:40e052c5b05bc15ee24c6839dc63dc681529fc554faa67d59f8920ca44306539","observation_id":"49b96ddb-2e88-4616-b950-5deefbf562d3","resolution":{"observed_at":"2026-07-02T11:56:56.061073Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2412.05169/citation-record","integrity":"/paper/2412.05169/integrity","json":"/paper/2412.05169/citation-record.json","paper":"/paper/2412.05169"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:21.012488Z","title":"Sharpness-aware minimization leads to low-rank features","venue":null,"work_id":"1d6410e8-898f-48a8-bdb3-02675a279586","year":2023},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.142653Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:6b6f71da0012ec91a857f23cac97a93a3a6e4edfdc1907db50804a2083f0203e","observation_id":"53d52d20-9b84-4975-aa58-2c6d5091879b","resolution":{"observed_at":"2026-08-11T20:57:21.016738Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.999506Z","title":"Invariant risk minimization, 2020","venue":null,"work_id":"9af6c3f0-0b73-4121-ac20-1ca49c8f1e89","year":2020},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.147305Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:15dbfa84aec5ca20aad04444d6caa6ce99b8bd8a750814739394731a1c956b2f","observation_id":"cd3c0d11-c818-4154-a6f6-bc1e3711e1ca","resolution":{"observed_at":"2026-08-11T20:57:21.003825Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.986010Z","title":"Why is sam robust to label noise? In International Conference on Learning Representations (ICLR), 2024","venue":null,"work_id":"230ca40e-29b9-495c-a2f6-bf7cf8bd1827","year":2024},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.151686Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:6d90eb76f60d6ccbbe6b79fada8d3a28897433800e1f775ff22d47c94597fc8a","observation_id":"326dcb55-61eb-448d-bd35-3df90565700b","resolution":{"observed_at":"2026-08-11T20:57:20.990739Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.156061Z","title":"A theory of learning from different domains","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.156061Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:e0761ba5c2f2789f6e65e01323b5b1cece4955caea18ab303a7a6ab79ab60884","observation_id":"cef66bce-7442-4f02-8686-8483c30fdf7b","resolution":{"observed_at":"2026-08-11T20:57:20.156061Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.972841Z","title":"Blackard and Denis J","venue":null,"work_id":"f09fe0de-3752-4c2a-bff3-236df7ac14ec","year":1999},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.160395Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:922fcbd9209d1e1c96b4a64a8aebdaef752400409f9cfdc2065b8679b91c8151","observation_id":"bfd2176b-a834-4e3b-8d48-faee902ba10a","resolution":{"observed_at":"2026-08-11T20:57:20.977016Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.164537Z","title":"Optimization methods for large-scale machine learning","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.164537Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:0c3a2ec23618891e21f08eea59ad04d1ef009a6df0bf0f6e6c2bcbddb118910d","observation_id":"f59bb4a3-51aa-4af8-b3b7-a54bb4170414","resolution":{"observed_at":"2026-08-11T20:57:20.164537Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.169292Z","title":"Swad: Domain generalization by seeking flat minima","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.169292Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:6aba5902359fbf6de1f2f72c066bc066aeb3c01e8c8f3b97012f6bd4ab97a017","observation_id":"e4cfda1f-0013-441a-9331-6e90097b9ec8","resolution":{"observed_at":"2026-08-11T20:57:20.169292Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.173209Z","title":"Entropy-sgd: Biasing gradient descent into wide valleys","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.173209Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:7dd9a6d26432d49e85f5857bc532ba2b9d52b98062c50c6087afd4bc35009057","observation_id":"3323ed48-6f9d-4c4b-873f-8e08938fa996","resolution":{"observed_at":"2026-08-11T20:57:20.173209Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.934811Z","title":"Why does sharpness-aware minimization generalize better than sgd? Advances in neural information processing systems, 36, 2024","venue":null,"work_id":"281ae29b-1dfb-40d1-b1c0-2ec8f967303f","year":2024},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.177633Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:0945a47b26b89813feba66da1a55a7f8690d858167e8fa4aaa377c4ed38aa51a","observation_id":"1bb1305b-c76e-46f4-8d73-83cfe3ed0f73","resolution":{"observed_at":"2026-08-11T20:57:20.939243Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.920669Z","title":"Sharp minima can generalize for deep nets","venue":null,"work_id":"d38503c7-a728-45d1-b2b8-66ae710fa6c0","year":2017},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.181588Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:26378ecfc8c64875419d6691f8be59cde7e0c60814db3e7439ffde900e92e14f","observation_id":"56f2178b-ecc2-452f-ab94-8c198ec91200","resolution":{"observed_at":"2026-08-11T20:57:20.924896Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.906431Z","title":"Efficient sharpness-aware minimization for improved training of neural networks","venue":null,"work_id":"e0fdb2f9-0211-4369-bc1f-6f1dfc6baa86","year":2022},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.185734Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:e39dd4c2f033674cc38f312d520e5d46e19a233822c878ef1e4523b8ccfb32f2","observation_id":"e804b8f4-de11-4a62-913e-ae4ffc1ae4dd","resolution":{"observed_at":"2026-08-11T20:57:20.911335Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.892393Z","title":"Sharpness-aware minimization for efficiently improving generalization","venue":null,"work_id":"f87bb46b-cd79-4187-b058-dbfb662058e8","year":2021},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.190169Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:dd7d2a77680811e1ad212a54d85ea9413e84505d17b5aa8b8f0e77b148abbccb","observation_id":"ea62cb9c-aa35-4b07-be02-cef497900085","resolution":{"observed_at":"2026-08-11T20:57:20.896912Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.194571Z","title":"Sachs, Brian Yin, Crystal Lee, Philipp Krähenbühl, and Alexei A","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.194571Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:41bf8ae427cdfa058ff381dcc13855ccfbb1d35468c13f02f7289797096e2b05","observation_id":"a1b8ac61-4540-4f4a-b74b-b99d171ca731","resolution":{"observed_at":"2026-08-11T20:57:20.194571Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.878730Z","title":"Gradual domain adaptation: Theory and algorithms","venue":null,"work_id":"571f8eb3-1714-4e2d-b06c-03deb4502b7a","year":2024},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.198841Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:6c19080f938c29b4ace9d453199f3e4137762e7aaf733d2a147207dbc08c0f8d","observation_id":"32d86d00-2373-47ef-9102-b77c04fc4bd0","resolution":{"observed_at":"2026-08-11T20:57:20.883111Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.202806Z","title":"Flat Minima","venue":null,"work_id":null,"year":1997},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.202806Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:6810ece287e54c5c07bc48002fc1df2a56b4f141ca43e6c4c0476f64c1401708","observation_id":"8b28a5c8-02a6-463b-8316-608541f4ae8b","resolution":{"observed_at":"2026-08-11T20:57:20.202806Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.865507Z","title":"Batch normalization: Accelerating deep network training by reducing internal covariate shift","venue":null,"work_id":"4f3cc507-3a44-4849-9e6a-ae54196b77c5","year":2015},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.206925Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:775b8473138761c962d774bb7c44531b287f5205803b4a8cb4f0167dc1328b2d","observation_id":"eba6f2b8-4817-4264-ab84-51fd29dc21f2","resolution":{"observed_at":"2026-08-11T20:57:20.869919Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.851904Z","title":"Averaging weights leads to wider optima and better generalization","venue":null,"work_id":"057999a9-8574-4724-913f-23adbc871ace","year":2018},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.210866Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:eae21c39e02cfeb39b26ea2dcc27a0d88785cd75e991d0ff6abf8825196dbae4","observation_id":"8068fa23-5c1d-4d43-9543-78bc266db910","resolution":{"observed_at":"2026-08-11T20:57:20.856797Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.838361Z","title":"Fantastic generalization measures and where to find them, 2020","venue":null,"work_id":"cc4bb0a4-2181-465d-b2d1-bc38e09deae9","year":2020},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.214793Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:39c23bc1bcfad093b142c8807681c761a6e69237c70bea60cce407f7c5067765","observation_id":"ad513d2d-8950-48d2-af93-21a1ef5f2352","resolution":{"observed_at":"2026-08-11T20:57:20.842946Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.218918Z","title":"When do flat minima optimizers work? Advances in Neural Information Processing Systems, 35: 0 16577--16595, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.218918Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:8a135408f23db2d5de6ff774fed097f7cb04b3b4cde5d7a458b07c9bd89823a6","observation_id":"e907466e-f605-436c-8d43-d677cb2de727","resolution":{"observed_at":"2026-08-11T20:57:20.218918Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.815585Z","title":"On large-batch training for deep learning: Generalization gap and sharp minima","venue":null,"work_id":"6ef52ff0-5710-42cf-91a9-2d194e806ffe","year":2017},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.223152Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:7f032651de4e91b1d2087a8e44e7cddb4b00fcd49ac7ddc16afe4918ab8ad722","observation_id":"16eeaee8-1484-47b8-810e-b5b5efbbc365","resolution":{"observed_at":"2026-08-11T20:57:20.820045Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.227029Z","title":"Fast and scalable bayesian deep learning by weight-perturbation in adam","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.227029Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:3b303542ad733711b6f3e1ad939af902666534dd218ce59657aed897b71481a6","observation_id":"d1a295b6-526b-4f72-a5f1-7d3a84d81d86","resolution":{"observed_at":"2026-08-11T20:57:20.227029Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.793766Z","title":"F isher SAM : Information geometry and sharpness aware minimisation","venue":null,"work_id":"8da11ef2-fda6-4c6d-aed1-35bd68575f88","year":2022},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.230965Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:c4892c99cb973a062d4d2ba583f6bd8346d175070208664cf56984ca8d63554d","observation_id":"e3dd3c6d-6e74-45ee-96d9-a59b8ca70b0b","resolution":{"observed_at":"2026-08-11T20:57:20.798270Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.780580Z","title":"Understanding self-training for gradual domain adaptation","venue":null,"work_id":"44206b8f-ba8a-42f8-9d31-d30ef33a288a","year":2020},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.234967Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:d3841f1032ae0e73053dc186ad25c3199eeb0c077bd9c79396bd7f3189a66699","observation_id":"0d041e57-4433-4c12-9548-5bdb7bb44a5c","resolution":{"observed_at":"2026-08-11T20:57:20.784784Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.239198Z","title":"Kuznetsov and M","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.239198Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:773fcd0e6bf9063f63cb901ee7e40744c636d5210d238d87bd68476e41c306f2","observation_id":"30464646-22c9-4922-82e6-a2a293b941f1","resolution":{"observed_at":"2026-08-11T20:57:20.239198Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.767297Z","title":"Discrepancy-based theory and algorithms for forecasting non-stationary time series, 2020 b","venue":null,"work_id":"90552bdb-13c3-4fec-9418-9eeda0d3ef3b","year":2020},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.243476Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:ab00425d02f5759d2431026aaae36f5abb03191c19e51579ed64bf28a45db7ff","observation_id":"68cb0938-6eeb-4a14-9914-285ec8fa88c5","resolution":{"observed_at":"2026-08-11T20:57:20.771771Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.753759Z","title":"Asam: Adaptive sharpness-aware minimization for scale-invariant learning of deep neural networks","venue":null,"work_id":"d8eb93a0-403b-4d1d-b1b1-7abf466f5b82","year":2021},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.247385Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:550f32a196ec59fd7864970385ebbef61096d837c0d4c0343a6b5d2f6ebf3555","observation_id":"2b2308ab-22c3-4c96-9b90-07f4a048e891","resolution":{"observed_at":"2026-08-11T20:57:20.758461Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.740694Z","title":"The mnist database of handwritten digits","venue":null,"work_id":"81c62acf-6a4c-4965-bb9d-e46e4fdc992e","year":1998},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.251379Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:fa22654f2c57b8bd17bd91b1517cf7674fe99c145de91819d2f46759e7ce8df8","observation_id":"c8e1327d-74f2-430c-9182-b9c3b57b38b1","resolution":{"observed_at":"2026-08-11T20:57:20.744921Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.726506Z","title":"Friendly sharpness-aware minimization","venue":null,"work_id":"33359902-d3aa-4865-8c7e-f6112d83ab40","year":2024},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.255728Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:256e42c86cd5e4856d6d90a649afcdf7b38ac719aaacaa4f52c1baea92d5ede4","observation_id":"c48f1f17-24d1-446d-a2e4-e5258e0f3a74","resolution":{"observed_at":"2026-08-11T20:57:20.731273Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.712679Z","title":"Fisher-rao metric, geometry, and complexity of neural networks","venue":null,"work_id":"877cb462-e228-42f0-8cbf-805c30f39a22","year":2019},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.259880Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:2ab9d8b1e3cd64e24be7cd0de0be84725e383a708c7c0bd2670e14580c97a491","observation_id":"edea56b6-6028-4b08-8867-dc6944d4bf77","resolution":{"observed_at":"2026-08-11T20:57:20.717302Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.696767Z","title":"Towards efficient and scalable sharpness-aware minimization","venue":null,"work_id":"b8987f25-675a-4552-b434-44b573548aeb","year":2022},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.263677Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:943d881198697556acc03d447ce272baa1a28f4f9de89eb89d9e5096dc003bf0","observation_id":"78ee4490-7101-47bd-9db0-b2cdf787968b","resolution":{"observed_at":"2026-08-11T20:57:20.702106Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.681650Z","title":"New insights and perspectives on the natural gradient method","venue":null,"work_id":"061f3ff1-615a-4c3f-9071-8f605c3b4716","year":2020},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.267606Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:95c38d2e30ef6e45d9146f2df5be2316b7395d3582f113ae10f281230412047e","observation_id":"3425d9dd-de40-4253-9dd0-a165d72af810","resolution":{"observed_at":"2026-08-11T20:57:20.686405Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.667053Z","title":"Normalization layers are all that sharpness-aware minimization needs","venue":null,"work_id":"6edf9ef6-a2ab-4f2f-bf95-8e7fb205db32","year":2024},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.271613Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:4bdc9597c58f955141146b1cc7b3c9a1160af9c41400a619653e48a002517a1d","observation_id":"77960d26-6c69-4e25-b5d1-c90820c7ac9d","resolution":{"observed_at":"2026-08-11T20:57:20.671751Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2210.12864","last_updated":"2022-10-23T21:49:58Z","snapshot_observed_at":"2026-08-13T13:58:42.451064Z","submitted_at":"2022-10-23T21:49:58Z","title":"K-SAM: Sharpness-Aware Minimization at the Speed of SGD","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.12864","snapshot_observed_at":"2026-08-11T20:57:20.275679Z","title":"K-sam: Sharpness-aware minimization at the speed of sgd, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.275679Z"},"links":{"cited_paper":"/paper/2210.12864","citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:c8c0cb11883539dd888e4c0068093767f22342bb00249bb2220c331a22039250","observation_id":"f52db1e5-fac7-4715-926b-4ef61a7d1cfd","resolution":{"observed_at":"2026-08-11T20:57:20.275679Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.653357Z","title":"Online learning via sequential complexities","venue":null,"work_id":"25d99c59-247f-465b-be8b-823ddf62aa18","year":2015},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.280259Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:8a5c61bf32d917c66da284fba5b2da37a28a8fbaa355e7630ffb215e15262d4d","observation_id":"d9e20b30-282c-44e8-900e-f91044d26eea","resolution":{"observed_at":"2026-08-11T20:57:20.657780Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.639147Z","title":null,"venue":null,"work_id":"b2b88f91-d353-47df-9243-965513036f9c","year":2020},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.284168Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:eaca8c3806874b95422fa28ce2467239e7787cbdb59125e8b1043aa97a1add4b","observation_id":"84e90a8e-3bb2-4182-bb82-d5dd9c71bcab","resolution":{"observed_at":"2026-08-11T20:57:20.643853Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.625630Z","title":"Sharpness-aware minimization enhances feature quality via balanced learning","venue":null,"work_id":"b8f53c9f-b9ea-41ca-8925-334a140d66a3","year":2024},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.288677Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:eae9ea3f2ad4bf207e4d3d4dc48b2eedeff849d3fb29aa7c878818060c848591","observation_id":"3ec86a82-7520-4044-b248-0f2590567a42","resolution":{"observed_at":"2026-08-11T20:57:20.629920Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.292565Z","title":"Dropout: A simple way to prevent neural networks from overfitting","venue":null,"work_id":null,"year":1929},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.292565Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:c13b21bb4347ce967958868ab850b3e00452c7fe313a17ce12271e25ddba26f9","observation_id":"88e1c4f1-f009-4d16-bc7a-98d29e08a9d2","resolution":{"observed_at":"2026-08-11T20:57:20.292565Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.602926Z","title":"Understanding gradual domain adaptation: Improved analysis, optimal path and beyond","venue":null,"work_id":"7d1f00f1-2984-452e-87b6-fbc12482c3a8","year":2022},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.296515Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:93f3475d0a043dbbaf62bd8a8b495372febd404c9a4761f25a58d1974843cacd","observation_id":"8e2e8929-5957-4a4e-928b-09e5d158d607","resolution":{"observed_at":"2026-08-11T20:57:20.607465Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.588347Z","title":"Sharpness minimization algorithms do not only minimize sharpness to achieve better generalization","venue":null,"work_id":"53fd74e5-04e5-4cfc-a54b-e114cee184c1","year":2023},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.300499Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:4f256f99e69d1bf4b6e00b9530a545d58c4b48c3d338ef9516e9a2edf47b0e19","observation_id":"2b70935b-95f2-47ac-aba0-ee2e3e783a51","resolution":{"observed_at":"2026-08-11T20:57:20.593099Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.304408Z","title":"Model soups: averaging weights of multiple fine-tuned models improves accuracy without increasing inference time","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.304408Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:144d91cd1a108290517eb641e032642fc5ed0a4a952c3169f88e2fbe09f684f5","observation_id":"ac65624e-ec7d-4eeb-bef5-45f853a0ff3b","resolution":{"observed_at":"2026-08-11T20:57:20.304408Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.308393Z","title":"Towards a theoretical framework of out-of-distribution generalization","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.308393Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:a29372cc41c8195cffa211978c3e5959e1a1eb24b723e04ea859b142e18a793c","observation_id":"6f6632b6-a929-4e6b-96cd-d624a0d39065","resolution":{"observed_at":"2026-08-11T20:57:20.308393Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.555378Z","title":"Flatness-aware minimization for domain generalization","venue":null,"work_id":"bf415020-3217-4d66-9aaf-86ef4cae3f57","year":2023},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.312397Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:5dbe7878a5c71b12c75dd0251a9782a9776a196be451be4e5b7889a83a74a175","observation_id":"8b0c315d-a86e-432e-b05b-24a32b885225","resolution":{"observed_at":"2026-08-11T20:57:20.559775Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.541447Z","title":"On learning invariant representations for domain adaptation","venue":null,"work_id":"defcef37-91db-4048-80f9-8807817fea03","year":2019},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.316430Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:87bc6abfaed2542f6b90dad630f6a8a5107365fee7b80f2fa7537252b11578ce","observation_id":"2cbe782e-8c93-49e7-a568-2078acdf5488","resolution":{"observed_at":"2026-08-11T20:57:20.546149Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.527188Z","title":"Fundamental limits and tradeoffs in invariant representation learning","venue":null,"work_id":"8fc51d34-faf2-436a-8c5f-6d9c5b9162f0","year":2022},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.320780Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:2d472eab419c482d599ef6cc8d80d85c574175a7259383092ceb7de11cdde128","observation_id":"5b190aff-9409-4a45-a693-296ca2df157a","resolution":{"observed_at":"2026-08-11T20:57:20.532010Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.512868Z","title":"Gradual domain adaptation via gradient flow","venue":null,"work_id":"b6858c7f-68c9-4a4e-8a96-784661e85b6d","year":2024},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.324877Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:a98a4f8e98cd47476bf4e3201708d6a331f5eeac8e2ab4f49cbcd059a1014968","observation_id":"41f2b8f6-5858-42b1-bd7c-87f0cf08efbb","resolution":{"observed_at":"2026-08-11T20:57:20.517623Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.496172Z","title":"Towards robust out-of-distribution generalization bounds via sharpness","venue":null,"work_id":"802980fb-f57b-4c3f-89a4-4d1e158d3b16","year":2024},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.329093Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:197cd75ad1877172ead4f27eb2ccd175dd505da4ace6fc122ff5a12019da02f2","observation_id":"150b1a79-28e1-4ad4-927c-51c3d77728e8","resolution":{"observed_at":"2026-08-11T20:57:20.502650Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T20:57:20.333112Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-11T20:57:20.333112Z"},"links":{"citing_paper":"/paper/2412.05169"},"observation_digest":"sha256:b2adc0081a8302398e70ffa20c21aad2096de7d1c573159cb125547aacdbd299","observation_id":"5eb017a5-5d11-40ae-aec9-e70cdd9d1c5a","resolution":{"observed_at":"2026-08-11T20:57:20.333112Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2412.05169","last_updated":"2024-12-06T16:41:44Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-11T20:48:13.051196Z","submitted_at":"2024-12-06T16:41:44Z","title":"Towards Understanding the Role of Sharpness-Aware Minimization Algorithms for Out-of-Distribution Generalization"},"reference_resolution":{"displayed":47,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":15,"verified_exact":0,"verified_fuzzy":32},"total_outbound_references":47},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"thesis":"As of 15 August 2026, this Paper Citation Record lists 47 of 47 outbound references and 2 inbound Pith citation observations for arXiv:2412.05169."}