{"as_of":"2026-08-08T21:10:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:225d9996a7ae1a566024d4979b80802f9ae505c45dca69195db29774d2a97654","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":53,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":53,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":53,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":53,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T13:30:24.271122Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":18,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-07T13:30:24.271122Z","title":"Handbook of convergence theorems for (stochastic) gradient methods","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.21651","last_updated":"2025-05-27T18:25:21Z","snapshot_observed_at":"2026-08-07T13:23:29.160561Z","submitted_at":"2025-05-27T18:25:21Z","title":"AutoSGD: Automatic Learning Rate Selection for Stochastic Gradient Descent","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:24.271122Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2505.21651"},"observation_digest":"sha256:e9fef70c48085f11246fb69bed4a15f250f8a46de8290e7bb618f88b8830a100","observation_id":"7d99c6f7-0b53-42b8-a1aa-b86f26760ca8","resolution":{"observed_at":"2026-08-07T13:30:24.271122Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-07T13:06:23.807122Z","title":"Garrigos and R","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.22922","last_updated":"2025-05-28T22:51:43Z","snapshot_observed_at":"2026-08-07T23:35:02.219408Z","submitted_at":"2025-05-28T22:51:43Z","title":"Scalable Parameter and Memory Efficient Pretraining for LLM: Recent Algorithmic Advances and Benchmarking","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T13:06:23.807122Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2505.22922"},"observation_digest":"sha256:9ad650ae40fc4559227d70bb147d53897f6b35363a64dbd096cbb8fbd0ca6653","observation_id":"698d426a-8526-4e44-ba34-6ff5d0488ef5","resolution":{"observed_at":"2026-08-07T13:06:23.807122Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2505.23737","last_updated":"2026-07-28T01:40:23Z","snapshot_observed_at":"2026-08-07T12:36:23.369583Z","submitted_at":"2025-05-29T17:58:01Z","title":"On the Convergence Analysis of Muon","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-19T13:19:37.035526Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2505.23737"},"observation_digest":"sha256:d36497cc65d2fbdb33dd2e533e25e6f1763aa152d24f07a710b6a67b6044dd41","observation_id":"544e4bdd-49e0-4f5d-818c-39c4569ff296","resolution":{"observed_at":"2026-05-19T13:22:19.199856Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-07T12:45:18.827093Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.23737","last_updated":"2026-07-28T01:40:23Z","snapshot_observed_at":"2026-08-07T12:36:23.369583Z","submitted_at":"2025-05-29T17:58:01Z","title":"On the Convergence Analysis of Muon","version":3},"reference_index":2011,"source":"pdf_text","source_observed_at":"2026-08-07T12:45:18.827093Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2505.23737"},"observation_digest":"sha256:349b2203ad63fd796f342902f342428c2a239891cc8341548464683e282172b0","observation_id":"328ca2cf-adf4-4ebd-a7fc-9cf9278dd6a1","resolution":{"observed_at":"2026-08-07T12:45:18.827093Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-07T10:56:21.716800Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04015","last_updated":"2025-06-04T14:46:18Z","snapshot_observed_at":"2026-08-07T10:46:59.839970Z","submitted_at":"2025-06-04T14:46:18Z","title":"GORACS: Group-level Optimal Transport-guided Coreset Selection for LLM-based Recommender Systems","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T10:56:21.716800Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2506.04015"},"observation_digest":"sha256:a9acc492c6860dc089beb3d1a3fb79124e986258104fcad9e981464f781682a1","observation_id":"45e1c47b-d304-4719-9c7a-ae8810cde4e7","resolution":{"observed_at":"2026-08-07T10:56:21.716800Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-07T10:58:55.491727Z","title":"and Gower, R","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.04126","last_updated":"2025-06-04T16:17:25Z","snapshot_observed_at":"2026-08-08T15:52:38.460566Z","submitted_at":"2025-06-04T16:17:25Z","title":"Incremental Gradient Descent with Small Epoch Counts is Surprisingly Slow on Ill-Conditioned Problems","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T10:58:55.491727Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2506.04126"},"observation_digest":"sha256:5606b1a0d16b51806a9693ef07dbc2818b6af851a68b6c23779820b69782d8d0","observation_id":"199ce0c3-7b7b-4681-81cf-351d7edea8aa","resolution":{"observed_at":"2026-08-07T10:58:55.491727Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-07T06:10:19.881347Z","title":"and Gower, R","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.06194","last_updated":"2025-06-06T15:53:35Z","snapshot_observed_at":"2026-08-07T05:56:21.438028Z","submitted_at":"2025-06-06T15:53:35Z","title":"Transformative or Conservative? Conservation laws for ResNets and Transformers","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T06:10:19.881347Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2506.06194"},"observation_digest":"sha256:6776247550c92a923ad4dcfa4e76f9b34996ebf9b685355ab21ea38a29ee76d2","observation_id":"72447bda-f0b9-4b19-9b60-85ab3ce312c0","resolution":{"observed_at":"2026-08-07T06:10:19.881347Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-07T04:55:07.182515Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.09711","last_updated":"2025-07-01T05:59:07Z","snapshot_observed_at":"2026-08-07T04:39:58.599636Z","submitted_at":"2025-06-11T13:24:18Z","title":"Non-Euclidean dual gradient ascent for entropically regularized linear and semidefinite programming","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T04:55:07.182515Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2506.09711"},"observation_digest":"sha256:682b373547ee0fc9d54aa081714b053048a2074729c48327248ee81bf548c9dd","observation_id":"617cb15a-46d4-4730-9626-b7d35b9579a1","resolution":{"observed_at":"2026-08-07T04:55:07.182515Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-06T22:24:37.682404Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.21833","last_updated":"2026-07-10T19:11:43Z","snapshot_observed_at":"2026-08-06T22:15:47.651351Z","submitted_at":"2025-06-27T00:47:03Z","title":"Memory Savings at What Cost? A Study of Alternatives to Backpropagation","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T22:24:37.682404Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2506.21833"},"observation_digest":"sha256:5e362b06f4b3d1eed9e143849010711b0944d967e03397c0ce4803443b87b73e","observation_id":"c7ab788e-7c16-4494-83ce-888e2d083663","resolution":{"observed_at":"2026-08-06T22:24:37.682404Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-06T21:57:26.261239Z","title":"Garrigos and R.M","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.23303","last_updated":"2025-06-29T15:47:10Z","snapshot_observed_at":"2026-08-07T22:06:54.289394Z","submitted_at":"2025-06-29T15:47:10Z","title":"On the boundedness of the sequence generated by minibatch stochastic gradient descent","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T21:57:26.261239Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2506.23303"},"observation_digest":"sha256:00cb1d36ecaf4995cc3b62c8148ecd4b2ce51196041a0a4d9f467c59820cc2fe","observation_id":"b2449ce3-0f05-40e2-8874-1415fc310549","resolution":{"observed_at":"2026-08-06T21:57:26.261239Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-06T21:42:45.007133Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.23756","last_updated":"2025-06-30T11:56:51Z","snapshot_observed_at":"2026-08-08T10:46:25.456077Z","submitted_at":"2025-06-30T11:56:51Z","title":"Optimized methods for composite optimization: a reduction perspective","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T21:42:45.007133Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2506.23756"},"observation_digest":"sha256:c7e1250a68386ae84ab0b24327251a3b3509b26c37c505aa083d8d6cc67d14d5","observation_id":"900b19bc-3ea6-4693-8941-ab9d58e7d8f5","resolution":{"observed_at":"2026-08-06T21:42:45.007133Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-06T18:27:04.741729Z","title":"Handbook of convergence theorems for (stochastic) gradient methods","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.08518","last_updated":"2025-07-11T12:05:03Z","snapshot_observed_at":"2026-08-06T18:14:14.394746Z","submitted_at":"2025-07-11T12:05:03Z","title":"Data Depth as a Risk","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-06T18:27:04.741729Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2507.08518"},"observation_digest":"sha256:62a9afab614b6cfd2703fdf588b0aaf542f055e358851c12bf18afb0bb04f1cd","observation_id":"90e18804-29e9-42d1-b08e-38c0d2c82e50","resolution":{"observed_at":"2026-08-06T18:27:04.741729Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-06T16:19:00.127615Z","title":"and Gower, R","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.14122","last_updated":"2025-07-18T17:53:12Z","snapshot_observed_at":"2026-08-07T19:28:33.594363Z","submitted_at":"2025-07-18T17:53:12Z","title":"Last-Iterate Complexity of SGD for Convex and Smooth Stochastic Problems","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-06T16:19:00.127615Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2507.14122"},"observation_digest":"sha256:4ab5fc34b3fdbb8103bcce20578eae03c762613adcdc2924b9e53398a397b069","observation_id":"de23c43e-b2ed-4039-ba5d-c8c7ed40feac","resolution":{"observed_at":"2026-08-06T16:19:00.127615Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-06T15:48:51.835026Z","title":"Garrigos and R","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15424","last_updated":"2025-07-21T09:24:49Z","snapshot_observed_at":"2026-08-07T11:37:10.917662Z","submitted_at":"2025-07-21T09:24:49Z","title":"Stochastic Quantum Hamiltonian Descent","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T15:48:51.835026Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2507.15424"},"observation_digest":"sha256:8767fb5ffba6efc0ec0e4220ea7d3819ebbdd2ce2083be030ecfaff82e9ad9e8","observation_id":"744272d5-fc88-4222-9352-b81651dd1072","resolution":{"observed_at":"2026-08-06T15:48:51.835026Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-06T06:11:16.496830Z","title":"Handbook of convergence theorems for (stochastic) gradient methods,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.00775","last_updated":"2026-06-04T16:16:15Z","snapshot_observed_at":"2026-08-06T06:11:14.689302Z","submitted_at":"2025-08-01T16:56:42Z","title":"Learning to optimize with guarantees: a complete characterization of linearly convergent algorithms","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T06:11:16.496830Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2508.00775"},"observation_digest":"sha256:2048cb9d3be61ef3e475930d6a7c5556e82ff47fc496546476006466d767f5b8","observation_id":"49e269f7-cf5e-4578-b7e0-66f94975e58c","resolution":{"observed_at":"2026-08-06T06:11:16.496830Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2508.09103","last_updated":"2026-04-03T14:28:20Z","snapshot_observed_at":"2026-07-06T22:11:52.522787Z","submitted_at":"2025-08-12T17:31:13Z","title":"Constrained free energy minimization for the design of thermal states and stabilizer thermodynamic systems","version":3},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-05-18T23:12:47.065365Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2508.09103"},"observation_digest":"sha256:56d59a84d0a8531aed9aec3ba2e2af977a577c65aba73acd28d8036d4fd014cb","observation_id":"0e045065-3d08-4c26-bac7-50393756b7cb","resolution":{"observed_at":"2026-05-18T23:12:52.857953Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T14:49:00.254125Z","title":"Handbook of convergence theorems for (stochastic) gradient methods","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.21022","last_updated":"2026-06-08T19:22:55Z","snapshot_observed_at":"2026-08-05T14:48:57.900801Z","submitted_at":"2025-08-28T17:24:59Z","title":"A Sketch-and-Project Analysis of Subsampled Natural Gradient Algorithms","version":3},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-05T14:49:00.254125Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2508.21022"},"observation_digest":"sha256:77a9d0fbfa4aa9fd5a2abda54da632f865fe563d49243006df1f038578bdd98b","observation_id":"b4394677-3d40-4776-9022-79941a13b3bd","resolution":{"observed_at":"2026-08-05T14:49:00.254125Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2509.02912","last_updated":"2026-04-05T02:48:52Z","snapshot_observed_at":"2026-08-03T04:49:12.991251Z","submitted_at":"2025-09-03T00:39:21Z","title":"Stochastic versus Deterministic in Stochastic Gradient Descent","version":3},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-18T20:15:50.901570Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2509.02912"},"observation_digest":"sha256:8c3d778db9fc3b0ba4f29ed1f22b61360536f46f994d6d8a9896a620231c4264","observation_id":"a923754c-d293-456b-b6fc-38c0d824e9d0","resolution":{"observed_at":"2026-05-18T20:16:50.172360Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T11:23:47.017667Z","title":"Handbook of convergence theorems for (stochastic) gradient methods","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.02970","last_updated":"2026-05-29T09:37:19Z","snapshot_observed_at":"2026-08-05T11:23:39.748947Z","submitted_at":"2025-09-03T03:14:58Z","title":"Delayed Momentum Aggregation: Communication-efficient Byzantine-robust Federated Learning with Partial Participation","version":3},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-05T11:23:47.017667Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2509.02970"},"observation_digest":"sha256:01054ea0afec718135d97fcb3a26b9ce555434a04782cdd1ef50dc8a100c2273","observation_id":"ec767131-6bda-4a23-a6bb-eae9d86ff5bf","resolution":{"observed_at":"2026-08-05T11:23:47.017667Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-04T19:11:48.753617Z","title":"Handbook of convergence theorems for (stochastic) gradient methods","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.09485","last_updated":"2025-09-12T01:27:15Z","snapshot_observed_at":"2026-08-07T14:36:54.609466Z","submitted_at":"2025-09-11T14:17:04Z","title":"Balancing Utility and Privacy: Dynamically Private SGD with Random Projection","version":2},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-04T19:11:48.753617Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2509.09485"},"observation_digest":"sha256:c6ab26bc3d1f451a3d55e08b4c80f8f09305a4bbd564b35985d5bf08b12e0f14","observation_id":"f2710097-b845-44d4-8a0e-cb8e6ffa8dd0","resolution":{"observed_at":"2026-08-04T19:11:48.753617Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2510.04686","last_updated":"2026-04-21T09:01:24Z","snapshot_observed_at":"2026-08-02T11:55:04.825901Z","submitted_at":"2025-10-06T10:56:41Z","title":"How does the optimizer implicitly bias the model merging loss landscape?","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-18T09:49:41.550981Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2510.04686"},"observation_digest":"sha256:3af6039843f33bf7e04128ce759a67dbfc6d8dcc192a49a4d3bbc79bc38fe3c2","observation_id":"55331fa9-9d44-4445-b665-19e2813685fc","resolution":{"observed_at":"2026-05-18T09:51:13.449239Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2510.07922","last_updated":"2026-05-01T03:08:45Z","snapshot_observed_at":"2026-07-06T22:32:07.945004Z","submitted_at":"2025-10-09T08:16:32Z","title":"SketchGuard: Scaling Byzantine-Robust Decentralized Federated Learning via Sketch-Based Screening","version":4},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-18T08:55:14.289595Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2510.07922"},"observation_digest":"sha256:78ff8a9f8d54a78f19b7a28dad4c68addad4d9ea11656a969dec00f5240b111d","observation_id":"5e2448c7-14a7-4b21-9eef-2b997c74a859","resolution":{"observed_at":"2026-05-18T08:56:08.513607Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-03T20:56:10.838262Z","title":"11235,arXiv:2301.11235","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2511.17976","last_updated":"2026-07-30T13:39:06Z","snapshot_observed_at":"2026-08-06T07:55:10.811529Z","submitted_at":"2025-11-22T08:35:40Z","title":"Accelerated optimization of measured relative entropies","version":2},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-03T20:56:10.838262Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2511.17976"},"observation_digest":"sha256:ccaefe33a1a0af589dacfb2039988b318cbc43999b76f774983ab6c7b3bc1c39","observation_id":"bfa6aaef-a7dc-4667-9056-3808783bb403","resolution":{"observed_at":"2026-08-03T20:56:10.838262Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-03T20:34:07.426307Z","title":"Behrooz Ghorbani, Shankar Krishnan, and Ying Xiao","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2511.19716","last_updated":"2026-08-04T17:24:58Z","snapshot_observed_at":"2026-08-07T23:11:20.886221Z","submitted_at":"2025-11-24T21:24:40Z","title":"Design Criteria for SGD Preconditioners: Local Conditioning, Noise Floors, and Basin Stability","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-03T20:34:07.426307Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2511.19716"},"observation_digest":"sha256:1e61073e878e6590960c4b8ad9699bd00ae1ef19242a9764fbd637a917c9aae6","observation_id":"13e9b960-fb01-4740-8d43-2600f0aeb5c0","resolution":{"observed_at":"2026-08-03T20:34:07.426307Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2512.18248","last_updated":"2026-05-11T05:28:48Z","snapshot_observed_at":"2026-07-30T07:13:55.806066Z","submitted_at":"2025-12-20T07:20:54Z","title":"On the Convergence Rate of LoRA Gradient Descent","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-16T20:25:58.250270Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2512.18248"},"observation_digest":"sha256:458928679885795c872912915e469d23bf9c55a2389c807dfa0cfd1924ae392c","observation_id":"2b744de1-8550-4f23-a8ed-d8d8d8f522f6","resolution":{"observed_at":"2026-05-16T20:28:24.084563Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-03T10:27:47.638277Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2601.10222","last_updated":"2026-06-02T13:23:42Z","snapshot_observed_at":"2026-08-03T21:23:39.674231Z","submitted_at":"2026-01-15T09:36:15Z","title":"Introduction to optimization methods for training SciML models","version":2},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-03T10:27:47.638277Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2601.10222"},"observation_digest":"sha256:d4b14f9067521b4fcec2105e70f67109b89a94e1751eeb6a427388df0c1bab6f","observation_id":"66b29a9f-f37e-48db-803f-1a75f369b3f7","resolution":{"observed_at":"2026-08-03T10:27:47.638277Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-03T06:50:09.070523Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2601.22003","last_updated":"2026-06-10T19:28:24Z","snapshot_observed_at":"2026-08-07T03:10:57.578271Z","submitted_at":"2026-01-29T17:13:25Z","title":"Efficient Stochastic Optimisation via Sequential Monte Carlo","version":2},"reference_index":2003,"source":"pdf_text","source_observed_at":"2026-08-03T06:50:09.070523Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2601.22003"},"observation_digest":"sha256:7aba573131c782922f26d3cd4363037578c4e22ffd6b27823c7c2fb2fd375a1d","observation_id":"f03df642-e748-46f7-b925-342af21cc48a","resolution":{"observed_at":"2026-08-03T06:50:09.070523Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-03T02:53:14.060774Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2602.09842","last_updated":"2026-05-26T14:40:02Z","snapshot_observed_at":"2026-08-06T04:10:02.481582Z","submitted_at":"2026-02-10T14:46:14Z","title":"Step-Size Stability in Stochastic Optimization: A Theoretical Perspective","version":2},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-03T02:53:14.060774Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2602.09842"},"observation_digest":"sha256:f7a817e01aeb141004981b37aef66f5a5f617b8a19fcdc9760bd9dcc1a8a9ed5","observation_id":"56c3f6fe-b240-47c3-a073-bbb6020c755e","resolution":{"observed_at":"2026-08-03T02:53:14.060774Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-03T01:09:06.485628Z","title":"Handbook of convergence theorems for (stochastic) gradient methods","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2602.10691","last_updated":"2026-07-15T07:42:51Z","snapshot_observed_at":"2026-08-08T13:45:32.808902Z","submitted_at":"2026-02-11T09:47:52Z","title":"Convergence Rates for Distribution Matching with Sliced Optimal Transport","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-03T01:09:06.485628Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2602.10691"},"observation_digest":"sha256:431b34f08f7dc7f05af84bec28c628bdc0a0ff2b3f6d4f2bf68dc7b75496fcce","observation_id":"d8442d59-db12-468b-bef2-788f64450555","resolution":{"observed_at":"2026-08-03T01:09:06.485628Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2602.18718","last_updated":"2026-05-19T13:16:25Z","snapshot_observed_at":"2026-07-29T19:02:40.084737Z","submitted_at":"2026-02-21T04:52:53Z","title":"Stochastic Gradient Variational Inference with Price's Gradient Estimator from Bures-Wasserstein to Parameter Space","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-21T12:14:32.642347Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2602.18718"},"observation_digest":"sha256:f84103d5964b6b0092cd45b9e159f9337f36ecbe750875acc1fb1fbe850513f6","observation_id":"66eab807-3379-40e7-b92c-3c90b645fbc3","resolution":{"observed_at":"2026-05-21T12:15:06.694304Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2603.10067","last_updated":"2026-05-21T20:35:00Z","snapshot_observed_at":"2026-07-06T22:48:35.345421Z","submitted_at":"2026-03-10T02:12:24Z","title":"HTMuon: Improving Muon via Heavy-Tailed Spectral Correction","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-25T06:50:29.893126Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2603.10067"},"observation_digest":"sha256:c03ffeb27ba991e9e82c875c02427e4a2d0b9f5cf6992b3aedc89cf12826cb72","observation_id":"8f298c13-e139-4fd0-9323-925cda206988","resolution":{"observed_at":"2026-05-25T06:55:26.291179Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2604.06909","last_updated":"2026-04-08T10:08:02Z","snapshot_observed_at":"2026-07-06T22:55:15.791334Z","submitted_at":"2026-04-08T10:08:02Z","title":"Mini-Batch Stochastic Krasnosel'ski\\u\\i-Mann Algorithm for Nonexpansive Fixed Point Problems","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-10T17:51:10.587034Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2604.06909"},"observation_digest":"sha256:c5d70de2a9dd25313d80bed1fe075ca42391513a32303b9bdaa389e7882a55a7","observation_id":"6d3f316c-663f-4dcc-9067-96d519977fa4","resolution":{"observed_at":"2026-05-11T06:00:58.487791Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2604.22581","last_updated":"2026-05-11T11:31:00Z","snapshot_observed_at":"2026-08-03T10:16:56.839436Z","submitted_at":"2026-04-24T14:13:29Z","title":"Stochastic Krasnoselskii-Mann Iterations: Convergence without Uniformly Bounded Variance","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-08T11:07:07.575910Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2604.22581"},"observation_digest":"sha256:411d9f67ce093f7ae688d29057d38ecb49c48cfe705fe5646682ab3598133ce9","observation_id":"2d3f6564-9dad-4d31-bd73-fcc9fe6d7d8d","resolution":{"observed_at":"2026-05-11T19:46:07.419003Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2604.22581","last_updated":"2026-05-11T11:31:00Z","snapshot_observed_at":"2026-08-03T10:16:56.839436Z","submitted_at":"2026-04-24T14:13:29Z","title":"Stochastic Krasnoselskii-Mann Iterations: Convergence without Uniformly Bounded Variance","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-12T02:52:40.976730Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2604.22581"},"observation_digest":"sha256:ccc2ca24884ec71fa443a7f84178534a590f9e9e29eab73d15d392315e471056","observation_id":"cdceee8d-d749-490f-b1db-30d0a30feecd","resolution":{"observed_at":"2026-05-12T07:26:30.085423Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2604.23017","last_updated":"2026-05-25T07:26:45Z","snapshot_observed_at":"2026-07-06T23:09:19.041332Z","submitted_at":"2026-04-24T21:08:39Z","title":"Complex Stochastic Gradient Descent and Directional Bias in Reproducing Kernel Hilbert Spaces","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-08T12:13:40.922801Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2604.23017"},"observation_digest":"sha256:6dd191ff9b792b849ebfd06b5ec1686c6b0d53a123bb8f72ca40121aac686d1d","observation_id":"8a22ac1c-b7b5-47db-9133-b7641b8c53f7","resolution":{"observed_at":"2026-05-11T19:21:07.761338Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2604.23017","last_updated":"2026-05-25T07:26:45Z","snapshot_observed_at":"2026-07-06T23:09:19.041332Z","submitted_at":"2026-04-24T21:08:39Z","title":"Complex Stochastic Gradient Descent and Directional Bias in Reproducing Kernel Hilbert Spaces","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-07-04T16:26:13.991595Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2604.23017"},"observation_digest":"sha256:83c6e238c830af715f0af2a6d5b8703ef0eb3efab564216d4c5902243759abd5","observation_id":"16a41696-132b-484e-ae80-4b9b3bf8ea75","resolution":{"observed_at":"2026-07-04T16:29:56.438417Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2604.25613","last_updated":"2026-04-28T13:23:15Z","snapshot_observed_at":"2026-08-05T08:31:13.718665Z","submitted_at":"2026-04-28T13:23:15Z","title":"One Coordinate at a Time: Convergence Guarantees for Rotosolve in Variational Quantum Algorithms","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-07T16:58:42.533922Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2604.25613"},"observation_digest":"sha256:3033a5287fd4bf3d6ac6eda5d7c946e0d6b8b45396bd5d91676d3863cc8585c7","observation_id":"ebff0ad4-eece-49ce-aace-0e1c5d2b1533","resolution":{"observed_at":"2026-05-11T23:26:16.987341Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2605.03313","last_updated":"2026-05-05T02:57:46Z","snapshot_observed_at":"2026-07-06T23:16:15.509235Z","submitted_at":"2026-05-05T02:57:46Z","title":"Distributed Learning with Adversarial Gradient Perturbations","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-07T17:32:12.469438Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2605.03313"},"observation_digest":"sha256:bac56a20f6573fde084685a362489713113b1514fc74e93d02955e53b39d3bc1","observation_id":"7974660c-7ce3-4e31-b3ad-f9f0b9da8f55","resolution":{"observed_at":"2026-05-11T23:16:38.024047Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2605.14663","last_updated":"2026-05-14T10:18:21Z","snapshot_observed_at":"2026-08-01T08:58:53.907936Z","submitted_at":"2026-05-14T10:18:21Z","title":"Optimal Asymptotic Rates for (Stochastic) Gradient Descent under the Local PL-Condition: A Geometric Approach","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-30T20:25:40.947453Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2605.14663"},"observation_digest":"sha256:fa16ba4786c091d5c315d500a9f3cb7c72ab71a159c4606acaf2748d4be255fc","observation_id":"ca3ee808-83d6-46a5-a11c-b9b36d489284","resolution":{"observed_at":"2026-07-01T14:35:47.508979Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2605.16017","last_updated":"2026-05-15T14:50:39Z","snapshot_observed_at":"2026-08-03T19:59:12.010055Z","submitted_at":"2026-05-15T14:50:39Z","title":"Accelerated Gradient Descent for Faster Convergence with Minimal Overhead","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-05-20T21:19:33.691048Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2605.16017"},"observation_digest":"sha256:5b4d98ef1740548a104dedb533d5a2e2955b2351add0114af263fd7704cee450","observation_id":"632a77a5-7b93-4314-8d36-786b9cccb42d","resolution":{"observed_at":"2026-05-20T21:23:44.646223Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2605.18675","last_updated":"2026-05-18T17:15:27Z","snapshot_observed_at":"2026-07-06T23:29:29.105126Z","submitted_at":"2026-05-18T17:15:27Z","title":"COOPO: Cyclic Offline-Online Policy Optimization Algorithm","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-05-20T13:11:16.568415Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2605.18675"},"observation_digest":"sha256:a0cfdd9d07ded9e29d62a90d2c4e9da6ffeecbc7da307a6c8d8174420297f0f9","observation_id":"3dc27a8b-1550-4f10-9b69-21f2f25cf622","resolution":{"observed_at":"2026-05-20T13:13:18.011051Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2605.19291","last_updated":"2026-05-19T03:10:33Z","snapshot_observed_at":"2026-07-06T23:30:02.207108Z","submitted_at":"2026-05-19T03:10:33Z","title":"Factor Augmented High-Dimensional SGD","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-05-20T03:31:04.529612Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2605.19291"},"observation_digest":"sha256:02c97be0e949eab7fc7d5176799fc28a694d87da8f93b0e6a1523e16957977d7","observation_id":"82d6e872-7d86-45fa-9e76-7572671ffc82","resolution":{"observed_at":"2026-05-20T03:33:01.622592Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2605.25034","last_updated":"2026-05-24T12:22:00Z","snapshot_observed_at":"2026-08-03T22:01:17.750102Z","submitted_at":"2026-05-24T12:22:00Z","title":"Randomized conjugate gradient least squares","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-29T23:49:53.122613Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2605.25034"},"observation_digest":"sha256:a6ede1beae13781685d94fa69acdc171801926afeb4a2fc6dad545ba9ad8d57b","observation_id":"2bd3db27-9633-4171-a764-620190520e56","resolution":{"observed_at":"2026-06-29T23:54:03.385731Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2605.25499","last_updated":"2026-05-25T06:58:51Z","snapshot_observed_at":"2026-08-04T11:46:17.852874Z","submitted_at":"2026-05-25T06:58:51Z","title":"Accelerated Dynamic Importance Weighting with Versatile Divergence-Minimizing Estimators","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-06-29T22:48:00.058391Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2605.25499"},"observation_digest":"sha256:a11ca16448e84f66d192a8e876f2184807c4bfe61f018ca3ba15ebde97d34d8f","observation_id":"1b0e3681-920a-4b03-9701-11d6239cee02","resolution":{"observed_at":"2026-06-29T22:54:01.204163Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2605.29304","last_updated":"2026-05-28T03:36:19Z","snapshot_observed_at":"2026-08-08T05:45:41.987512Z","submitted_at":"2026-05-28T03:36:19Z","title":"On subspace-constrained preconditioning for randomized iterative methods","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-06-29T06:21:37.788799Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2605.29304"},"observation_digest":"sha256:0eb8a39032a3f448cefb65a5c36d2a041c754185fa982ef6fa547f15aa1efe19","observation_id":"6695fc90-989c-437a-baa7-8fb5559334dc","resolution":{"observed_at":"2026-06-29T06:23:08.773250Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2606.02365","last_updated":"2026-06-01T15:13:28Z","snapshot_observed_at":"2026-08-02T10:19:39.354615Z","submitted_at":"2026-06-01T15:13:28Z","title":"FOAM: Frequency and Operator Error-Based Adaptive Damping Method for Reducing Staleness-Oriented Error for Shampoo","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-06-28T15:49:18.685160Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2606.02365"},"observation_digest":"sha256:f6539d67a531aa2678d8a7df3c4da1ad45f09f841e4614fb0f0e9ca0c97e93e9","observation_id":"1a808a56-e34a-4b8b-8ef7-2bbf15fe32fd","resolution":{"observed_at":"2026-07-01T22:06:16.042460Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2606.27171","last_updated":"2026-06-25T15:39:19Z","snapshot_observed_at":"2026-08-05T08:11:17.186534Z","submitted_at":"2026-06-25T15:39:19Z","title":"Stochastic Gradient Optimization with Model-Assisted Sampling","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-26T05:06:04.159392Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2606.27171"},"observation_digest":"sha256:6b5514de7f6bf23c5e08e3d762829b3246af0675fa6e579632539c9535e42f63","observation_id":"d1c42a26-9462-4ae7-99c2-56c373c2af72","resolution":{"observed_at":"2026-07-04T13:39:50.339832Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2606.28973","last_updated":"2026-06-27T15:25:34Z","snapshot_observed_at":"2026-07-07T00:02:57.622708Z","submitted_at":"2026-06-27T15:25:34Z","title":"Sharp $O(1/k)$ convergence rate for the Sinkhorn algorithm via a local analysis","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-06-30T08:38:03.985543Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2606.28973"},"observation_digest":"sha256:a46208b8927c79df698fee802fefbabec573805e02d09e674f3be239e06eac33","observation_id":"dbb51a75-21fa-4ffc-bd86-09dffbcc52d1","resolution":{"observed_at":"2026-06-30T08:44:27.699688Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2606.29593","last_updated":"2026-06-28T20:27:52Z","snapshot_observed_at":"2026-08-07T04:15:53.433457Z","submitted_at":"2026-06-28T20:27:52Z","title":"How AI settled the complexity of the oldest SGD algorithm","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-06-30T07:19:07.518524Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2606.29593"},"observation_digest":"sha256:0fb16bed5963b74551bf997e2767aca3f0b7901e2af77eaf665e25dd28fe62fe","observation_id":"9bd78cbd-250d-4bf3-9ae1-c4686a3402a6","resolution":{"observed_at":"2026-06-30T07:24:21.394540Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2606.30310","last_updated":"2026-06-29T13:53:29Z","snapshot_observed_at":"2026-08-07T11:22:53.515369Z","submitted_at":"2026-06-29T13:53:29Z","title":"Highly Data Parallelizable Estimation of the Sliced-Wasserstein Distance Using Cumulative Distribution Functions","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-06-30T04:05:59.704550Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2606.30310"},"observation_digest":"sha256:61a6e6cf6172ee8fea4ed6be26a2e29fd4bfdcdb39258c39e61b91afb47e70ae","observation_id":"95dd39ea-cd2d-4dbf-a2e3-6cbaec5d2b4f","resolution":{"observed_at":"2026-06-30T04:14:18.916838Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2606.32005","last_updated":"2026-06-30T17:38:22Z","snapshot_observed_at":"2026-08-03T17:58:37.538515Z","submitted_at":"2026-06-30T17:38:22Z","title":"Random Reshuffling Dominates Stochastic Gradient Descent","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-07-01T03:50:15.120009Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2606.32005"},"observation_digest":"sha256:9bd11eae344567b2b4154aa77458f11859665a48b77b29e861a9997b14081ecc","observation_id":"ba7afe21-f279-4a8a-93f5-8f4ed6f3413b","resolution":{"observed_at":"2026-07-01T11:55:42.295919Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":"2301.11235","doi":"10.48550/arxiv.2301.11235","metadata_source":"arxiv_reference","pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Handbook of convergence theorems for (stochastic) gradient methods.arXiv preprint arXiv:2301.11235","venue":"arXiv (Cornell University)","work_id":"dc7b33e2-4ddd-47b1-ab42-c2de5a7bbff6","year":2024},"citing_paper":{"arxiv_id":"2607.00665","last_updated":"2026-07-01T09:13:24Z","snapshot_observed_at":"2026-08-03T18:18:39.283181Z","submitted_at":"2026-07-01T09:13:24Z","title":"Effective dynamics of the Sinkhorn algorithm in the regime of low entropy regularization","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-07-02T08:05:55.418491Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2607.00665"},"observation_digest":"sha256:085330832c1b5f99242bbe0c3b5a3c3439c1dcf81b4ad020fe7743c53a10b00f","observation_id":"9b8f8eed-f8a1-4468-a5e1-5111f3cc020d","resolution":{"observed_at":"2026-07-02T08:06:47.591600Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11235","snapshot_observed_at":"2026-07-11T20:46:05.467029Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.04233","last_updated":"2026-07-05T11:11:40Z","snapshot_observed_at":"2026-08-06T18:57:36.945231Z","submitted_at":"2026-07-05T11:11:40Z","title":"Unified convergence analysis for gradient descent optimization methods in the training of deep neural networks","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-07-11T20:46:05.467029Z"},"links":{"cited_paper":"/paper/2301.11235","citing_paper":"/paper/2607.04233"},"observation_digest":"sha256:9b12dba411d56fe9722d625eab26c5ed379c3eb9e4914bb5ce136db9c94c1ce2","observation_id":"d8bccd32-aa24-4f55-85b3-77628c9dc6b1","resolution":{"observed_at":"2026-07-11T20:46:05.467029Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2301.11235/citation-record","integrity":"/paper/2301.11235/integrity","json":"/paper/2301.11235/citation-record.json","paper":"/paper/2301.11235"},"outbound":[],"paper":{"arxiv_id":"2301.11235","last_updated":"2024-03-09T13:28:29Z","latest_version":3,"primary_category":"math.OC","snapshot_observed_at":"2026-07-06T14:44:56.245065Z","submitted_at":"2023-01-26T17:18:36Z","title":"Handbook of Convergence Theorems for (Stochastic) Gradient Methods"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 53 inbound Pith citation observations for arXiv:2301.11235."}