{"as_of":"2026-08-08T12:09:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:83c5dc63da36cd7378d99b3b5ea441d2d1e9422eec1ef4a9fca31d85d9a88d68","coverage":[{"denominator":26,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":26,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-02T22:19:25.000465Z","state":"measured"},{"denominator":26,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":26,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2602.17743/citation-record","integrity":"/paper/2602.17743/integrity","json":"/paper/2602.17743/citation-record.json","paper":"/paper/2602.17743"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T22:19:22.145224Z","title":"Language models are few-shot learners.Advances in neural information processing systems, 33:1877–1901, 2020","venue":null,"work_id":null,"year":1901},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:22.145224Z"},"links":{"citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:9acedc8261a64a3cedaf925369779c2fc395c969753d24f91242a90517b75015","observation_id":"e43a5379-88f1-4767-96c9-ec5a1bee1ced","resolution":{"observed_at":"2026-08-02T22:19:22.145224Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2202.12837","last_updated":"2022-10-20T14:04:10Z","snapshot_observed_at":"2026-08-01T15:23:39.998892Z","submitted_at":"2022-02-25T17:25:19Z","title":"Rethinking the Role of Demonstrations: What Makes In-Context Learning Work?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2202.12837","snapshot_observed_at":"2026-08-02T22:19:22.230499Z","title":"Rethinking the role of demonstrations: What makes in-context learning work?arXiv preprint arXiv:2202.12837, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:22.230499Z"},"links":{"cited_paper":"/paper/2202.12837","citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:5ead59c379fc65f957cef6822240f80c26e57418efb87aa4a4f253488bae37e0","observation_id":"df79976d-786c-4099-9f57-a49e1a3d8b3d","resolution":{"observed_at":"2026-08-02T22:19:22.230499Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2111.02080","last_updated":"2022-07-21T07:44:13Z","snapshot_observed_at":"2026-07-30T03:45:35.558793Z","submitted_at":"2021-11-03T09:12:33Z","title":"An Explanation of In-context Learning as Implicit Bayesian Inference","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2111.02080","snapshot_observed_at":"2026-08-02T22:19:22.307467Z","title":"An explanation of in-context learning as implicit bayesian inference.arXiv preprint arXiv:2111.02080, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:22.307467Z"},"links":{"cited_paper":"/paper/2111.02080","citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:789a1340a9747f2d35e68cb4ee0b71f5c82784ca21c06b966f5723fca623b967","observation_id":"3024d273-a56e-45dc-a911-9098cffffad5","resolution":{"observed_at":"2026-08-02T22:19:22.307467Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T22:19:22.381339Z","title":"In-context learning is provably bayesian inference: A generalization theory for meta-learning.arXiv preprint arXiv:2510.10981, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:22.381339Z"},"links":{"citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:ac06bbb124249b28869059cbd981d0480ad1b82cc42693fe281f3e1dc1c48a68","observation_id":"378fb731-d62e-46c1-9f5f-614ee7b68720","resolution":{"observed_at":"2026-08-02T22:19:22.381339Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T22:19:22.482123Z","title":"Transformers learn in-context by gradient descent","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:22.482123Z"},"links":{"citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:50605dda8dee776b62bff3310e99deaecda59a89a10df8e2ad309ad68852f909","observation_id":"ad3919cb-52e5-4da8-aebc-177ed9aeacad","resolution":{"observed_at":"2026-08-02T22:19:22.482123Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T22:19:22.579367Z","title":"Transformers learn to implement preconditioned gradient descent for in-context learning.Advances in Neural Information Processing Systems, 36:45614–45650, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:22.579367Z"},"links":{"citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:cfbd5080d89fde75641633dc725058c2ab82b5f6217078ed02337729e7e21d48","observation_id":"d46aef11-b7c3-47bf-b0d0-6ce4e92d58c5","resolution":{"observed_at":"2026-08-02T22:19:22.579367Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.04295","last_updated":"2024-08-30T11:57:47Z","snapshot_observed_at":"2026-08-04T23:34:13.332065Z","submitted_at":"2024-07-05T06:57:30Z","title":"Jailbreak Attacks and Defenses Against Large Language Models: A Survey","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.04295","snapshot_observed_at":"2026-08-02T22:19:22.686038Z","title":"Jailbreak attacks and defenses against large language models: A survey.arXiv preprint arXiv:2407.04295, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:22.686038Z"},"links":{"cited_paper":"/paper/2407.04295","citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:96595ddcb42b9106423859a1cf45e423fc74a5610ee78ed6244e9f4a93260525","observation_id":"c769a4bc-d07b-442c-908b-5837fa51655b","resolution":{"observed_at":"2026-08-02T22:19:22.686038Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T22:19:22.793582Z","title":"Jailbroken: How does llm safety training fail?Advances in Neural Information Processing Systems, 36:80079–80110, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:22.793582Z"},"links":{"citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:24f62a9c28ae0bfd5e6fb77659cf603560d1eeaafe13fcdcb86e273c973be317","observation_id":"bd50ed31-c6ea-4eef-a5bf-343b9adbf7a5","resolution":{"observed_at":"2026-08-02T22:19:22.793582Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.15043","last_updated":"2023-12-20T20:48:57Z","snapshot_observed_at":"2026-07-06T15:59:23.019044Z","submitted_at":"2023-07-27T17:49:12Z","title":"Universal and Transferable Adversarial Attacks on Aligned Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.15043","snapshot_observed_at":"2026-08-02T22:19:22.870827Z","title":"Universal and transferable adversarial attacks on aligned language models.arXiv preprint arXiv:2307.15043, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:22.870827Z"},"links":{"cited_paper":"/paper/2307.15043","citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:63ff2d5e5768df09b98ee00e24b8e3841113e94553ade59484f84c2d78d19db2","observation_id":"e8dddf92-1504-4676-8c36-e273aede687c","resolution":{"observed_at":"2026-08-02T22:19:22.870827Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2510.23254","last_updated":"2026-05-07T10:47:24Z","snapshot_observed_at":"2026-07-06T22:34:08.878520Z","submitted_at":"2025-10-27T12:16:49Z","title":"Optimal In-context Adaptivity and Distributional Robustness of Transformers","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2510.23254","snapshot_observed_at":"2026-08-02T22:19:22.958702Z","title":"Provable test-time adaptivity and distributional robustness of in-context learning.arXiv preprint arXiv:2510.23254, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:22.958702Z"},"links":{"cited_paper":"/paper/2510.23254","citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:6480600b7a6ecc05ad197a5a25d7ed904ebf231bbb5c00a82fab2ff2d5c46e13","observation_id":"554dd639-fa1a-4a66-b734-2a7c586f0e60","resolution":{"observed_at":"2026-08-02T22:19:22.958702Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1710.10571","last_updated":"2020-05-01T07:29:34Z","snapshot_observed_at":"2026-08-07T21:20:40.257232Z","submitted_at":"2017-10-29T07:27:57Z","title":"Certifying Some Distributional Robustness with Principled Adversarial Training","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1710.10571","snapshot_observed_at":"2026-08-02T22:19:23.064805Z","title":"Certifying some distributional robustness with principled adversarial training.arXiv preprint arXiv:1710.10571, 2017","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:23.064805Z"},"links":{"cited_paper":"/paper/1710.10571","citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:78eaf1c6a09bdb3a66b26950a873d910f8ce6aa5502c56d1e4b5efbdacabf031","observation_id":"46cc1d35-cf3f-4d3a-a030-5919ea97bc5e","resolution":{"observed_at":"2026-08-02T22:19:23.064805Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T22:19:23.163869Z","title":"A distributionally robust perspective on uncertainty quantification and chance constrained programming.Mathematical Programming, 151(1):35–62, 2015","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:23.163869Z"},"links":{"citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:191992bc562f092c021bf79e00a6259c2ef058027b0b7a4d75e4dec4f5b522a2","observation_id":"fc867cc7-3193-4070-9cbc-e21070571ed0","resolution":{"observed_at":"2026-08-02T22:19:23.163869Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T22:19:23.274566Z","title":"What can transformers learn in-context? a case study of simple function classes.Advances in neural information processing systems, 35:30583–30598, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:23.274566Z"},"links":{"citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:8164fdd57fc955ff726a6d9fdfd5a723efe3e2dbb7b5bca6f62c008aabc15c64","observation_id":"59a7d9f1-8e4b-4a14-a65d-fd2cf51ed0a8","resolution":{"observed_at":"2026-08-02T22:19:23.274566Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T22:19:23.376413Z","title":"Transformers as statisticians: Provable in-context learning with in-context algorithm selection.Advances in neural information processing systems, 36:57125–57211, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:23.376413Z"},"links":{"citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:c58b1b3ca02a83d580205db20ec6bdda96719c9ade910d3c100de96983b47223","observation_id":"42928dde-37a8-4c4d-b1e8-35309ffc8c0d","resolution":{"observed_at":"2026-08-02T22:19:23.376413Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T22:19:23.452948Z","title":"Transformers as algorithms: Generalization and stability in in-context learning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:23.452948Z"},"links":{"citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:27322a25a83289e85cb9f6f504fcc66e0bdb89d9d7fbc6c6e40dc4d54c4321bd","observation_id":"21fe781d-d9c1-428a-9a7c-70864dc70904","resolution":{"observed_at":"2026-08-02T22:19:23.452948Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T22:19:23.526724Z","title":"Pretraining task diversity and the emergence of non-bayesian in-context learning for regression.Advances in neural information processing systems, 36:14228– 14246, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:23.526724Z"},"links":{"citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:ef22f901704aa1f977a1bdce66edba0fe94527a2a6b8bd692f7dd7a42dc2ec65","observation_id":"2e797193-9bca-4dee-a27f-c8e57dd96615","resolution":{"observed_at":"2026-08-02T22:19:23.526724Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T22:19:23.617496Z","title":"Piecewise linear regression via a difference of convex functions","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:23.617496Z"},"links":{"citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:ba72cc42c36116eb9218f766835b4e3eabddf6574a396aee5ae25a933963c084","observation_id":"b37ec43a-b3c4-48f8-92cf-e6c039b372bf","resolution":{"observed_at":"2026-08-02T22:19:23.617496Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T22:19:23.725937Z","title":"Finite-sample analysis of m-estimators using self-concordance.Electronic Journal of Statistics, 15:326–391, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:23.725937Z"},"links":{"citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:be8a2378f0fb2f21f54530571f46921ea1ba23b96701d9eecc4b8b305b1bb873","observation_id":"ca0ddc94-32f6-4b25-9546-cbc9bd32219e","resolution":{"observed_at":"2026-08-02T22:19:23.725937Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.15661","last_updated":"2023-05-17T21:08:32Z","snapshot_observed_at":"2026-08-07T19:33:04.292595Z","submitted_at":"2022-11-28T18:59:51Z","title":"What learning algorithm is in-context learning? Investigations with linear models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.15661","snapshot_observed_at":"2026-08-02T22:19:23.802187Z","title":"What learning algorithm is in-context learning? investigations with linear models.arXiv preprint arXiv:2211.15661, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:23.802187Z"},"links":{"cited_paper":"/paper/2211.15661","citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:dc5258d465f28ec3334b75434eaa361843abc9d6eb09db96b3d4ddfb7d703af7","observation_id":"3fd457b1-f9a9-4528-a7c7-f9761bf77bf4","resolution":{"observed_at":"2026-08-02T22:19:23.802187Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.09927","last_updated":"2023-10-19T20:31:32Z","snapshot_observed_at":"2026-07-06T15:43:31.881201Z","submitted_at":"2023-06-16T15:50:03Z","title":"Trained Transformers Learn Linear Models In-Context","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.09927","snapshot_observed_at":"2026-08-02T22:19:23.982015Z","title":"Trained transformers learn linear models in-context.arXiv preprint arXiv:2306.09927, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:23.982015Z"},"links":{"cited_paper":"/paper/2306.09927","citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:be23d43deafdfc7db4bf383f04e24c562443634214840de72898276aa97c46ce","observation_id":"f2201271-7f22-4384-8742-54b34a4cdf71","resolution":{"observed_at":"2026-08-02T22:19:23.982015Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T22:19:24.217308Z","title":"Wasserstein distributionally robust optimization: Theory and applications in machine learning","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:24.217308Z"},"links":{"citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:e0051d800316af0521a34e61264d192a6b374ab4e1d631fab6b14a434099c3f2","observation_id":"41d126c2-40b3-470d-a99c-8309fd4456b4","resolution":{"observed_at":"2026-08-02T22:19:24.217308Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T22:19:24.369296Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:24.369296Z"},"links":{"citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:251cfb3e81039f1707a3dfd479967d3a5a9fd7f3a6ec5699397ed7c5d2ecd7d7","observation_id":"28325ea6-3178-449a-b6fd-38b93767890a","resolution":{"observed_at":"2026-08-02T22:19:24.369296Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T22:19:24.488775Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:24.488775Z"},"links":{"citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:d56a8c5f546d04be3ae9662a57d237a4b5135b066fcdceabe42ef1b092054db8","observation_id":"a00a9756-9f51-4f64-a711-3cca25672556","resolution":{"observed_at":"2026-08-02T22:19:24.488775Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T22:19:24.722634Z","title":"Larger models (bigger m) are smoother, making them less sensitive to shifts","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:24.722634Z"},"links":{"citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:484b670d3314e9a06ff4cad056d916a93ef88ec432dab5f0601fcae16aba6505","observation_id":"5e2cf2d5-0939-4918-86ed-93f0826426e8","resolution":{"observed_at":"2026-08-02T22:19:24.722634Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T22:19:24.900156Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:24.900156Z"},"links":{"citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:afed86f923b5a1ac1f779e72a821cb4b28ed21e30f50d0872b3661f975e266f3","observation_id":"ac8f255c-2c38-40fb-8ac7-64c95f8c3d0b","resolution":{"observed_at":"2026-08-02T22:19:24.900156Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T22:19:25.000465Z","title":"The p d/mand1/ √ Nscaling laws remain unchanged","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-02T22:19:25.000465Z"},"links":{"citing_paper":"/paper/2602.17743"},"observation_digest":"sha256:f39c5612eca950345358f27aea036d5eeece20e988f09075bed2998707a7992a","observation_id":"4a76e75d-7288-4392-bdde-6e9d07294aaf","resolution":{"observed_at":"2026-08-02T22:19:25.000465Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2602.17743","last_updated":"2026-07-20T14:54:26Z","latest_version":2,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-07T21:21:07.188072Z","submitted_at":"2026-02-19T12:37:00Z","title":"Bigger Is Safer: Provable Robustness in In-Context Learning Scales with Capacity"},"reference_resolution":{"displayed":26,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":26,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":26},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 26 of 26 outbound references and 0 inbound Pith citation observations for arXiv:2602.17743."}