{"as_of":"2026-08-07T09:42:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:98cccd232bee353fedc6828177f5518e9f49f7b29d9e31be51c3a12e88f636fc","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":25,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":25,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":25,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":25,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T20:15:24.011287Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":7,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2303.18223","last_updated":"2026-03-18T05:34:39Z","snapshot_observed_at":"2026-08-06T23:27:24.356320Z","submitted_at":"2023-03-31T17:28:46Z","title":"A Survey of Large Language Models","version":19},"reference_index":236,"source":"pdf_text","source_observed_at":"2026-05-10T22:46:39.268353Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2303.18223"},"observation_digest":"sha256:537a8f057ed86da1dee8f4412a045a826c40b4341186f3e9e0a5141b69520e36","observation_id":"c894b174-2f1a-4007-92fb-7003f519a3fe","resolution":{"observed_at":"2026-05-10T22:46:40.287347Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2309.12284","last_updated":"2024-05-03T17:36:07Z","snapshot_observed_at":"2026-08-02T15:00:50.388422Z","submitted_at":"2023-09-21T17:45:42Z","title":"MetaMath: Bootstrap Your Own Mathematical Questions for Large Language Models","version":4},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-05-13T10:07:53.748795Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2309.12284"},"observation_digest":"sha256:984f56d84eb9c2d9cd8094f1772d5cedecbba146a0510b58d0ca460235451b8e","observation_id":"2cabdda0-928c-42f8-8775-028b8a8be8ec","resolution":{"observed_at":"2026-05-13T10:07:53.863501Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2411.05527","last_updated":"2026-05-04T11:38:18Z","snapshot_observed_at":"2026-07-06T19:47:25.840369Z","submitted_at":"2024-11-08T12:35:58Z","title":"How Good is Your Wikipedia? Auditing Data Quality for Low-resource and Multilingual NLP","version":3},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-05-23T17:36:18.451771Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2411.05527"},"observation_digest":"sha256:26ec752f4e132b438f19c469a4e367d0a709e20f654188633cbc80797dd8431d","observation_id":"80cb290a-d72a-4454-980d-5e17b1e77046","resolution":{"observed_at":"2026-05-23T17:38:15.898675Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2504.21850","last_updated":"2026-05-14T18:32:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-04-30T17:57:22Z","title":"Visual Compositional Tuning","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-22T17:39:09.890605Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2504.21850"},"observation_digest":"sha256:4a019a7cb28817782bef603d9f0dd7d61c5d5f16e14b3b91ddf2e5fa3809589f","observation_id":"9fa03547-68b8-47b7-b02f-0e39ab57bd3c","resolution":{"observed_at":"2026-05-22T17:41:53.142850Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-06T20:15:24.011287Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.03648","last_updated":"2025-07-04T15:25:04Z","snapshot_observed_at":"2026-08-06T20:02:40.466269Z","submitted_at":"2025-07-04T15:25:04Z","title":"Disentangling the Roles of Representation and Selection in Data Pruning","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-06T20:15:24.011287Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2507.03648"},"observation_digest":"sha256:77b58cf48e01701121c5ed936eb73d98853ba97b3cabb2521deeed3d1f9d42f3","observation_id":"ca62f42a-f897-4865-96b0-edae64c870fa","resolution":{"observed_at":"2026-08-06T20:15:24.011287Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-06T16:53:11.656872Z","title":"When less is more: Investigating data pruning for pretraining llms at scale","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.12466","last_updated":"2025-07-16T17:59:45Z","snapshot_observed_at":"2026-08-06T16:42:43.163399Z","submitted_at":"2025-07-16T17:59:45Z","title":"Language Models Improve When Pretraining Data Matches Target Tasks","version":1},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-08-06T16:53:11.656872Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2507.12466"},"observation_digest":"sha256:3b53c789729f00f204efcf61b77d6aa9cedd4e46b8d09c51b0f34c124fe7683a","observation_id":"6d07eb10-06d0-43e4-8429-1cb456aa0180","resolution":{"observed_at":"2026-08-06T16:53:11.656872Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T20:00:05.194995Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.11551","last_updated":"2025-08-18T06:38:38Z","snapshot_observed_at":"2026-08-05T19:58:29.896587Z","submitted_at":"2025-08-15T15:53:09Z","title":"ADMIRE-BayesOpt: Accelerated Data MIxture RE-weighting for Language Models with Bayesian Optimization","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-05T20:00:05.194995Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2508.11551"},"observation_digest":"sha256:3aac04e3015e5be49e862cbc386471f7699e742bb293871883a3ba6bbf57a64e","observation_id":"f1ed7484-fbb0-434d-9557-1e0cf813a131","resolution":{"observed_at":"2026-08-05T20:00:05.194995Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2602.18584","last_updated":"2026-05-16T01:08:46Z","snapshot_observed_at":"2026-08-02T09:38:59.455163Z","submitted_at":"2026-02-20T19:44:24Z","title":"GIST: Targeted Data Selection for Instruction Tuning via Coupled Optimization Geometry","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-21T12:26:14.261351Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2602.18584"},"observation_digest":"sha256:c9720cfeac4168d03f4e5a689807f16cc574d6353b9a9a6bd32677986a80aa57","observation_id":"2e348e6e-7e5a-4f43-9d9d-a883ad6a9fef","resolution":{"observed_at":"2026-05-21T12:30:07.825970Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-03T02:34:19.737306Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2603.17205","last_updated":"2026-07-30T22:57:15Z","snapshot_observed_at":"2026-08-06T17:53:38.735441Z","submitted_at":"2026-03-17T23:11:45Z","title":"OPERA: Online Data Pruning for Efficient Retrieval Model Adaptation","version":3},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-03T02:34:19.737306Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2603.17205"},"observation_digest":"sha256:0fd109d1a138bf7b68b6ad7a6cc94957322eb52561099cfa629d8190eb309bfb","observation_id":"02f8e2ee-dc95-42cb-a6ab-97b1549804ad","resolution":{"observed_at":"2026-08-03T02:34:19.737306Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2604.02345","last_updated":"2026-02-11T17:02:38Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-02-11T17:02:38Z","title":"UI-Oceanus: Scaling GUI Agents with Synthetic Environmental Dynamics","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-05-16T02:59:27.807789Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2604.02345"},"observation_digest":"sha256:5dfcfcf40f37ef10f49667d75016ac3d01f7745a91d4deeb3a2c20027b47bcf2","observation_id":"f79798d1-047b-49c3-a62f-082c3e17873d","resolution":{"observed_at":"2026-05-16T03:00:31.597557Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2604.07940","last_updated":"2026-04-09T08:00:22Z","snapshot_observed_at":"2026-07-06T22:57:09.435331Z","submitted_at":"2026-04-09T08:00:22Z","title":"A Systematic Framework for Tabular Data Disentanglement","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-10T18:27:39.046005Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2604.07940"},"observation_digest":"sha256:2ab42e66c6a75c12cdf26a74716ca17ca2da599776c578f60c737bc0ddb421b5","observation_id":"bec684eb-a500-4351-a15f-776c0af3754b","resolution":{"observed_at":"2026-05-11T00:30:55.343421Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2604.11810","last_updated":"2026-04-09T14:08:01Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-04-09T14:08:01Z","title":"GRACE: A Dynamic Coreset Selection Framework for Large Language Model Optimization","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-05-10T18:06:46.131725Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2604.11810"},"observation_digest":"sha256:26633efddfb98be62c5969a78fd40b33c6f834c20d66421336c6cd397c2d9b2a","observation_id":"a3ddf1a1-2cb4-416e-b88d-752280d835eb","resolution":{"observed_at":"2026-05-10T20:30:49.021001Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2604.17396","last_updated":"2026-04-19T11:59:58Z","snapshot_observed_at":"2026-08-01T18:23:32.699442Z","submitted_at":"2026-04-19T11:59:58Z","title":"Representation-Guided Parameter-Efficient LLM Unlearning","version":1},"reference_index":135,"source":"arxiv_source","source_observed_at":"2026-05-10T06:01:46.885030Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2604.17396"},"observation_digest":"sha256:d59dab2c2dd24545d878a595b5e62257628e0e83db97b53ae42e58fa28cdf410","observation_id":"3a677dd1-e712-45c9-805e-2f19681436fb","resolution":{"observed_at":"2026-05-10T06:06:19.197284Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2605.00369","last_updated":"2026-06-05T01:05:58Z","snapshot_observed_at":"2026-07-06T23:13:47.082550Z","submitted_at":"2026-05-01T03:12:16Z","title":"InvEvolve: Evolving White-Box Inventory Policies via Large Language Models with Performance Guarantees","version":2},"reference_index":159,"source":"arxiv_source","source_observed_at":"2026-05-09T19:50:39.734124Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2605.00369"},"observation_digest":"sha256:4a9c7e060b69db10ec38c413586de6cac8bc3f3eb2132172363c71fcc9344455","observation_id":"42d6e4f0-05f0-4d7e-b871-71d586e13301","resolution":{"observed_at":"2026-05-11T15:31:07.878353Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2605.00369","last_updated":"2026-06-05T01:05:58Z","snapshot_observed_at":"2026-07-06T23:13:47.082550Z","submitted_at":"2026-05-01T03:12:16Z","title":"InvEvolve: Evolving White-Box Inventory Policies via Large Language Models with Performance Guarantees","version":3},"reference_index":159,"source":"arxiv_source","source_observed_at":"2026-05-12T02:38:45.322351Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2605.00369"},"observation_digest":"sha256:f9f36883fe61f16d813c5e37d905f3660021d0911f4aefe36be4ef148f926a93","observation_id":"06f0d75f-3540-4396-bc2f-322592cb7c9d","resolution":{"observed_at":"2026-05-12T02:41:17.710831Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2605.05227","last_updated":"2026-04-19T14:23:23Z","snapshot_observed_at":"2026-08-02T14:54:36.818075Z","submitted_at":"2026-04-19T14:23:23Z","title":"Rethinking Data Curation in LLM Training: Online Reweighting Offers Better Generalization than Offline Methods","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-10T07:09:21.652035Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2605.05227"},"observation_digest":"sha256:7e8617c4b438983fd6a09dbe2a80713581200a2580e6f4d46179fafc55726d05","observation_id":"9b194a4b-747d-408c-854b-b7680421d8ca","resolution":{"observed_at":"2026-05-10T07:11:53.392676Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2605.12906","last_updated":"2026-05-25T14:11:40Z","snapshot_observed_at":"2026-07-06T23:24:32.412966Z","submitted_at":"2026-05-13T02:33:04Z","title":"Data Difficulty and the Generalization--Extrapolation Tradeoff in LLM Fine-Tuning","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-05-14T20:02:58.318276Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2605.12906"},"observation_digest":"sha256:755a2f6f654189e9d3a01356bc9068bc33e085b1f91e4877c244800526db5c2e","observation_id":"b00f21c8-e6b6-4814-ac7b-59860a533759","resolution":{"observed_at":"2026-05-14T20:07:54.465343Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2605.19762","last_updated":"2026-05-19T12:37:01Z","snapshot_observed_at":"2026-08-02T21:57:08.320932Z","submitted_at":"2026-05-19T12:37:01Z","title":"What Really Improves Mathematical Reasoning: Structured Reasoning Signals Beyond Pure Code","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-05-20T05:06:46.360174Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2605.19762"},"observation_digest":"sha256:9f30e7de24a4bbd1671e7ed7077bb7e766223ce1d74fd21e2f2fad167bec3ebe","observation_id":"d5555e70-5a61-4bd6-98b2-0169e925a5cf","resolution":{"observed_at":"2026-05-20T05:08:05.057764Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2605.22389","last_updated":"2026-05-21T12:21:41Z","snapshot_observed_at":"2026-07-06T23:32:44.699825Z","submitted_at":"2026-05-21T12:21:41Z","title":"Unified Data Selection for LLM Reasoning","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-05-22T05:33:20.930156Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2605.22389"},"observation_digest":"sha256:04ef399cb356590c4655d9c7274e3efc7212e267d0e31b20c5def3f8ba4a2030","observation_id":"db30ffd2-cf6e-4225-839f-d856a63944c6","resolution":{"observed_at":"2026-05-22T05:34:40.137379Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2605.28631","last_updated":"2026-05-27T15:38:09Z","snapshot_observed_at":"2026-07-06T23:38:09.179597Z","submitted_at":"2026-05-27T15:38:09Z","title":"Single-Rollout Hidden-State Dynamics for Training-Free RLVR Data Selection","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-06-29T14:08:40.968105Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2605.28631"},"observation_digest":"sha256:0d5bde2f6b9a33b3e8224ac27dd8f3595cc0bca8280761166ee4379c2f1b5893","observation_id":"ab1d612f-3fa9-44d0-923c-f018258698cf","resolution":{"observed_at":"2026-06-29T14:13:30.110956Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2606.08574","last_updated":"2026-06-07T11:11:51Z","snapshot_observed_at":"2026-08-03T13:39:16.935988Z","submitted_at":"2026-06-07T11:11:51Z","title":"OrderDP: A Theoretically Guaranteed Lossless Dynamic Data Pruning Framework","version":1},"reference_index":111,"source":"arxiv_source","source_observed_at":"2026-06-27T18:26:32.883834Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2606.08574"},"observation_digest":"sha256:74bd648eb302ea8f54bd4001db01a9a20de5c5132b6ed05ef83f1b86d68c8278","observation_id":"9977346f-1710-4a74-a843-032e2e0f72e6","resolution":{"observed_at":"2026-07-02T23:07:27.218492Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2606.23611","last_updated":"2026-06-22T17:11:15Z","snapshot_observed_at":"2026-07-06T23:58:20.975630Z","submitted_at":"2026-06-22T17:11:15Z","title":"Data Selection Through Iterative Self-Filtering for Vision-Language Settings","version":1},"reference_index":195,"source":"arxiv_source","source_observed_at":"2026-06-26T09:22:47.537137Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2606.23611"},"observation_digest":"sha256:7abf9682a37ebe350dbe859a53e9b853219a2a2014ceaa78e75c9eebf91c6d8a","observation_id":"8b75c3a5-0934-4ce1-a63d-ccd01945265c","resolution":{"observed_at":"2026-07-04T09:49:44.946514Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2606.24998","last_updated":"2026-06-23T16:02:40Z","snapshot_observed_at":"2026-07-31T12:51:11.444875Z","submitted_at":"2026-06-23T16:02:40Z","title":"Internal Data Repetition Destroys Language Models","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-06-26T00:12:56.745617Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2606.24998"},"observation_digest":"sha256:634dd4a02ada4323adfa4643b86c2d35f896c76dcdc1e0553e1ed7f2e7f4b7f0","observation_id":"ddfb4caa-9c82-4503-86fc-8301fe8837a6","resolution":{"observed_at":"2026-07-04T16:49:57.884375Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":"2309.04564","doi":"10.48550/arxiv.2309.04564","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"When less is more: Investigating data pruning for pretraining llms at scale.arXiv preprint arXiv:2309.04564","venue":"arXiv (Cornell University)","work_id":"cc2c01ad-f12f-43a3-ac00-8df29526cffd","year":2023},"citing_paper":{"arxiv_id":"2606.26091","last_updated":"2026-06-24T17:59:02Z","snapshot_observed_at":"2026-08-02T03:35:37.854415Z","submitted_at":"2026-06-24T17:59:02Z","title":"On-Policy Self-Distillation with Sampled Demonstrations Reduces Output Diversity","version":1},"reference_index":216,"source":"arxiv_source","source_observed_at":"2026-06-25T19:23:56.452083Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2606.26091"},"observation_digest":"sha256:3efbd46a8bdfcb1f2c139a4a5183ab496a0c0d9ab91b1191858b27cf0a9a78c5","observation_id":"870d7ae7-8728-4ea3-8419-a661c6fd1a08","resolution":{"observed_at":"2026-07-04T20:50:12.707543Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.04564","snapshot_observed_at":"2026-07-31T11:19:34.319148Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.24585","last_updated":"2026-07-27T15:51:41Z","snapshot_observed_at":"2026-08-02T21:43:59.103624Z","submitted_at":"2026-07-27T15:51:41Z","title":"From Data to Device: ELMOD An Efficient German-First 2.7B Language Model for Mobile Inference","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-07-31T11:19:34.319148Z"},"links":{"cited_paper":"/paper/2309.04564","citing_paper":"/paper/2607.24585"},"observation_digest":"sha256:c48fe262e4b4b506e1f99cac2ee48168fa9f379909512d0c29d888575d9a70ac","observation_id":"94a9d6c8-a587-4be7-8fed-9f93f61025bc","resolution":{"observed_at":"2026-07-31T11:19:34.319148Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2309.04564/citation-record","integrity":"/paper/2309.04564/integrity","json":"/paper/2309.04564/citation-record.json","paper":"/paper/2309.04564"},"outbound":[],"paper":{"arxiv_id":"2309.04564","last_updated":"2023-09-08T19:34:05Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-07T03:50:04.516601Z","submitted_at":"2023-09-08T19:34:05Z","title":"When Less is More: Investigating Data Pruning for Pretraining LLMs at Scale"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 25 inbound Pith citation observations for arXiv:2309.04564."}