{"as_of":"2026-08-04T09:59:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:8bc15426829300c73a2ad3122f1af482f6ccb7d214464180dafc5e2d4e1ee50c","coverage":[{"denominator":50,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":50,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-08T18:45:52.380042Z","state":"measured"},{"denominator":52,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":52,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-04T06:34:03.388597+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-02T21:10:10.548489Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-02T21:17:24.096157Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"cited_work":{"arxiv_id":"2605.02364","doi":null,"metadata_source":"pith","pith_arxiv_id":"2605.02364","snapshot_observed_at":"2026-07-02T21:17:24.096157Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","venue":"cs.CL","work_id":"326180e6-e50c-4a2b-9c79-304730099ad8","year":2026},"citing_paper":{"arxiv_id":"2606.28551","last_updated":"2026-06-30T20:49:34Z","snapshot_observed_at":"2026-08-02T23:02:11.703134Z","submitted_at":"2026-06-26T19:11:29Z","title":"DataComp-VLM: Improved Open Datasets for Vision-Language Models","version":1},"reference_index":174,"source":"pdf_text","source_observed_at":"2026-06-30T01:16:16.834861Z"},"links":{"cited_paper":"/paper/2605.02364","citing_paper":"/paper/2606.28551"},"observation_digest":"sha256:14f552d60b358ba11fcbac8c180d861fd88accbb9891ac2f758fc8decdd358e9","observation_id":"823ca312-4935-4b5a-a72d-dabc32919ab4","resolution":{"observed_at":"2026-07-01T15:45:47.700401Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"cited_work":{"arxiv_id":"2605.02364","doi":null,"metadata_source":"pith","pith_arxiv_id":"2605.02364","snapshot_observed_at":"2026-07-02T21:17:24.096157Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","venue":"cs.CL","work_id":"326180e6-e50c-4a2b-9c79-304730099ad8","year":2026},"citing_paper":{"arxiv_id":"2606.28551","last_updated":"2026-06-30T20:49:34Z","snapshot_observed_at":"2026-08-02T23:02:11.703134Z","submitted_at":"2026-06-26T19:11:29Z","title":"DataComp-VLM: Improved Open Datasets for Vision-Language Models","version":2},"reference_index":174,"source":"pdf_text","source_observed_at":"2026-07-02T21:10:10.548489Z"},"links":{"cited_paper":"/paper/2605.02364","citing_paper":"/paper/2606.28551"},"observation_digest":"sha256:3c2539b3d47d2676ca6aeb4362d58049266772366306d166595b4af6931abe72","observation_id":"0794ea55-ab27-4175-99b4-29d30f256602","resolution":{"observed_at":"2026-07-02T21:17:24.099070Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2605.02364/citation-record","integrity":"/paper/2605.02364/integrity","json":"/paper/2605.02364/citation-record.json","paper":"/paper/2605.02364"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T13:07:05.804778Z","title":"Scaling Learning Algorithms Towards","venue":null,"work_id":"bb2761cc-98d0-411b-92f6-803773d64460","year":null},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:3a89617c01fb596b66f99a2dff1c946abd0a9837c2c3821037c3d8697ac2e784","observation_id":"94c0a1a3-a34c-4f18-a4a3-81a8f870df79","resolution":{"observed_at":"2026-05-26T04:07:19.396697Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T13:07:05.766846Z","title":"and Osindero, Simon and Teh, Yee Whye , journal =","venue":null,"work_id":"0a5921e3-ac4e-46f1-85ae-866119a87be0","year":null},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:498b05f159295ee63f2210c6bc0b73cfae86ed524e0c3d3ffc6850b0e0c6b35a","observation_id":"e706b4a9-17e4-4b06-8541-f3de161a9435","resolution":{"observed_at":"2026-05-26T04:07:19.385586Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T21:57:45.913036Z","title":"2016 , publisher=","venue":null,"work_id":"cf0899e0-53ee-4591-aae4-f38fa5ac12ad","year":2016},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:9625b300726de8ee025495038d1c434a51612ef1026dabb1fe1962a8f6f9e5c5","observation_id":"c3faed76-48c7-4728-ad91-051d775fdd2a","resolution":{"observed_at":"2026-05-26T04:07:19.363426Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T22:57:43.199137Z","title":"2024 , eprint=","venue":null,"work_id":"56de79bb-a384-4567-8274-19e8bfa69568","year":2024},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:5eb20d80b5fc70955825386213735f0b3310b16b1e47c48670b62990a9e43d2b","observation_id":"fcb5fc7c-5c94-4441-9b3f-1e4749fa75ac","resolution":{"observed_at":"2026-05-26T04:07:19.382075Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T17:16:24.111187Z","title":"2023 , eprint=","venue":null,"work_id":"329bd184-e5b5-42bc-848f-e4d9af020cd5","year":2023},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:405c51a543cfddb7819cf7ee5cdadca027f7f93cac8a7959e202a808db503fda","observation_id":"1f84c3a7-b28d-4796-888d-e8bd434ec562","resolution":{"observed_at":"2026-05-26T04:07:19.356618Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2022.emnlp-main.616","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T00:40:18.626667Z","title":"Few-shot Learning with Multilingual Generative Language Models","venue":null,"work_id":"74c8f60a-d0c4-40dd-9dcd-a2d9195f5e09","year":2022},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:9f6c4a917124468216f0c72ca2279ab43ef6b625613dd21d6c1e2cf6d67e351e","observation_id":"c262169c-c1ec-40d1-bc5c-bbf51cbe0c29","resolution":{"observed_at":"2026-05-08T18:49:24.763775Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2025 , eprint=","venue":null,"work_id":"beaf5983-cbf0-476e-bd8c-0d1d59c9a212","year":2025},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:c5b1d0b7b93c1e2aa86c07c31c69ec5e9e9a74fb85496c8231c0d81464ddfc53","observation_id":"88200595-8201-41ce-917b-9adfe763c10a","resolution":{"observed_at":"2026-05-26T04:07:19.374978Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Transactions on Machine Learning Research , issn=","venue":null,"work_id":"8800f577-3868-44d9-a911-36f6aed194ce","year":2022},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:fd4cb7af86f73e3cce1e3a00094a73be8b65aa05b9712f08bced0680e18cbbfe","observation_id":"44b984cd-fac3-4ef1-9d4c-c793f5a72728","resolution":{"observed_at":"2026-05-26T04:07:19.320265Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2023 , eprint=","venue":null,"work_id":"f731a16e-7c33-40e0-8503-ece9b50152be","year":2023},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:54bdb7184b7fcf480c18319adebef59f2a9a0b81aa793b13b01f259647f3ef73","observation_id":"1f3303ee-e963-43a5-945a-a993ad73e00c","resolution":{"observed_at":"2026-05-26T04:07:19.390092Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Proceedings of the 41st International Conference on Machine Learning , articleno =","venue":null,"work_id":"2beb90e6-b3c9-42c1-90c0-652f0daab6fe","year":2024},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:ae5e615e09215f1e86c9c0877d8a691ce82284e0d409cc774416235170c3343a","observation_id":"47c0fd26-1584-4619-80e2-0856cc32a1db","resolution":{"observed_at":"2026-05-26T04:07:19.345413Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2024 , eprint=","venue":null,"work_id":"118a1e5a-1fea-42b2-9d62-f0e74d87b158","year":2024},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:c9e6efab18435ccb50d4a0ff127eee276ea7849fdceeb8425fa34b48ff122852","observation_id":"659719f8-3c27-4233-8a22-7bf220eeeb17","resolution":{"observed_at":"2026-05-26T04:07:19.393209Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2001.08361","last_updated":"2020-01-23T03:59:20Z","snapshot_observed_at":"2026-07-06T08:52:12.656082Z","submitted_at":"2020-01-23T03:59:20Z","title":"Scaling Laws for Neural Language Models","version":1},"cited_work":{"arxiv_id":"2001.08361","doi":"10.1145/3616855.3635845","metadata_source":"pith","pith_arxiv_id":"2001.08361","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Scaling Laws for Neural Language Models","venue":"cs.LG","work_id":"b7dd8749-9c45-4977-ab9b-64478dce1ae8","year":2020},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"cited_paper":"/paper/2001.08361","citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:da945d4616e5c8ef37c9c02dec427d569e46229c12a09452a4ad601c4d2c1e34","observation_id":"1e33909b-8832-4e07-944f-eaabc9316599","resolution":{"observed_at":"2026-05-09T06:15:37.209432Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:46:15.066696Z","title":"2025 , eprint=","venue":null,"work_id":"82c90b7d-a7ea-43ce-adc5-798edeb7ce7e","year":2025},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:8033d68e62101f202a4059a3cfaa0065e138454a1cd3636eddc71e6cb474aec5","observation_id":"5fd30999-ffed-4b0e-bebd-75152f9ad89d","resolution":{"observed_at":"2026-05-26T04:07:19.330041Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T19:11:21.546039Z","title":"2024 , eprint=","venue":null,"work_id":"c266f192-9505-4b5f-a900-4fd9ef4aac80","year":2024},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:11063227fef5351773aa10a90109b47325c38181fbbf50f53596d7033f4a3ea1","observation_id":"d0c4fe74-7c8b-4624-b75e-2a27fd6bc27e","resolution":{"observed_at":"2026-05-26T04:07:19.326675Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2023 , eprint=","venue":null,"work_id":"6dccd328-8f9a-4d05-a8ba-f4ecb56419aa","year":2023},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:2e55a53e54214de4b7d0dfb35954e615da0a347c68696a298e046cb0dda3927a","observation_id":"9713b7ac-116a-4289-b17e-fe8879fa51ce","resolution":{"observed_at":"2026-05-26T04:07:19.389519Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T02:54:28.083826Z","title":"and Sifre, Laurent , title =","venue":null,"work_id":"4d76fcc0-f1da-4d6f-a1ea-6817d49cc734","year":2022},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:90fd2ed0aa62af2a70f32bb5f546b38a21df39247f14863cf0380dd7e616beea","observation_id":"073be856-1375-4dff-b1ed-443c6991b39d","resolution":{"observed_at":"2026-05-26T04:07:19.378818Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T04:46:46.775981Z","title":"2025 , eprint=","venue":null,"work_id":"26c7b6ed-f86e-4ed8-b9ed-b1783d90255b","year":2025},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:3301aefda875882ec4628ad47e879b0b1a792278da76e7d4f95ae6fc750a5a7a","observation_id":"d3956e38-917f-4e6b-8a33-7641ef34fb70","resolution":{"observed_at":"2026-05-26T04:07:19.393008Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Proceedings of the 41st International Conference on Machine Learning , articleno =","venue":null,"work_id":"850e90da-9642-4b65-80c2-5e2dc05aaeec","year":2024},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:3bf9bcaa803b6b21e68f3637404f08920513eb11317da601f565ab6bd39182d8","observation_id":"cafc5f2a-c80b-4fed-8925-55742b1c98d5","resolution":{"observed_at":"2026-05-26T04:07:19.370493Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2024 , eprint=","venue":null,"work_id":"488208b2-9472-4be5-9948-b69d7fac14f0","year":2024},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:0fb1900cfe6c8951adbb469f915bbe020d1763af29d6c4b98e3911f2da17900a","observation_id":"4f87c13c-ca9c-44c9-b6c2-0cdc351bb17b","resolution":{"observed_at":"2026-05-26T04:07:19.371553Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2024 , eprint=","venue":null,"work_id":"ea48e479-0ebe-43ab-8190-ea1ba2c1b6c0","year":2024},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:d0ff0b01a3cd81d1de602367e70a5f7e0da4b6380415b2c4a8b9bec35c3d07f2","observation_id":"b46eca6f-8557-43f1-a480-55312c00c5e7","resolution":{"observed_at":"2026-05-26T04:07:19.327100Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2024 , eprint=","venue":null,"work_id":"6cb45b20-b7f6-47cd-8dda-067040461f6a","year":2024},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:43856a2faf072f483e1468cae73199cf9bfaaffa07180025b90665f3775fc894","observation_id":"84f28b52-ac7b-4b9a-ac98-3c3d73c68d57","resolution":{"observed_at":"2026-05-26T04:07:19.308014Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2025 , eprint=","venue":null,"work_id":"7807e554-0477-4c10-bf9b-a64710125e37","year":2025},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:bc65fec50c37bb97d1b38b535246afc2536718b21735b830c8c31ebdc54a3ebf","observation_id":"63f11eea-1e93-4808-bdca-b5502167c992","resolution":{"observed_at":"2026-05-26T04:07:19.313061Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2025 , eprint=","venue":null,"work_id":"49d85b53-013d-4519-8126-72dd68481916","year":2025},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:19ab62890f0bd8503bf49980bc078e92578e900a1ed8c2a2529062c492a9582d","observation_id":"d983311e-82c2-423f-9f76-c1a8d6d4815d","resolution":{"observed_at":"2026-05-26T04:07:19.304451Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2025 , eprint=","venue":null,"work_id":"fff16024-41c5-4b61-b920-a4b767f2b454","year":2025},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:19a5389bc5cc4a7c7f2c4ecb7bd8f43da248e4ef81887d5142e1bdfaa1970ece","observation_id":"0ffebefa-a16f-4c4f-8d1c-f45787aedbac","resolution":{"observed_at":"2026-05-26T04:07:19.357701Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"and Barak, Boaz and Le Scao, Teven and Piktus, Aleksandra and Tazi, Nouamane and Pyysalo, Sampo and Wolf, Thomas and Raffel, Colin , title =","venue":null,"work_id":"0307c144-a01d-4f82-8b60-094bfce46c9f","year":2023},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:36bba94711d2f08af73aab7bbb5a5b79cfd741f66aafb6470aba3e1d92b1836b","observation_id":"a590d4f8-86a3-42d1-bdf8-3df67d78325a","resolution":{"observed_at":"2026-05-26T04:07:19.383667Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2022 , eprint=","venue":null,"work_id":"f0d0ecf1-e20b-40f2-a8ae-e389a8070b69","year":2022},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:eda0c0fb203c3c089aff0f2033b9870e658fd794f4b92690b595857780ffec90","observation_id":"cbd4f636-075a-4e01-b5ac-c390aa82af32","resolution":{"observed_at":"2026-05-26T04:07:19.396381Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2022.acl-long.577","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T08:26:58.844496Z","title":"Deduplicating Training Data Makes Language Models Better","venue":null,"work_id":"59f1bbd1-2ddb-4173-821f-545a910ec2aa","year":2022},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:83b79a5a0fc39be375af589b53ddbcdc0c48bbe2dfcfa62ad76fc4bd9f02e874","observation_id":"cb6ef5e6-38c1-4783-accc-67670304de54","resolution":{"observed_at":"2026-05-08T18:49:25.464214Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-07-14T17:50:54.596138+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-14T17:50:54.596138+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2101.03961","last_updated":"2022-06-16T20:36:07Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2021-01-11T16:11:52Z","title":"Switch Transformers: Scaling to Trillion Parameter Models with Simple and Efficient Sparsity","version":3},"cited_work":{"arxiv_id":"2101.03961","doi":"10.1214/18-ejs1395","metadata_source":"pith","pith_arxiv_id":"2101.03961","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Switch Transformers: Scaling to Trillion Parameter Models with Simple and Efficient Sparsity","venue":"cs.LG","work_id":"f43c4955-a965-4897-a11b-c4b25d2aeaa8","year":2021},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"cited_paper":"/paper/2101.03961","citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:7289e2f383105a0e2f538e7c5467b9c37afa3bbe073b1fa07714e230e55f96d8","observation_id":"a19ae7db-ef35-4c00-8331-239a0679cfe6","resolution":{"observed_at":"2026-05-12T23:57:11.134962Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-06T11:32:22.010994Z","title":"2025 , eprint=","venue":null,"work_id":"58376947-1a64-4892-a13c-2eb8e0ab2c9f","year":2025},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:2c7de8be6691c6f26a805435b9cebcaaefc82eb5facf8fd88f93d85b7eb0ac6d","observation_id":"1f525ac2-bf3c-4768-aefe-6b2da1622f63","resolution":{"observed_at":"2026-05-26T04:07:19.360602Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1712.00409","last_updated":"2017-12-01T17:13:14Z","snapshot_observed_at":"2026-07-06T06:12:18.811024Z","submitted_at":"2017-12-01T17:13:14Z","title":"Deep Learning Scaling is Predictable, Empirically","version":1},"cited_work":{"arxiv_id":"1712.00409","doi":null,"metadata_source":"pith","pith_arxiv_id":"1712.00409","snapshot_observed_at":"2026-07-04T18:40:03.351817Z","title":"Deep Learning Scaling is Predictable, Empirically","venue":"cs.LG","work_id":"3638ccb4-3a4f-460e-8b6f-867a65922801","year":2017},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"cited_paper":"/paper/1712.00409","citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:532fb83434c78f450ef7f759bbc12ce6870f4fb6084f0f43c7fac40977ca9b89","observation_id":"c150fcac-23b5-4f74-a24b-2a7647603359","resolution":{"observed_at":"2026-05-12T04:01:58.491474Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T04:46:46.700327Z","title":"2024 , eprint=","venue":null,"work_id":"94860f33-c1e9-46de-b7ef-cdfde74468a5","year":2024},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:d52f13395c09d228c73df04294f3057303912f69ff06e4b2f95968367d5c3786","observation_id":"39142426-38ac-4710-a00c-0b180ec0084b","resolution":{"observed_at":"2026-05-26T04:07:19.347681Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2112.11446","last_updated":"2022-01-21T18:39:38Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2021-12-08T19:41:47Z","title":"Scaling Language Models: Methods, Analysis & Insights from Training Gopher","version":2},"cited_work":{"arxiv_id":"2112.11446","doi":"10.48550/arxiv.2112.11446","metadata_source":"pith","pith_arxiv_id":"2112.11446","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Scaling Language Models: Methods, Analysis & Insights from Training Gopher","venue":"cs.CL","work_id":"47ce8be9-e500-407d-af41-ac2d132215eb","year":2021},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"cited_paper":"/paper/2112.11446","citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:87eaf3fedc16d5a8aebaf01535391589724525c5f070615e67ed6005573d7b56","observation_id":"140bedea-67ec-4e37-8888-eb67179131da","resolution":{"observed_at":"2026-05-11T19:13:40.454049Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Language Models are Few-Shot Learners , url =","venue":null,"work_id":"149e3b3b-8b9e-43f7-be39-d7d01edb52db","year":null},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:2252277e1ff27c106776b91dae0d8505b1f9f6add2224b2025f192bfd48f34ea","observation_id":"8d3a83b8-4359-49be-9583-5fe1435f0e42","resolution":{"observed_at":"2026-05-26T04:07:19.338732Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Language Models are Unsupervised Multitask Learners , url =","venue":null,"work_id":"b810f583-3a11-4acf-bcf4-a9fb99c1ab96","year":null},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:41e2bf36a082955b1d119f183598b7da8a0ae44da0a75b400660d3d9f7846c04","observation_id":"c120ddd6-9746-4470-b320-b27336f6c371","resolution":{"observed_at":"2026-05-26T04:07:19.374156Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"title =","venue":null,"work_id":"8a2e4289-87d8-44f8-907f-9f1aa7b54e9d","year":null},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:13a9956c1f9cc3d798bbb3543809a605bbf0f44c3945ef297e5138e30406abf8","observation_id":"1a9b33bd-b65f-48a5-9d70-f98d3689cb7c","resolution":{"observed_at":"2026-05-26T04:07:19.349096Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"booktitle =","venue":null,"work_id":"6ebb2978-cfb6-4ed9-98bf-af2184b19c42","year":null},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:3eb50046d19e404a6b6d9200fee7d9edf20a446ce9194545bc9ee1eecd0d49f4","observation_id":"fa3cd43b-1746-4991-b602-74afb2637979","resolution":{"observed_at":"2026-05-26T04:07:19.366872Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T22:57:43.494382Z","title":"Attention is All you Need , url =","venue":null,"work_id":"fa4d7178-75f8-4139-9b39-41030a2917db","year":null},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:234c8eeaac8621bfbd7f654a5357c7e416dcb2fb5ec45a6b469429ca085b9e6f","observation_id":"e195405d-9dd7-446f-9abe-44aea9a7aa05","resolution":{"observed_at":"2026-05-26T04:07:19.284073Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"RoFormer: Enhanced transformer with Rotary Position Embedding , journal =","venue":null,"work_id":"633e2bcd-bb08-431b-8aab-08e5c981ddf3","year":2024},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:1e7f9241918c2a607cac4291aa07efb4854d0b7ee096ae8076f480180dbbd526","observation_id":"ee1e290e-31e2-4da9-abfe-55d74c0d2e78","resolution":{"observed_at":"2026-05-08T18:49:25.446433Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T14:46:18.236076Z","title":"2020 , eprint=","venue":null,"work_id":"1a5d14cc-ecd6-416e-822b-ebcb198535ce","year":2020},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:cb04075d285f9886d111709808f08e0d665aa45b65978f4cf98291f16ad0295d","observation_id":"f78e4d51-2b34-446e-9017-a2a8ccee330e","resolution":{"observed_at":"2026-05-26T04:07:19.360279Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.02954","last_updated":"2024-01-05T18:59:13Z","snapshot_observed_at":"2026-08-02T13:11:16.882565Z","submitted_at":"2024-01-05T18:59:13Z","title":"DeepSeek LLM: Scaling Open-Source Language Models with Longtermism","version":1},"cited_work":{"arxiv_id":"2401.02954","doi":"10.48550/arxiv.2401.02954","metadata_source":"pith","pith_arxiv_id":"2401.02954","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"DeepSeek LLM: Scaling Open-Source Language Models with Longtermism","venue":"cs.CL","work_id":"01b10587-025b-499d-8ba3-7a538d24c2d6","year":2024},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"cited_paper":"/paper/2401.02954","citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:c61def9a7934a23866a2cdcc315f68a3b3e097de220cac1c2361edc664716203","observation_id":"e17b215d-5d03-43fc-adb3-45ed150e63e9","resolution":{"observed_at":"2026-05-11T06:08:06.759074Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.08540","last_updated":"2024-06-14T20:21:05Z","snapshot_observed_at":"2026-07-06T17:43:56.860733Z","submitted_at":"2024-03-13T13:54:00Z","title":"Language models scale reliably with over-training and on downstream tasks","version":2},"cited_work":{"arxiv_id":"2403.08540","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.08540","snapshot_observed_at":"2026-07-04T20:30:07.554623Z","title":"Kanishk Gandhi, Denise Lee, Gabriel Grand, Muxin Liu, Winson Cheng, Archit Sharma, and Noah D Goodman","venue":null,"work_id":"b3ccca34-2e12-48b4-ad09-521ec9797b0c","year":2024},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"cited_paper":"/paper/2403.08540","citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:0e61b0224d5e5a96aa2fa673041162b2a02eee6a5eba702fa674ebca213d563a","observation_id":"60cbf363-5845-4a46-b395-5a261572d9ac","resolution":{"observed_at":"2026-05-09T06:15:37.244820Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16511","last_updated":"2025-04-26T00:13:08Z","snapshot_observed_at":"2026-07-06T21:13:27.382529Z","submitted_at":"2025-04-23T08:36:50Z","title":"QuaDMix: Quality-Diversity Balanced Data Selection for Efficient LLM Pretraining","version":2},"cited_work":{"arxiv_id":"2504.16511","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.16511","snapshot_observed_at":"2026-07-03T16:48:39.391752Z","title":"Quadmix: Quality- diversity balanced data selection for efficient llm pretraining","venue":null,"work_id":"281ee593-55a6-4abe-92a7-f73b0b1be16e","year":2025},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"cited_paper":"/paper/2504.16511","citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:232be4f4b56aca18f3031d6c51c63b7a0eea51739f2c2281293e5da8b0f21103","observation_id":"810dd919-b834-4c7b-aa5e-a33660d5f189","resolution":{"observed_at":"2026-05-09T06:15:37.219533Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T04:46:46.719407Z","title":"2018 , eprint=","venue":null,"work_id":"9880c11f-6c71-4e36-9736-affdf3d20344","year":2018},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:d76409fc1f7fcdbb7eae8df8a54b0b3798d00ba055406d293c82a210519b8dcf","observation_id":"8abad15e-4670-4cfd-b2c4-b81eae569004","resolution":{"observed_at":"2026-05-26T04:07:19.377339Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Measuring Massive Multitask Language Understanding","venue":null,"work_id":"d66ba62a-6e0b-4af9-800d-a5e13136f55b","year":null},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:0e28ba86d24221492f58fd3162a6c90ec1e451c20ee620a112f94c1ab8f69fee","observation_id":"26a2614c-d487-436c-b7ab-c9e0c382aa2e","resolution":{"observed_at":"2026-05-26T04:07:19.334332Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/p17-1147","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:46:14.489252Z","title":"TriviaQA: A large scale distantly supervised challenge dataset for reading comprehension","venue":"Proceedings of the 55th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)","work_id":"d05a9c57-9d88-473a-aa65-efb13f9dee25","year":2017},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:c2134cfdaa5031e9315addca4ed0f3edbf0b1fcc450d3597893f031ac97e82cc","observation_id":"eb1eaa39-5ee9-436b-b82a-a5e67b0dde7e","resolution":{"observed_at":"2026-05-08T18:49:24.755652Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-01T09:08:04.800216+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-01T09:08:04.800216+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/p19-1472","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T12:07:03.406813Z","title":"URL https:// doi.org/10.18653/v1/p19-1472","venue":"Proceedings of the 57th Annual Meeting of the Association for Computational Linguistics","work_id":"11bfc949-547c-40f3-a86d-953eb9b2154c","year":2019},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:fc3a0aa8f343375b9d05f92ebf50d136d6236a25a16c58aa42e5995913535171","observation_id":"caa0ee54-30c7-47e8-9583-f65d2283c90c","resolution":{"observed_at":"2026-05-08T18:49:24.759595Z","resolver_source":"doi","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-01T09:08:05.023254+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-01T09:08:05.023254+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":"e1a46fe2-9bdd-44c0-bf80-fe1ca53bc51d","year":null},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:0138b77e93f836b49cc9ea0782b342ca186855b4d7881375ab3fd232a34ef22e","observation_id":"7de3fc29-33bb-40d2-8d76-1d6fee50c1df","resolution":{"observed_at":"2026-05-26T04:07:19.363956Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.13216","last_updated":"2025-06-16T08:16:03Z","snapshot_observed_at":"2026-07-06T21:42:44.024437Z","submitted_at":"2025-06-16T08:16:03Z","title":"Capability Salience Vector: Fine-grained Alignment of Loss and Capabilities for Downstream Task Scaling Law","version":1},"cited_work":{"arxiv_id":"2506.13216","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.13216","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"arXiv preprint arXiv:2506.13216 , year=","venue":null,"work_id":"87de1bca-8023-4e10-ae67-1e5dd491b05d","year":null},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"cited_paper":"/paper/2506.13216","citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:5603da40e1519ad5563fea65b333dd3ef1fce9c888493099535201f39886ffb7","observation_id":"9e7d1f0a-44e0-42cb-82f6-574c6dd46af9","resolution":{"observed_at":"2026-05-09T06:15:37.238284Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.01492","last_updated":"2025-01-23T17:35:43Z","snapshot_observed_at":"2026-07-06T18:39:43.472708Z","submitted_at":"2024-07-01T17:31:03Z","title":"RegMix: Data Mixture as Regression for Language Model Pre-training","version":2},"cited_work":{"arxiv_id":"2407.01492","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2407.01492","snapshot_observed_at":"2026-07-04T16:19:57.817054Z","title":"Regmix: Data mixture as regression for language model pre-training","venue":null,"work_id":"e02a1110-2b86-4d6b-ae4e-7086743692b5","year":2024},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"cited_paper":"/paper/2407.01492","citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:34491163eb52d167a1b6c6de13ac827728b781b920c291be6cf695782c93e5b3","observation_id":"ef85db39-dadc-4467-a575-73fab92bf041","resolution":{"observed_at":"2026-05-09T06:15:37.272265Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":"de4fa55e-700f-4f15-a10b-b8ed3943e46a","year":null},"citing_paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-05-08T18:45:52.380042Z"},"links":{"citing_paper":"/paper/2605.02364"},"observation_digest":"sha256:b99a91685d4f24952d63ab95c0d7576a6e71dc7eb08b3a5fa01309d938fd2440","observation_id":"3508e138-148c-48f4-8c5d-50aed512980c","resolution":{"observed_at":"2026-05-26T04:07:19.380780Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2605.02364","last_updated":"2026-05-04T09:07:54Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-07-06T23:15:27.816549Z","submitted_at":"2026-05-04T09:07:54Z","title":"InfoLaw: Information Scaling Laws for Large Language Models with Quality-Weighted Mixture Data and Repetition"},"reference_resolution":{"displayed":50,"state_counts":{"malformed_identifier":0,"metadata_mismatch":8,"parse_uncertain":0,"unresolved":0,"verified_exact":6,"verified_fuzzy":36},"total_outbound_references":50},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"thesis":"As of 4 August 2026, this Paper Citation Record lists 50 of 50 outbound references and 2 inbound Pith citation observations for arXiv:2605.02364."}