{"as_of":"2026-08-05T11:10:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:e106fe39a8014dd03ec9bfaf1d4a9188303e8aad79ce617dd163df7f9b8b60d1","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":7,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":7,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-05T06:32:48.257954+00:00","state":"measured"},{"denominator":7,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":7,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-29T14:14:25.876963Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":19,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2109.07740","last_updated":"2021-09-16T06:15:20Z","snapshot_observed_at":"2026-07-06T11:48:14.757537Z","submitted_at":"2021-09-16T06:15:20Z","title":"Scaling Laws for Neural Machine Translation","version":1},"cited_work":{"arxiv_id":"2109.07740","doi":"10.48550/arxiv.2109.07740","metadata_source":"arxiv_reference","pith_arxiv_id":"2109.07740","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Deep learning versus kernel learning: an empirical study of loss landscape geometry and the time evolution of the neural tangent kernel","venue":"arXiv (Cornell University)","work_id":"aac9427a-e822-4e50-9d69-8595bce3696f","year":2021},"citing_paper":{"arxiv_id":"2304.01373","last_updated":"2023-05-31T17:54:07Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-04-03T20:58:15Z","title":"Pythia: A Suite for Analyzing Large Language Models Across Training and Scaling","version":2},"reference_index":124,"source":"arxiv_source","source_observed_at":"2026-05-15T17:45:17.540282Z"},"links":{"cited_paper":"/paper/2109.07740","citing_paper":"/paper/2304.01373"},"observation_digest":"sha256:6516180dc00f195a486a9a88525c78122cd67fb74886b75d10b16c7dddb956c3","observation_id":"7c279b5b-bb6c-4555-80a1-db7b76935550","resolution":{"observed_at":"2026-05-15T17:45:17.791463Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2109.07740","last_updated":"2021-09-16T06:15:20Z","snapshot_observed_at":"2026-07-06T11:48:14.757537Z","submitted_at":"2021-09-16T06:15:20Z","title":"Scaling Laws for Neural Machine Translation","version":1},"cited_work":{"arxiv_id":"2109.07740","doi":"10.48550/arxiv.2109.07740","metadata_source":"arxiv_reference","pith_arxiv_id":"2109.07740","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Deep learning versus kernel learning: an empirical study of loss landscape geometry and the time evolution of the neural tangent kernel","venue":"arXiv (Cornell University)","work_id":"aac9427a-e822-4e50-9d69-8595bce3696f","year":2021},"citing_paper":{"arxiv_id":"2305.16264","last_updated":"2025-06-28T00:00:06Z","snapshot_observed_at":"2026-07-06T15:33:27.761070Z","submitted_at":"2023-05-25T17:18:55Z","title":"Scaling Data-Constrained Language Models","version":5},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-18T01:35:21.150772Z"},"links":{"cited_paper":"/paper/2109.07740","citing_paper":"/paper/2305.16264"},"observation_digest":"sha256:e89889fd9f78eb962882c2a20d96f15cd623348165450e1c9cd24de7f9ce490c","observation_id":"e98c2853-cd4b-45dc-ab15-872aa187c3b3","resolution":{"observed_at":"2026-05-18T01:35:21.619989Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2109.07740","last_updated":"2021-09-16T06:15:20Z","snapshot_observed_at":"2026-07-06T11:48:14.757537Z","submitted_at":"2021-09-16T06:15:20Z","title":"Scaling Laws for Neural Machine Translation","version":1},"cited_work":{"arxiv_id":"2109.07740","doi":"10.48550/arxiv.2109.07740","metadata_source":"arxiv_reference","pith_arxiv_id":"2109.07740","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Deep learning versus kernel learning: an empirical study of loss landscape geometry and the time evolution of the neural tangent kernel","venue":"arXiv (Cornell University)","work_id":"aac9427a-e822-4e50-9d69-8595bce3696f","year":2021},"citing_paper":{"arxiv_id":"2308.08998","last_updated":"2023-08-21T10:23:42Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-08-17T14:12:48Z","title":"Reinforced Self-Training (ReST) for Language Modeling","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-13T07:59:55.849296Z"},"links":{"cited_paper":"/paper/2109.07740","citing_paper":"/paper/2308.08998"},"observation_digest":"sha256:9b499413befe64657550c8cbb57c8ecf12f4d38195a3ecb20a10dd0457752420","observation_id":"c03dddc1-7ccd-4665-8a17-5cecf311a4ec","resolution":{"observed_at":"2026-05-13T07:59:56.008047Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2109.07740","last_updated":"2021-09-16T06:15:20Z","snapshot_observed_at":"2026-07-06T11:48:14.757537Z","submitted_at":"2021-09-16T06:15:20Z","title":"Scaling Laws for Neural Machine Translation","version":1},"cited_work":{"arxiv_id":"2109.07740","doi":"10.48550/arxiv.2109.07740","metadata_source":"arxiv_reference","pith_arxiv_id":"2109.07740","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Deep learning versus kernel learning: an empirical study of loss landscape geometry and the time evolution of the neural tangent kernel","venue":"arXiv (Cornell University)","work_id":"aac9427a-e822-4e50-9d69-8595bce3696f","year":2021},"citing_paper":{"arxiv_id":"2405.00592","last_updated":"2025-06-30T15:11:57Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-05-01T15:59:00Z","title":"Scaling and renormalization in high-dimensional regression","version":4},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-24T01:54:48.781227Z"},"links":{"cited_paper":"/paper/2109.07740","citing_paper":"/paper/2405.00592"},"observation_digest":"sha256:c138d6cf32ae35024e169a41234c0aeeacdeeac22e47233995610e5eca87f9b0","observation_id":"2fe0ca98-a977-49cc-8ecf-9e831ad29a84","resolution":{"observed_at":"2026-05-24T01:55:55.094304Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2109.07740","last_updated":"2021-09-16T06:15:20Z","snapshot_observed_at":"2026-07-06T11:48:14.757537Z","submitted_at":"2021-09-16T06:15:20Z","title":"Scaling Laws for Neural Machine Translation","version":1},"cited_work":{"arxiv_id":"2109.07740","doi":"10.48550/arxiv.2109.07740","metadata_source":"arxiv_reference","pith_arxiv_id":"2109.07740","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Deep learning versus kernel learning: an empirical study of loss landscape geometry and the time evolution of the neural tangent kernel","venue":"arXiv (Cornell University)","work_id":"aac9427a-e822-4e50-9d69-8595bce3696f","year":2021},"citing_paper":{"arxiv_id":"2405.14782","last_updated":"2026-05-31T00:04:33Z","snapshot_observed_at":"2026-08-02T05:44:43.939336Z","submitted_at":"2024-05-23T16:50:49Z","title":"Lessons from the Trenches on Reproducible Evaluation of Language Models","version":2},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-05-16T18:44:49.519995Z"},"links":{"cited_paper":"/paper/2109.07740","citing_paper":"/paper/2405.14782"},"observation_digest":"sha256:208cb3df781079344621fe7dabec543eab412e8480bf60beb02269db4cac95af","observation_id":"7450ff08-f62a-402d-92dd-ab5e24db16d1","resolution":{"observed_at":"2026-05-16T18:44:49.725793Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2109.07740","last_updated":"2021-09-16T06:15:20Z","snapshot_observed_at":"2026-07-06T11:48:14.757537Z","submitted_at":"2021-09-16T06:15:20Z","title":"Scaling Laws for Neural Machine Translation","version":1},"cited_work":{"arxiv_id":"2109.07740","doi":"10.48550/arxiv.2109.07740","metadata_source":"arxiv_reference","pith_arxiv_id":"2109.07740","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Deep learning versus kernel learning: an empirical study of loss landscape geometry and the time evolution of the neural tangent kernel","venue":"arXiv (Cornell University)","work_id":"aac9427a-e822-4e50-9d69-8595bce3696f","year":2021},"citing_paper":{"arxiv_id":"2502.05074","last_updated":"2025-11-10T21:45:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-02-07T16:45:40Z","title":"Two-Point Deterministic Equivalence for Stochastic Gradient Dynamics in Linear Models","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-23T04:16:04.110552Z"},"links":{"cited_paper":"/paper/2109.07740","citing_paper":"/paper/2502.05074"},"observation_digest":"sha256:b8260b696b11d3d9805179bbe97e3795d67c327b29cadcea18611d77e3ef6a3e","observation_id":"6b523248-fc7c-495e-9405-61a0e3e4437b","resolution":{"observed_at":"2026-05-23T04:17:30.967158Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2109.07740","last_updated":"2021-09-16T06:15:20Z","snapshot_observed_at":"2026-07-06T11:48:14.757537Z","submitted_at":"2021-09-16T06:15:20Z","title":"Scaling Laws for Neural Machine Translation","version":1},"cited_work":{"arxiv_id":"2109.07740","doi":"10.48550/arxiv.2109.07740","metadata_source":"arxiv_reference","pith_arxiv_id":"2109.07740","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Deep learning versus kernel learning: an empirical study of loss landscape geometry and the time evolution of the neural tangent kernel","venue":"arXiv (Cornell University)","work_id":"aac9427a-e822-4e50-9d69-8595bce3696f","year":2021},"citing_paper":{"arxiv_id":"2605.27989","last_updated":"2026-05-27T05:23:45Z","snapshot_observed_at":"2026-08-03T00:57:30.869718Z","submitted_at":"2026-05-27T05:23:45Z","title":"Law of Neural Interaction: Depth-Width Shape, Interaction Efficiency, and Generalization","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-06-29T14:14:25.876963Z"},"links":{"cited_paper":"/paper/2109.07740","citing_paper":"/paper/2605.27989"},"observation_digest":"sha256:95bbfb10a223cb434be8d6f2bc81d5a60f47adaf953246e95697563a099758a8","observation_id":"cbe6605a-7ba2-4d33-a680-6302b2832113","resolution":{"observed_at":"2026-06-29T14:23:30.992298Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2109.07740/citation-record","integrity":"/paper/2109.07740/integrity","json":"/paper/2109.07740/citation-record.json","paper":"/paper/2109.07740"},"outbound":[],"paper":{"arxiv_id":"2109.07740","last_updated":"2021-09-16T06:15:20Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-07-06T11:48:14.757537Z","submitted_at":"2021-09-16T06:15:20Z","title":"Scaling Laws for Neural Machine Translation"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"thesis":"As of 5 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 7 inbound Pith citation observations for arXiv:2109.07740."}