{"as_of":"2026-08-08T07:20:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:ebfc235bc2b19334e44c587f4a1225ebd93cd5d8f96013e4693071834323a064","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":7,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":7,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":7,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":7,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T14:17:13.520349Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-22T14:06:38.082725Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2308.02019","last_updated":"2023-10-24T17:58:42Z","snapshot_observed_at":"2026-07-06T16:02:19.253303Z","submitted_at":"2023-08-03T20:20:01Z","title":"Baby Llama: knowledge distillation from an ensemble of teachers trained on a small dataset with no performance penalty","version":2},"cited_work":{"arxiv_id":"2308.02019","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2308.02019","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Baby llama: knowledge distillation from an ensemble of teachers trained on a small dataset with no performance penalty","venue":null,"work_id":"01a2fa23-3c60-474c-a7e3-893279f7b893","year":2023},"citing_paper":{"arxiv_id":"2401.05459","last_updated":"2024-05-08T06:16:23Z","snapshot_observed_at":"2026-08-02T13:57:57.119489Z","submitted_at":"2024-01-10T09:25:45Z","title":"Personal LLM Agents: Insights and Survey about the Capability, Efficiency and Security","version":2},"reference_index":256,"source":"pdf_text","source_observed_at":"2026-05-17T00:57:26.303195Z"},"links":{"cited_paper":"/paper/2308.02019","citing_paper":"/paper/2401.05459"},"observation_digest":"sha256:4fa585144e612de4b9c90dc0dcfdeb7771428b12cb6e53c5d0d62c81d0c22a56","observation_id":"8eb6c18a-5c4c-4d40-b3e7-89d8725f6d07","resolution":{"observed_at":"2026-05-17T00:57:26.737571Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.02019","last_updated":"2023-10-24T17:58:42Z","snapshot_observed_at":"2026-07-06T16:02:19.253303Z","submitted_at":"2023-08-03T20:20:01Z","title":"Baby Llama: knowledge distillation from an ensemble of teachers trained on a small dataset with no performance penalty","version":2},"cited_work":{"arxiv_id":"2308.02019","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2308.02019","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Baby llama: knowledge distillation from an ensemble of teachers trained on a small dataset with no performance penalty","venue":null,"work_id":"01a2fa23-3c60-474c-a7e3-893279f7b893","year":2023},"citing_paper":{"arxiv_id":"2404.14294","last_updated":"2024-07-19T04:47:36Z","snapshot_observed_at":"2026-07-06T18:03:47.096406Z","submitted_at":"2024-04-22T15:53:08Z","title":"A Survey on Efficient Inference for Large Language Models","version":3},"reference_index":134,"source":"pdf_text","source_observed_at":"2026-05-15T02:39:33.007894Z"},"links":{"cited_paper":"/paper/2308.02019","citing_paper":"/paper/2404.14294"},"observation_digest":"sha256:bdedb6f4de31973e326316f2c84e8893f39af99f5c5dbe9df6f97ca604481aca","observation_id":"8aec52a6-4426-4d97-88f8-17b905a716ff","resolution":{"observed_at":"2026-05-15T02:39:33.549278Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.02019","last_updated":"2023-10-24T17:58:42Z","snapshot_observed_at":"2026-07-06T16:02:19.253303Z","submitted_at":"2023-08-03T20:20:01Z","title":"Baby Llama: knowledge distillation from an ensemble of teachers trained on a small dataset with no performance penalty","version":2},"cited_work":{"arxiv_id":"2308.02019","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2308.02019","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Baby llama: knowledge distillation from an ensemble of teachers trained on a small dataset with no performance penalty","venue":null,"work_id":"01a2fa23-3c60-474c-a7e3-893279f7b893","year":2023},"citing_paper":{"arxiv_id":"2505.16120","last_updated":"2026-05-04T12:54:07Z","snapshot_observed_at":"2026-08-02T10:48:03.601811Z","submitted_at":"2025-05-22T01:52:15Z","title":"LLM-Powered AI Agent Systems and Their Applications in Industry","version":2},"reference_index":89,"source":"pdf_text","source_observed_at":"2026-05-22T14:05:54.535411Z"},"links":{"cited_paper":"/paper/2308.02019","citing_paper":"/paper/2505.16120"},"observation_digest":"sha256:341806f705c56df0ae4e6ad7ef6024a6eedd6eb4d7121becdbf7cc0c8b29e6e4","observation_id":"fae18880-b6dc-4a66-9424-60474cca7f56","resolution":{"observed_at":"2026-05-22T14:06:38.085837Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.02019","last_updated":"2023-10-24T17:58:42Z","snapshot_observed_at":"2026-07-06T16:02:19.253303Z","submitted_at":"2023-08-03T20:20:01Z","title":"Baby Llama: knowledge distillation from an ensemble of teachers trained on a small dataset with no performance penalty","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.02019","snapshot_observed_at":"2026-08-07T14:17:13.520349Z","title":"Llama: Open and efficient foundation language models.arXiv preprint arXiv:2308.02019,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.19529","last_updated":"2025-05-29T16:57:36Z","snapshot_observed_at":"2026-08-07T22:25:26.276499Z","submitted_at":"2025-05-26T05:29:47Z","title":"Small Language Models: Architectures, Techniques, Evaluation, Problems and Future Adaptation","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T14:17:13.520349Z"},"links":{"cited_paper":"/paper/2308.02019","citing_paper":"/paper/2505.19529"},"observation_digest":"sha256:ba4715313a63e8f76edbf3e8ac45cfeb7233d9d13740ca589b0825bfb947841a","observation_id":"7e41aa38-ff46-4ba9-8ab9-2384f277e640","resolution":{"observed_at":"2026-08-07T14:17:13.520349Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.02019","last_updated":"2023-10-24T17:58:42Z","snapshot_observed_at":"2026-07-06T16:02:19.253303Z","submitted_at":"2023-08-03T20:20:01Z","title":"Baby Llama: knowledge distillation from an ensemble of teachers trained on a small dataset with no performance penalty","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.02019","snapshot_observed_at":"2026-08-07T12:39:01.792404Z","title":"Baby llama: knowledge distillation from an en- semble of teachers trained on a small dataset with no performance penalty","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.24181","last_updated":"2025-05-30T03:43:24Z","snapshot_observed_at":"2026-08-07T12:28:41.940043Z","submitted_at":"2025-05-30T03:43:24Z","title":"SCOUT: Teaching Pre-trained Language Models to Enhance Reasoning via Flow Chain-of-Thought","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T12:39:01.792404Z"},"links":{"cited_paper":"/paper/2308.02019","citing_paper":"/paper/2505.24181"},"observation_digest":"sha256:2e9c973586f704f8d394a71840bfb6acd5e5b1efbaac090d1d718fd8d487deac","observation_id":"7ef4a804-9d0b-4784-8126-ff010bf97736","resolution":{"observed_at":"2026-08-07T12:39:01.792404Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.02019","last_updated":"2023-10-24T17:58:42Z","snapshot_observed_at":"2026-07-06T16:02:19.253303Z","submitted_at":"2023-08-03T20:20:01Z","title":"Baby Llama: knowledge distillation from an ensemble of teachers trained on a small dataset with no performance penalty","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.02019","snapshot_observed_at":"2026-08-06T23:57:26.323704Z","title":"arXiv preprint arXiv:2308.02019 (2023)","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.15681","last_updated":"2026-06-25T14:33:27Z","snapshot_observed_at":"2026-08-06T23:49:22.433306Z","submitted_at":"2025-06-18T17:59:49Z","title":"GenRecal: Generation after Recalibration from Large to Small Vision-Language Models","version":4},"reference_index":97,"source":"pdf_text","source_observed_at":"2026-08-06T23:57:26.323704Z"},"links":{"cited_paper":"/paper/2308.02019","citing_paper":"/paper/2506.15681"},"observation_digest":"sha256:a03b92dd36f0ffb97807236a3a17ae3f6b80550b78974d92cd3a7f20e480200d","observation_id":"e8975f7d-f1f9-43a4-bbe4-5d7b3c9f2384","resolution":{"observed_at":"2026-08-06T23:57:26.323704Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.02019","last_updated":"2023-10-24T17:58:42Z","snapshot_observed_at":"2026-07-06T16:02:19.253303Z","submitted_at":"2023-08-03T20:20:01Z","title":"Baby Llama: knowledge distillation from an ensemble of teachers trained on a small dataset with no performance penalty","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.02019","snapshot_observed_at":"2026-08-06T19:29:09.101214Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.05617","last_updated":"2025-07-08T02:54:15Z","snapshot_observed_at":"2026-08-07T03:43:31.678279Z","submitted_at":"2025-07-08T02:54:15Z","title":"Flipping Knowledge Distillation: Leveraging Small Models' Expertise to Enhance LLMs in Text Matching","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-06T19:29:09.101214Z"},"links":{"cited_paper":"/paper/2308.02019","citing_paper":"/paper/2507.05617"},"observation_digest":"sha256:8afbf874df8457435479a22e9e2659ee842da70336ce20164805eebf7fd2b039","observation_id":"3648a555-298b-4291-a034-82a2532b8fa9","resolution":{"observed_at":"2026-08-06T19:29:09.101214Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2308.02019/citation-record","integrity":"/paper/2308.02019/integrity","json":"/paper/2308.02019/citation-record.json","paper":"/paper/2308.02019"},"outbound":[],"paper":{"arxiv_id":"2308.02019","last_updated":"2023-10-24T17:58:42Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-07-06T16:02:19.253303Z","submitted_at":"2023-08-03T20:20:01Z","title":"Baby Llama: knowledge distillation from an ensemble of teachers trained on a small dataset with no performance penalty"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 7 inbound Pith citation observations for arXiv:2308.02019."}