{"as_of":"2026-08-10T05:07:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:18a994059d2c5ca6af5fda42e6f36649ad4e9f791bd990338160f86231cb32ac","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":13,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":13,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":13,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":13,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-09T14:27:57.423803Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":1,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2403.15796","last_updated":"2025-01-15T02:48:59Z","snapshot_observed_at":"2026-08-04T02:07:01.235387Z","submitted_at":"2024-03-23T11:03:31Z","title":"Understanding Emergent Abilities of Language Models from the Loss Perspective","version":3},"cited_work":{"arxiv_id":"2403.15796","doi":"10.48550/arxiv.2403.15796","metadata_source":"arxiv_reference","pith_arxiv_id":"2403.15796","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2403.15796 , year=","venue":"arXiv (Cornell University)","work_id":"871bc70b-ac57-442c-892b-f8d06c9abab2","year":2025},"citing_paper":{"arxiv_id":"2404.06395","last_updated":"2024-06-03T08:54:38Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-04-09T15:36:50Z","title":"MiniCPM: Unveiling the Potential of Small Language Models with Scalable Training Strategies","version":3},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-13T18:00:53.389420Z"},"links":{"cited_paper":"/paper/2403.15796","citing_paper":"/paper/2404.06395"},"observation_digest":"sha256:66a9df5efe8d4f3fb8c556bef62f7fdb17f53ba067b407cc3ffd42bc3294d371","observation_id":"c31e2163-919e-4684-958e-f1391d816486","resolution":{"observed_at":"2026-05-13T18:00:53.457262Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.15796","last_updated":"2025-01-15T02:48:59Z","snapshot_observed_at":"2026-08-04T02:07:01.235387Z","submitted_at":"2024-03-23T11:03:31Z","title":"Understanding Emergent Abilities of Language Models from the Loss Perspective","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.15796","snapshot_observed_at":"2026-08-09T14:27:57.423803Z","title":"Understanding emergent abilities of language models from the loss perspective","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.01804","last_updated":"2025-02-03T20:33:20Z","snapshot_observed_at":"2026-08-10T01:51:21.193055Z","submitted_at":"2025-02-03T20:33:20Z","title":"Soup-of-Experts: Pretraining Specialist Models via Parameters Averaging","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-09T14:27:57.423803Z"},"links":{"cited_paper":"/paper/2403.15796","citing_paper":"/paper/2502.01804"},"observation_digest":"sha256:e7fbe4d2fbb55f5d3c855fcb9cd5de652c60b1201684cb5052b8c31a66dd67b1","observation_id":"e84ceb55-0f61-4ce1-b195-7ba48fdcab56","resolution":{"observed_at":"2026-08-09T14:27:57.423803Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.15796","last_updated":"2025-01-15T02:48:59Z","snapshot_observed_at":"2026-08-04T02:07:01.235387Z","submitted_at":"2024-03-23T11:03:31Z","title":"Understanding Emergent Abilities of Language Models from the Loss Perspective","version":3},"cited_work":{"arxiv_id":"2403.15796","doi":"10.48550/arxiv.2403.15796","metadata_source":"arxiv_reference","pith_arxiv_id":"2403.15796","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2403.15796 , year=","venue":"arXiv (Cornell University)","work_id":"871bc70b-ac57-442c-892b-f8d06c9abab2","year":2025},"citing_paper":{"arxiv_id":"2502.02737","last_updated":"2025-02-04T21:43:16Z","snapshot_observed_at":"2026-08-10T04:07:56.506438Z","submitted_at":"2025-02-04T21:43:16Z","title":"SmolLM2: When Smol Goes Big -- Data-Centric Training of a Small Language Model","version":1},"reference_index":168,"source":"arxiv_source","source_observed_at":"2026-05-13T17:30:02.803757Z"},"links":{"cited_paper":"/paper/2403.15796","citing_paper":"/paper/2502.02737"},"observation_digest":"sha256:9adcd8d7c1b166fd76d1987e11e8c624cb50c0f63030a793aacac6b370e3bcc6","observation_id":"3a6eefdf-0aad-4e04-a859-e9f4de283a6b","resolution":{"observed_at":"2026-05-13T17:30:02.937523Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.15796","last_updated":"2025-01-15T02:48:59Z","snapshot_observed_at":"2026-08-04T02:07:01.235387Z","submitted_at":"2024-03-23T11:03:31Z","title":"Understanding Emergent Abilities of Language Models from the Loss Perspective","version":3},"cited_work":{"arxiv_id":"2403.15796","doi":"10.48550/arxiv.2403.15796","metadata_source":"arxiv_reference","pith_arxiv_id":"2403.15796","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2403.15796 , year=","venue":"arXiv (Cornell University)","work_id":"871bc70b-ac57-442c-892b-f8d06c9abab2","year":2025},"citing_paper":{"arxiv_id":"2502.12120","last_updated":"2026-05-20T13:58:03Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-02-17T18:45:25Z","title":"LLMs on the Line: Data Determines Loss-to-Loss Scaling Laws","version":3},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-05-23T02:47:37.492619Z"},"links":{"cited_paper":"/paper/2403.15796","citing_paper":"/paper/2502.12120"},"observation_digest":"sha256:697374d65a197692ae560f3fdf055ff3ef881e1a882f92a7592a54fe3b20dd2b","observation_id":"1b5153ba-a9e3-4f99-8924-8726896eb432","resolution":{"observed_at":"2026-05-23T02:52:27.109203Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.15796","last_updated":"2025-01-15T02:48:59Z","snapshot_observed_at":"2026-08-04T02:07:01.235387Z","submitted_at":"2024-03-23T11:03:31Z","title":"Understanding Emergent Abilities of Language Models from the Loss Perspective","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.15796","snapshot_observed_at":"2026-08-07T00:42:00.296033Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13216","last_updated":"2025-06-16T08:16:03Z","snapshot_observed_at":"2026-08-09T06:43:04.058790Z","submitted_at":"2025-06-16T08:16:03Z","title":"Capability Salience Vector: Fine-grained Alignment of Loss and Capabilities for Downstream Task Scaling Law","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-07T00:42:00.296033Z"},"links":{"cited_paper":"/paper/2403.15796","citing_paper":"/paper/2506.13216"},"observation_digest":"sha256:3b5303af3d6bc626f058689a2cacd958e7b26c7da380437ff056c95d440ce66f","observation_id":"8dae9bd0-9a5c-419a-ba21-37e4a47687ba","resolution":{"observed_at":"2026-08-07T00:42:00.296033Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.15796","last_updated":"2025-01-15T02:48:59Z","snapshot_observed_at":"2026-08-04T02:07:01.235387Z","submitted_at":"2024-03-23T11:03:31Z","title":"Understanding Emergent Abilities of Language Models from the Loss Perspective","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.15796","snapshot_observed_at":"2026-08-06T18:16:47.823217Z","title":"Understanding emergent abilities of language models from the loss perspective.arXiv preprint arXiv:2403.15796, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.08920","last_updated":"2026-06-05T11:04:35Z","snapshot_observed_at":"2026-08-08T00:11:49.840114Z","submitted_at":"2025-07-11T17:02:25Z","title":"AMix-1: A Pathway to Test-Time Scalable Protein Foundation Model","version":4},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-06T18:16:47.823217Z"},"links":{"cited_paper":"/paper/2403.15796","citing_paper":"/paper/2507.08920"},"observation_digest":"sha256:d439435d4bf64bfbcf58fae7ea51df470cc69004f93ec935f4228b6a31380672","observation_id":"d7fbd434-b4e6-48a5-9caf-54111cd0313b","resolution":{"observed_at":"2026-08-06T18:16:47.823217Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.15796","last_updated":"2025-01-15T02:48:59Z","snapshot_observed_at":"2026-08-04T02:07:01.235387Z","submitted_at":"2024-03-23T11:03:31Z","title":"Understanding Emergent Abilities of Language Models from the Loss Perspective","version":3},"cited_work":{"arxiv_id":"2403.15796","doi":"10.48550/arxiv.2403.15796","metadata_source":"arxiv_reference","pith_arxiv_id":"2403.15796","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2403.15796 , year=","venue":"arXiv (Cornell University)","work_id":"871bc70b-ac57-442c-892b-f8d06c9abab2","year":2025},"citing_paper":{"arxiv_id":"2605.17767","last_updated":"2026-05-21T20:45:44Z","snapshot_observed_at":"2026-08-04T06:32:13.198687Z","submitted_at":"2026-05-18T02:37:50Z","title":"Feature Learning in Linear-Width Two-Layer Networks: Two vs. One Step of Gradient Descent","version":1},"reference_index":213,"source":"arxiv_source","source_observed_at":"2026-05-20T01:29:14.555216Z"},"links":{"cited_paper":"/paper/2403.15796","citing_paper":"/paper/2605.17767"},"observation_digest":"sha256:66e5d7d5570ca436a8e30f8ff048e019138e8492cc70f05c0bc4f1ec23204f5f","observation_id":"743a70bc-e96e-4fe1-9ae3-8f5d5265b47f","resolution":{"observed_at":"2026-05-20T01:32:56.072235Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.15796","last_updated":"2025-01-15T02:48:59Z","snapshot_observed_at":"2026-08-04T02:07:01.235387Z","submitted_at":"2024-03-23T11:03:31Z","title":"Understanding Emergent Abilities of Language Models from the Loss Perspective","version":3},"cited_work":{"arxiv_id":"2403.15796","doi":"10.48550/arxiv.2403.15796","metadata_source":"arxiv_reference","pith_arxiv_id":"2403.15796","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2403.15796 , year=","venue":"arXiv (Cornell University)","work_id":"871bc70b-ac57-442c-892b-f8d06c9abab2","year":2025},"citing_paper":{"arxiv_id":"2605.17767","last_updated":"2026-05-21T20:45:44Z","snapshot_observed_at":"2026-08-04T06:32:13.198687Z","submitted_at":"2026-05-18T02:37:50Z","title":"Feature Learning in Linear-Width Two-Layer Networks: Two vs. One Step of Gradient Descent","version":2},"reference_index":213,"source":"arxiv_source","source_observed_at":"2026-05-25T06:39:16.246591Z"},"links":{"cited_paper":"/paper/2403.15796","citing_paper":"/paper/2605.17767"},"observation_digest":"sha256:c7956d96655a715b1f3ebaa8fc4872f84b32e3c2a7aa2569f4f48de7dfccc1fd","observation_id":"baebb411-75cc-44f2-bddd-99bcf941aa9b","resolution":{"observed_at":"2026-05-25T06:40:24.962878Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.15796","last_updated":"2025-01-15T02:48:59Z","snapshot_observed_at":"2026-08-04T02:07:01.235387Z","submitted_at":"2024-03-23T11:03:31Z","title":"Understanding Emergent Abilities of Language Models from the Loss Perspective","version":3},"cited_work":{"arxiv_id":"2403.15796","doi":"10.48550/arxiv.2403.15796","metadata_source":"arxiv_reference","pith_arxiv_id":"2403.15796","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2403.15796 , year=","venue":"arXiv (Cornell University)","work_id":"871bc70b-ac57-442c-892b-f8d06c9abab2","year":2025},"citing_paper":{"arxiv_id":"2605.19195","last_updated":"2026-05-18T23:51:02Z","snapshot_observed_at":"2026-07-06T23:29:57.003383Z","submitted_at":"2026-05-18T23:51:02Z","title":"The Thermodynamic Costs of Simple Linear Regression","version":1},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-05-20T07:03:21.987679Z"},"links":{"cited_paper":"/paper/2403.15796","citing_paper":"/paper/2605.19195"},"observation_digest":"sha256:cc6fd98fc4a0339771169bac7ff65314f617602d22b6d9953f0addd97420af4d","observation_id":"f42199a4-1035-428f-9c13-99fc0995c091","resolution":{"observed_at":"2026-05-20T07:03:23.237151Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.15796","last_updated":"2025-01-15T02:48:59Z","snapshot_observed_at":"2026-08-04T02:07:01.235387Z","submitted_at":"2024-03-23T11:03:31Z","title":"Understanding Emergent Abilities of Language Models from the Loss Perspective","version":3},"cited_work":{"arxiv_id":"2403.15796","doi":"10.48550/arxiv.2403.15796","metadata_source":"arxiv_reference","pith_arxiv_id":"2403.15796","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2403.15796 , year=","venue":"arXiv (Cornell University)","work_id":"871bc70b-ac57-442c-892b-f8d06c9abab2","year":2025},"citing_paper":{"arxiv_id":"2606.22873","last_updated":"2026-06-25T18:44:01Z","snapshot_observed_at":"2026-08-02T23:29:21.699637Z","submitted_at":"2026-06-22T05:37:43Z","title":"SingGuard: A Policy-Adaptive Multimodal LLM Guardrail with Dynamic Reasoning","version":2},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-06-26T09:19:50.623741Z"},"links":{"cited_paper":"/paper/2403.15796","citing_paper":"/paper/2606.22873"},"observation_digest":"sha256:429a22809bd2b64bd29cf6101d2eb6245d80573e0d483a2b4deeda8bd1247274","observation_id":"69de36b4-5800-460a-810d-8f6305cf4edb","resolution":{"observed_at":"2026-07-04T09:59:44.688100Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.15796","last_updated":"2025-01-15T02:48:59Z","snapshot_observed_at":"2026-08-04T02:07:01.235387Z","submitted_at":"2024-03-23T11:03:31Z","title":"Understanding Emergent Abilities of Language Models from the Loss Perspective","version":3},"cited_work":{"arxiv_id":"2403.15796","doi":"10.48550/arxiv.2403.15796","metadata_source":"arxiv_reference","pith_arxiv_id":"2403.15796","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2403.15796 , year=","venue":"arXiv (Cornell University)","work_id":"871bc70b-ac57-442c-892b-f8d06c9abab2","year":2025},"citing_paper":{"arxiv_id":"2606.22873","last_updated":"2026-06-25T18:44:01Z","snapshot_observed_at":"2026-08-02T23:29:21.699637Z","submitted_at":"2026-06-22T05:37:43Z","title":"SingGuard: A Policy-Adaptive Multimodal LLM Guardrail with Dynamic Reasoning","version":3},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-06-29T01:18:19.195007Z"},"links":{"cited_paper":"/paper/2403.15796","citing_paper":"/paper/2606.22873"},"observation_digest":"sha256:3ccca6a584b4ee1774becb0db6f00d14cb9450e946606da8957aa0cd45e4b750","observation_id":"527d4801-de94-4f6e-bbff-b3961bfd114c","resolution":{"observed_at":"2026-07-01T18:55:59.667241Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.15796","last_updated":"2025-01-15T02:48:59Z","snapshot_observed_at":"2026-08-04T02:07:01.235387Z","submitted_at":"2024-03-23T11:03:31Z","title":"Understanding Emergent Abilities of Language Models from the Loss Perspective","version":3},"cited_work":{"arxiv_id":"2403.15796","doi":"10.48550/arxiv.2403.15796","metadata_source":"arxiv_reference","pith_arxiv_id":"2403.15796","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2403.15796 , year=","venue":"arXiv (Cornell University)","work_id":"871bc70b-ac57-442c-892b-f8d06c9abab2","year":2025},"citing_paper":{"arxiv_id":"2606.30815","last_updated":"2026-06-29T18:42:03Z","snapshot_observed_at":"2026-08-06T19:03:34.282099Z","submitted_at":"2026-06-29T18:42:03Z","title":"When transformers learn \"impossible\" languages, what do they learn?","version":1},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-07-01T02:13:58.839175Z"},"links":{"cited_paper":"/paper/2403.15796","citing_paper":"/paper/2606.30815"},"observation_digest":"sha256:b425f7cd53bb5857f2e2bbf08989ad225d36b6d207ed4012c4c6e86b7a9173bc","observation_id":"586b36fa-c9c2-4a24-a13e-89909a3a01da","resolution":{"observed_at":"2026-07-01T02:15:14.469922Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.15796","last_updated":"2025-01-15T02:48:59Z","snapshot_observed_at":"2026-08-04T02:07:01.235387Z","submitted_at":"2024-03-23T11:03:31Z","title":"Understanding Emergent Abilities of Language Models from the Loss Perspective","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.15796","snapshot_observed_at":"2026-08-04T04:31:48.703352Z","title":"Understanding Emergent Abilities of Language Models from the Loss Perspective.arXiv e-prints, page arXiv:2403.15796, March 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.02573","last_updated":"2026-08-03T17:49:40Z","snapshot_observed_at":"2026-08-08T09:09:06.432222Z","submitted_at":"2026-08-03T17:49:40Z","title":"Foundation Models for Astrophysics","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-04T04:31:48.703352Z"},"links":{"cited_paper":"/paper/2403.15796","citing_paper":"/paper/2608.02573"},"observation_digest":"sha256:cd5bac9a18c8159442008fe75c50d8f8d3959f1452382a9845e7fcc1a062a467","observation_id":"2e698ea0-b7b7-4b00-8086-b84d28f5ab9f","resolution":{"observed_at":"2026-08-04T04:31:48.703352Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2403.15796/citation-record","integrity":"/paper/2403.15796/integrity","json":"/paper/2403.15796/citation-record.json","paper":"/paper/2403.15796"},"outbound":[],"paper":{"arxiv_id":"2403.15796","last_updated":"2025-01-15T02:48:59Z","latest_version":3,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-04T02:07:01.235387Z","submitted_at":"2024-03-23T11:03:31Z","title":"Understanding Emergent Abilities of Language Models from the Loss Perspective"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 13 inbound Pith citation observations for arXiv:2403.15796."}