{"as_of":"2026-08-09T05:25:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:6032e7e679f3dad877f2d02403fb70dd76c76258f008e9f3b68e57edb8f5e760","coverage":[{"denominator":38,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":38,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T14:32:34.342638Z","state":"measured"},{"denominator":44,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":44,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":6,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":6,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T00:13:37.810334Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.18642","snapshot_observed_at":"2026-08-07T00:13:37.810334Z","title":"Skip-thinking: Chunk-wise chain-of-thought distillation enable smaller language models to reason better and faster.arXiv preprint arXiv:2505.18642, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.14728","last_updated":"2025-06-17T17:08:32Z","snapshot_observed_at":"2026-08-07T14:16:24.457276Z","submitted_at":"2025-06-17T17:08:32Z","title":"AgentDistill: Training-Free Agent Distillation with Generalizable MCP Boxes","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T00:13:37.810334Z"},"links":{"cited_paper":"/paper/2505.18642","citing_paper":"/paper/2506.14728"},"observation_digest":"sha256:0a19ef991eda256b35d1c659e92da8622433b9857631bafbc8c038b92edcd2cd","observation_id":"d0a7bef4-9d7f-4a31-ae40-eb5e419f8910","resolution":{"observed_at":"2026-08-07T00:13:37.810334Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"cited_work":{"arxiv_id":"2505.18642","doi":"10.48550/arxiv.2505.18642","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.18642","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Skip-thinking: Chunk-wise chain- of-thought distillation enable smaller language mod- els to reason better and faster","venue":"arXiv (Cornell University)","work_id":"aa551551-718d-4a6c-a242-f6b83dd0060f","year":2025},"citing_paper":{"arxiv_id":"2605.10195","last_updated":"2026-05-14T07:42:56Z","snapshot_observed_at":"2026-08-03T09:42:12.047946Z","submitted_at":"2026-05-11T08:45:17Z","title":"Breaking the Reward Barrier: Accelerating Tree-of-Thought Reasoning via Speculative Exploration","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-12T03:51:52.375703Z"},"links":{"cited_paper":"/paper/2505.18642","citing_paper":"/paper/2605.10195"},"observation_digest":"sha256:a31aab02e74f877e97c1752fcefd9d2b248f7dc6f058d60d73e58ce90e965e8d","observation_id":"1a47abdb-7c81-4328-9ae7-6f72296d93fc","resolution":{"observed_at":"2026-05-12T06:51:30.000222Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"cited_work":{"arxiv_id":"2505.18642","doi":"10.48550/arxiv.2505.18642","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.18642","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Skip-thinking: Chunk-wise chain- of-thought distillation enable smaller language mod- els to reason better and faster","venue":"arXiv (Cornell University)","work_id":"aa551551-718d-4a6c-a242-f6b83dd0060f","year":2025},"citing_paper":{"arxiv_id":"2605.10195","last_updated":"2026-05-14T07:42:56Z","snapshot_observed_at":"2026-08-03T09:42:12.047946Z","submitted_at":"2026-05-11T08:45:17Z","title":"Breaking the Reward Barrier: Accelerating Tree-of-Thought Reasoning via Speculative Exploration","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-15T05:11:32.053440Z"},"links":{"cited_paper":"/paper/2505.18642","citing_paper":"/paper/2605.10195"},"observation_digest":"sha256:75179d1732618a53254fff66d3b90c6bab9db9cc8cd9895f4c5aec6e2e74da45","observation_id":"58bcff32-b9c2-4a4a-92fd-09f45af9bde2","resolution":{"observed_at":"2026-05-15T05:15:03.384390Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"cited_work":{"arxiv_id":"2505.18642","doi":"10.48550/arxiv.2505.18642","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.18642","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Skip-thinking: Chunk-wise chain- of-thought distillation enable smaller language mod- els to reason better and faster","venue":"arXiv (Cornell University)","work_id":"aa551551-718d-4a6c-a242-f6b83dd0060f","year":2025},"citing_paper":{"arxiv_id":"2606.06840","last_updated":"2026-06-05T02:32:24Z","snapshot_observed_at":"2026-08-05T15:03:33.373071Z","submitted_at":"2026-06-05T02:32:24Z","title":"Characterize Then Distill: Mechanistic Reasoning in Large Output Spaces","version":1},"reference_index":79,"source":"arxiv_source","source_observed_at":"2026-06-27T22:22:52.690010Z"},"links":{"cited_paper":"/paper/2505.18642","citing_paper":"/paper/2606.06840"},"observation_digest":"sha256:f253546dad6692a6f54e5e7bb9d5c03d2579faa4cc1a69d6a7b50e247c46c0fa","observation_id":"8559c368-79af-44b5-9d67-77b885838708","resolution":{"observed_at":"2026-06-27T22:31:21.381039Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"cited_work":{"arxiv_id":"2505.18642","doi":"10.48550/arxiv.2505.18642","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.18642","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Skip-thinking: Chunk-wise chain- of-thought distillation enable smaller language mod- els to reason better and faster","venue":"arXiv (Cornell University)","work_id":"aa551551-718d-4a6c-a242-f6b83dd0060f","year":2025},"citing_paper":{"arxiv_id":"2606.20937","last_updated":"2026-06-18T21:01:42Z","snapshot_observed_at":"2026-08-02T17:33:10.336044Z","submitted_at":"2026-06-18T21:01:42Z","title":"Learning through Internalization","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-06-26T17:43:18.915404Z"},"links":{"cited_paper":"/paper/2505.18642","citing_paper":"/paper/2606.20937"},"observation_digest":"sha256:68d490891f5d9e42a0d15568dc12029d9fcbcb7fcba6530af46285b77103cea4","observation_id":"cafc8e93-8a31-4e4a-83b7-b5d0cd19a3a1","resolution":{"observed_at":"2026-07-04T03:39:31.035917Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"cited_work":{"arxiv_id":"2505.18642","doi":"10.48550/arxiv.2505.18642","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.18642","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Skip-thinking: Chunk-wise chain- of-thought distillation enable smaller language mod- els to reason better and faster","venue":"arXiv (Cornell University)","work_id":"aa551551-718d-4a6c-a242-f6b83dd0060f","year":2025},"citing_paper":{"arxiv_id":"2607.02234","last_updated":"2026-07-02T14:33:07Z","snapshot_observed_at":"2026-07-07T00:07:42.752664Z","submitted_at":"2026-07-02T14:33:07Z","title":"Purified OPSD: On-Policy Self-Distillation Without Losing How to Think","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-03T13:56:13.827493Z"},"links":{"cited_paper":"/paper/2505.18642","citing_paper":"/paper/2607.02234"},"observation_digest":"sha256:307dcbed5958925984bf896234f89f0ae1feebf098b3453de6d28f5af130eb78","observation_id":"b0a8c6a0-2948-479a-8305-dde54ced0341","resolution":{"observed_at":"2026-07-03T13:58:21.079234Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2505.18642/citation-record","integrity":"/paper/2505.18642/integrity","json":"/paper/2505.18642/citation-record.json","paper":"/paper/2505.18642"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2404.09170","last_updated":"2025-08-07T03:25:23Z","snapshot_observed_at":"2026-07-06T17:59:55.100798Z","submitted_at":"2024-04-14T07:19:27Z","title":"Distilling Reasoning Ability from Large Language Models with Adaptive Thinking","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.09170","snapshot_observed_at":"2026-08-07T14:32:31.503098Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:31.503098Z"},"links":{"cited_paper":"/paper/2404.09170","citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:0b47a977e0ef3cd4687c65f379dbedfe84509c0c049c3f469c7bdf0f75dced22","observation_id":"6d99eafc-6c8c-44e2-8c1a-0a68672eb086","resolution":{"observed_at":"2026-08-07T14:32:31.503098Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2024.findings-acl.409","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:34.658406Z","title":null,"venue":null,"work_id":"80206773-d481-4e10-8042-50b0fa2f4a7f","year":2024},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:31.558018Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:4b089af0c34a315fb58ae2238ac928efdf97d159467321c659b3cf3aa1229bd9","observation_id":"9e174d96-a622-4671-b835-5a5824b05088","resolution":{"observed_at":"2026-08-07T14:32:34.724426Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.15402","last_updated":"2024-06-06T01:58:54Z","snapshot_observed_at":"2026-07-06T16:24:10.053246Z","submitted_at":"2023-09-27T04:53:10Z","title":"Navigate through Enigmatic Labyrinth A Survey of Chain of Thought Reasoning: Advances, Frontiers and Future","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.15402","snapshot_observed_at":"2026-08-07T14:32:31.643348Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:31.643348Z"},"links":{"cited_paper":"/paper/2309.15402","citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:2715e1d71d25c8e0c4c7b08af1ff2fdcc48e8809253863ae0ea8c6b86a31d998","observation_id":"d01ca7da-75aa-419f-800c-cb6e3fc01c68","resolution":{"observed_at":"2026-08-07T14:32:31.643348Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2110.14168","last_updated":"2021-11-18T00:23:45Z","snapshot_observed_at":"2026-08-07T01:45:38.840969Z","submitted_at":"2021-10-27T04:49:45Z","title":"Training Verifiers to Solve Math Word Problems","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2110.14168","snapshot_observed_at":"2026-08-07T14:32:31.717800Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:31.717800Z"},"links":{"cited_paper":"/paper/2110.14168","citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:48e737b394a2d762375f95144d20a8d50bb96ba49efd6a12baafdd561f44a84b","observation_id":"3dc65d92-8081-4e92-acb2-71449ba4d41b","resolution":{"observed_at":"2026-08-07T14:32:31.717800Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.14838","last_updated":"2024-05-23T17:54:14Z","snapshot_observed_at":"2026-07-06T18:18:51.814306Z","submitted_at":"2024-05-23T17:54:14Z","title":"From Explicit CoT to Implicit CoT: Learning to Internalize CoT Step by Step","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.14838","snapshot_observed_at":"2026-08-07T14:32:31.805993Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:31.805993Z"},"links":{"cited_paper":"/paper/2405.14838","citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:228e6e7020f38c4f8921768acd4e694d2b43396666eb9ffdd078f1cacd50173e","observation_id":"01d02b20-c56a-4130-8431-2789940c381e","resolution":{"observed_at":"2026-08-07T14:32:31.805993Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.01460","last_updated":"2023-11-02T17:59:49Z","snapshot_observed_at":"2026-08-05T13:45:13.159592Z","submitted_at":"2023-11-02T17:59:49Z","title":"Implicit Chain of Thought Reasoning via Knowledge Distillation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.01460","snapshot_observed_at":"2026-08-07T14:32:31.908544Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:31.908544Z"},"links":{"cited_paper":"/paper/2311.01460","citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:4fb1c1802612f7e1580fc4c64d066834c2ab38bd27dd027e9ad0d6b82adda398","observation_id":"5246b3fc-bbc6-44cb-9fc4-33fe1c71fd3e","resolution":{"observed_at":"2026-08-07T14:32:31.908544Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-07-06T18:55:11.576666Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-07T14:32:31.984908Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:31.984908Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:78691ba8ee1427f8d63645b01b82bcaddaa4433d94a2afaf76b57ec9ad4c139b","observation_id":"b883cf0f-9e8e-4b88-8de7-0421223d25bb","resolution":{"observed_at":"2026-08-07T14:32:31.984908Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2012.08795","last_updated":"2020-12-16T08:43:15Z","snapshot_observed_at":"2026-08-08T15:46:57.928453Z","submitted_at":"2020-12-16T08:43:15Z","title":"Study on the Large Batch Size Training of Neural Networks Based on the Second Order Gradient","version":1},"cited_work":{"arxiv_id":"2012.08795","doi":null,"metadata_source":"pith","pith_arxiv_id":"2012.08795","snapshot_observed_at":"2026-08-07T14:32:34.825870Z","title":"Study on the Large Batch Size Training of Neural Networks Based on the Second Order Gradient","venue":"cs.LG","work_id":"500c4f06-33af-43a9-8513-c10707761709","year":2020},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:32.054455Z"},"links":{"cited_paper":"/paper/2012.08795","citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:7266705f7ea1f4b91a3a233de281c9d388c2dca190ca3de49458c4f709e78cf9","observation_id":"27ab3df6-aaa9-45ca-ac77-3db2d1225503","resolution":{"observed_at":"2026-08-07T14:32:34.886189Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:32.144209Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:32.144209Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:6d66ffac357a04f90dacaf388543abacb0f054455677c7a8ae898d81c5ab0a78","observation_id":"ede6edc7-7b29-41b9-8227-6a59ee9f94c2","resolution":{"observed_at":"2026-08-07T14:32:32.144209Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.02226","last_updated":"2024-04-21T03:39:21Z","snapshot_observed_at":"2026-08-06T21:01:03.193755Z","submitted_at":"2023-10-03T17:32:41Z","title":"Think before you speak: Training Language Models With Pause Tokens","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.02226","snapshot_observed_at":"2026-08-07T14:32:32.230744Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:32.230744Z"},"links":{"cited_paper":"/paper/2310.02226","citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:3d3bb54834772bcd0a43366f9efc89419fc8e9de1b654412fd03a6d88d719078","observation_id":"4a505697-7948-4a32-933e-b00d2a86094f","resolution":{"observed_at":"2026-08-07T14:32:32.230744Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.06769","last_updated":"2025-11-03T00:53:34Z","snapshot_observed_at":"2026-08-07T06:05:27.895209Z","submitted_at":"2024-12-09T18:55:56Z","title":"Training Large Language Models to Reason in a Continuous Latent Space","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.06769","snapshot_observed_at":"2026-08-07T14:32:32.350794Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:32.350794Z"},"links":{"cited_paper":"/paper/2412.06769","citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:be09a12b498a7964ebabefc5b68111057fd8cd84f044c400fa9a74f4a2197b55","observation_id":"cd1e36d7-70e2-4dda-a8e4-f9d5149e7380","resolution":{"observed_at":"2026-08-07T14:32:32.350794Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:32.435023Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:32.435023Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:9e004c642b5987743e19057880f7288b0246471a339b467fe9898fc760101ae1","observation_id":"eb079afb-7239-466c-a1e7-7a26c74298e7","resolution":{"observed_at":"2026-08-07T14:32:32.435023Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:32.548924Z","title":null,"venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:32.548924Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:b6d446ebb9c40ccc366b4a71b1e33fa02bc3e78f9327b268dcecd076beed42a3","observation_id":"5df18576-50b5-4414-abf9-f6104210ba18","resolution":{"observed_at":"2026-08-07T14:32:32.548924Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:32.669252Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:32.669252Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:2279f5c2f5666d9a24745990b15e48d8e9b1ddb07608bd6654f73affd05693cc","observation_id":"52f65f3d-c399-4d4b-831c-2ded076b7cb2","resolution":{"observed_at":"2026-08-07T14:32:32.669252Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:32.742397Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:32.742397Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:b30d351157df4f5f819662a317ac2c277d72f532e54137e5f69feada8bed06a3","observation_id":"53d320f2-3534-4f51-bd81-6a71a886d9b9","resolution":{"observed_at":"2026-08-07T14:32:32.742397Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:35.537541Z","title":null,"venue":null,"work_id":"581cad03-cd49-4344-9bfa-ae1eb187e122","year":2018},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:32.820855Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:e4e2584bb8df93ce06fac6b15f574731edcc634fed3d2143a4def6f4cfb72419","observation_id":"d1c58408-d82b-47b0-816c-6088de5b4e4e","resolution":{"observed_at":"2026-08-07T14:32:35.608584Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:35.389974Z","title":null,"venue":null,"work_id":"3fd87685-8ede-4ccd-b009-21b4faa4bd37","year":2017},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:32.918084Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:c39fced38764a446c100c3396b2e81cd9b70ea12d9de7dcaff54ead5a38b09d5","observation_id":"307afb48-6d99-44f1-b2a5-75ff227f334f","resolution":{"observed_at":"2026-08-07T14:32:35.435159Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:32.980107Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:32.980107Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:90da0fdd1f1ff054e145197d47269c3eba4a19ccfdaaecc3b8932bddbbb0538e","observation_id":"672d769c-9c79-4fec-98cd-612ee9fefe8e","resolution":{"observed_at":"2026-08-07T14:32:32.980107Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:33.044801Z","title":null,"venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:33.044801Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:2414dc5df1e73db041caec2cf494fd4bec02fe0da54b87732830f5b00a88243f","observation_id":"27bb550c-5b77-43cc-82a9-6ea65d9376fc","resolution":{"observed_at":"2026-08-07T14:32:33.044801Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.12832","last_updated":"2023-12-20T08:28:36Z","snapshot_observed_at":"2026-07-06T17:05:48.182547Z","submitted_at":"2023-12-20T08:28:36Z","title":"Turning Dust into Gold: Distilling Complex Reasoning Capabilities from LLMs by Leveraging Negative Data","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.12832","snapshot_observed_at":"2026-08-07T14:32:33.106718Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:33.106718Z"},"links":{"cited_paper":"/paper/2312.12832","citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:20793349f2caddbbf17d3e99eed6b97e5fbecf4e7a828fb6b71e9fd26327a62b","observation_id":"0f2aafa4-3aa2-4c22-8346-d19f8bd74dc2","resolution":{"observed_at":"2026-08-07T14:32:33.106718Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2024.naacl-long.376","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:34.520599Z","title":null,"venue":null,"work_id":"e3cc3dac-84ae-4f1d-9db7-905bd4df3de4","year":2024},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:33.170679Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:a6918869e2898f9668b2c08f002937bad8fed46cf1c5623dbe36efe27402a541","observation_id":"41393e79-4797-460e-9e6c-2126c43369e1","resolution":{"observed_at":"2026-08-07T14:32:34.575129Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2212.08410","last_updated":"2023-06-01T12:17:01Z","snapshot_observed_at":"2026-07-06T14:31:38.001143Z","submitted_at":"2022-12-16T11:24:42Z","title":"Teaching Small Language Models to Reason","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2212.08410","snapshot_observed_at":"2026-08-07T14:32:33.263968Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:33.263968Z"},"links":{"cited_paper":"/paper/2212.08410","citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:96a0c437947ad14cab0cc984bd280a15e75de44a853ccff1038ce50b028c6b34","observation_id":"8bfb1776-7273-40cf-a107-84bab7478474","resolution":{"observed_at":"2026-08-07T14:32:33.263968Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:33.367034Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:33.367034Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:518b30a8d724effd9524e5f10acbe4207369df9185d54bee800f44a29ab984c1","observation_id":"7d6b26e7-b3b5-45dd-a6ed-daeef7fe0e2f","resolution":{"observed_at":"2026-08-07T14:32:33.367034Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:33.462368Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:33.462368Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:0b9f81025d6e8b8332e47247131b8409a00ca7e9e843dd26bfc5d49eaf00654e","observation_id":"b1ae57a8-44a1-4d30-a538-603394c46291","resolution":{"observed_at":"2026-08-07T14:32:33.462368Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:33.589676Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:33.589676Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:50eafd1b639e6d7fc932b55ff140fe6edab767e6a25ca8eab6e2ee04e08e1625","observation_id":"289b59a8-b5a3-4acd-9c4c-0ce5829e32ba","resolution":{"observed_at":"2026-08-07T14:32:33.589676Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:33.666753Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:33.666753Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:0363e48437f9bed16efc724e8a40a46e521789685ca28235ebccea4ae7f1ccf9","observation_id":"5ea62235-4267-456b-94f8-34b09d876634","resolution":{"observed_at":"2026-08-07T14:32:33.666753Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2022.findings-naacl.169","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:34.440145Z","title":null,"venue":null,"work_id":"4c788c86-d27a-4d46-9cdd-1f8e54860a71","year":2022},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:33.726022Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:f94f9cfda04f52c610f7c445c1800c3720ff299dd519d50fead4b7f25c89abfb","observation_id":"0eab5223-d4dd-483b-ba5f-ab988a715394","resolution":{"observed_at":"2026-08-07T14:32:34.463223Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:33.789910Z","title":null,"venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:33.789910Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:3e1e1ac119239efc0fecfcc7c8774bd2b166871ad0ac3041404d938924a24585","observation_id":"f881ddcd-2956-437b-b2d8-1024ea47a6c7","resolution":{"observed_at":"2026-08-07T14:32:33.789910Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2206.04615","last_updated":"2023-06-12T17:51:15Z","snapshot_observed_at":"2026-07-06T13:19:12.109592Z","submitted_at":"2022-06-09T17:05:34Z","title":"Beyond the Imitation Game: Quantifying and extrapolating the capabilities of language models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2206.04615","snapshot_observed_at":"2026-08-07T14:32:33.838327Z","title":"Brown, Adam Santoro, and et al","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:33.838327Z"},"links":{"cited_paper":"/paper/2206.04615","citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:fd2f7f8babb38fe3bcbccc9e021e499fc8233697dad4bfe861063bc37258a1e8","observation_id":"ab553435-55ff-4ee2-9c69-259ebd20c697","resolution":{"observed_at":"2026-08-07T14:32:33.838327Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.09288","last_updated":"2023-07-19T17:08:59Z","snapshot_observed_at":"2026-08-07T12:56:43.323460Z","submitted_at":"2023-07-18T14:31:57Z","title":"Llama 2: Open Foundation and Fine-Tuned Chat Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.09288","snapshot_observed_at":"2026-08-07T14:32:33.878979Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:33.878979Z"},"links":{"cited_paper":"/paper/2307.09288","citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:4b6463ba02e4de4eac44e7a11e24707db4ccdba782ec2d482924656c8adcbeed","observation_id":"27fd3f08-dd0d-48ff-9103-935f0696944e","resolution":{"observed_at":"2026-08-07T14:32:33.878979Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:35.305628Z","title":null,"venue":null,"work_id":"c009da68-66f8-44ff-9078-7090eaf463ab","year":2024},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:33.938172Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:80661d598af7406410447f6bb57bd2160f6a7f67576c2910e22edb1bd3a8da04","observation_id":"781107d4-e822-4d69-86ad-a559f57f20e9","resolution":{"observed_at":"2026-08-07T14:32:35.328236Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:34.005045Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:34.005045Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:3bd625fc497263c201748325a6120661ab589d720e3cc51d00405b8eed44a5dd","observation_id":"4e6467aa-68ea-4545-9612-498d34d34c75","resolution":{"observed_at":"2026-08-07T14:32:34.005045Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:35.258657Z","title":null,"venue":null,"work_id":"93ff755b-3dd9-45b4-9c7f-696379317f1c","year":2022},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:34.072748Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:b67024dff9b36d35ff7ed2983b613fb14aa4d3cba74f9e0edecd2b1a3e815078","observation_id":"7d3374d9-f9ad-4d59-8c8a-8691c63aa5e3","resolution":{"observed_at":"2026-08-07T14:32:35.277994Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:35.126978Z","title":null,"venue":null,"work_id":"8ed50b1b-af16-4695-aa1b-35fa09923f65","year":2024},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:34.129695Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:345664c76878db48b1dfc5ac2e4e1eda7c4ae6c20cd880cc7e0c2dea8c422607","observation_id":"52fcadf2-f589-42e5-bf27-c8f9d486b1f7","resolution":{"observed_at":"2026-08-07T14:32:35.183805Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:35.014727Z","title":null,"venue":null,"work_id":"5ab6b7a2-d206-48ee-9e99-a0a3fd4b4d5c","year":2024},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:34.164762Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:cb1f19d831bbf270decf7e220dd8815a0d90090b170c54e8335ed41bd5610180","observation_id":"a3c0d7c2-2c54-46de-b531-9f5f9753b93d","resolution":{"observed_at":"2026-08-07T14:32:35.058522Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.13888","last_updated":"2024-03-20T08:37:42Z","snapshot_observed_at":"2026-08-08T03:59:52.512778Z","submitted_at":"2023-05-23T10:11:56Z","title":"PaD: Program-aided Distillation Can Teach Small Models Reasoning Better than Chain-of-thought Fine-tuning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.13888","snapshot_observed_at":"2026-08-07T14:32:34.203188Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:34.203188Z"},"links":{"cited_paper":"/paper/2305.13888","citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:3a2a543ab271d9fbaba4e3ab4a42fb62a7fd13397554072d5685127e25baa144","observation_id":"d6d1e800-8918-4325-9b43-ad8eeb96c8b2","resolution":{"observed_at":"2026-08-07T14:32:34.203188Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:34.256637Z","title":"online\" 'onlinestring :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:34.256637Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:c459fba6d73dbebcc6d0a0230d857beddde4e1603b3b0fe92b046eb1a8a25ea0","observation_id":"29e66df1-1453-4a51-bd6d-7b450692136d","resolution":{"observed_at":"2026-08-07T14:32:34.256637Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:32:34.342638Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-07T14:32:34.342638Z"},"links":{"citing_paper":"/paper/2505.18642"},"observation_digest":"sha256:0e5c5587e958457db004b438b68df071e4b9cbc0d6b649c41ab493d029f3cbf4","observation_id":"d102e5ff-a891-4022-9c70-2f5091c7e05c","resolution":{"observed_at":"2026-08-07T14:32:34.342638Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2505.18642","last_updated":"2025-05-24T11:04:52Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-07T15:16:08.846967Z","submitted_at":"2025-05-24T11:04:52Z","title":"Skip-Thinking: Chunk-wise Chain-of-Thought Distillation Enable Smaller Language Models to Reason Better and Faster"},"reference_resolution":{"displayed":38,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":34,"verified_exact":4,"verified_fuzzy":0},"total_outbound_references":38},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 38 of 38 outbound references and 6 inbound Pith citation observations for arXiv:2505.18642."}