{"as_of":"2026-08-05T14:19:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:7c7ad2ec420fb125ff34eb1de2202e8a6f1b806a15a54dc5c09dcd0fd63e21f8","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":8,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":8,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-05T06:32:48.257954+00:00","state":"measured"},{"denominator":8,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":8,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-03T05:27:44.529687Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-07-06T19:29:18.770662Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-03T05:27:44.529687Z","title":"From sparse dependence to sparse attention: unveiling how chain-of-thought enhances transformer sample efficiency.arXiv preprint arXiv:2410.05459,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2602.02470","last_updated":"2026-05-31T23:31:59Z","snapshot_observed_at":"2026-08-03T06:29:57.525520Z","submitted_at":"2026-02-02T18:50:57Z","title":"Breaking the Reversal Curse in Autoregressive Language Models via Identity Bridge","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-03T05:27:44.529687Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2602.02470"},"observation_digest":"sha256:c83215aa64362adf38a3ee1bede0664c137c6e3689619a0e1be66ba42e2b0bf0","observation_id":"42f0197e-30c9-46bb-a399-aea76b88e17d","resolution":{"observed_at":"2026-08-03T05:27:44.529687Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-07-06T19:29:18.770662Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-02T23:11:56.681936Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2602.14872","last_updated":"2026-06-29T01:04:09Z","snapshot_observed_at":"2026-08-02T23:11:47.613570Z","submitted_at":"2026-02-16T16:03:08Z","title":"On the Emergence of Implicit Curriculum in RLVR Learning Dynamics","version":3},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-02T23:11:56.681936Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2602.14872"},"observation_digest":"sha256:235ad1eee2ea4a3d37538bbdaa545fbfd0358b909f69401a2fe1c325f0027a65","observation_id":"4f8fbc10-55d4-40ce-82d7-6c4d511cdeff","resolution":{"observed_at":"2026-08-02T23:11:56.681936Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-07-06T19:29:18.770662Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":"2410.05459","doi":"10.48550/arxiv.2410.05459","metadata_source":"arxiv_reference","pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"From sparse dependence to sparse attention: Unveiling how chain-of-thought enhances transformer sample efficiency.ArXiv, abs/2410.05459","venue":"arXiv (Cornell University)","work_id":"12dec1cc-a66a-4ef5-bfe7-061e33762461","year":2024},"citing_paper":{"arxiv_id":"2604.22951","last_updated":"2026-07-08T19:29:03Z","snapshot_observed_at":"2026-07-12T23:17:30.297545Z","submitted_at":"2026-04-24T18:49:08Z","title":"The Power of Power Law: Asymmetry Enables Compositional Reasoning","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-05-08T11:49:49.787123Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2604.22951"},"observation_digest":"sha256:166d3a2692f890dddfad6b46d2d4e1fa6ecd7dd6a3c338b15b61210515a4c2a5","observation_id":"34b2affc-a800-43ae-86ad-2d78e95b882a","resolution":{"observed_at":"2026-05-11T19:31:08.381738Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-07-06T19:29:18.770662Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-07-12T18:26:05.728364Z","title":"From sparse dependence to sparse attention: unveiling how chain-of-thought enhances transformer sample efficiency","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2604.22951","last_updated":"2026-07-08T19:29:03Z","snapshot_observed_at":"2026-07-12T23:17:30.297545Z","submitted_at":"2026-04-24T18:49:08Z","title":"The Power of Power Law: Asymmetry Enables Compositional Reasoning","version":2},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-07-12T18:26:05.728364Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2604.22951"},"observation_digest":"sha256:7cfaac187f556d5bb2d5a41a21d656d82521b7bf23c1110d7191c6f8798b05a3","observation_id":"8a1e30ca-2190-41fe-84d1-9b468517f84d","resolution":{"observed_at":"2026-07-12T18:26:05.728364Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-07-06T19:29:18.770662Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":"2410.05459","doi":"10.48550/arxiv.2410.05459","metadata_source":"arxiv_reference","pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"From sparse dependence to sparse attention: Unveiling how chain-of-thought enhances transformer sample efficiency.ArXiv, abs/2410.05459","venue":"arXiv (Cornell University)","work_id":"12dec1cc-a66a-4ef5-bfe7-061e33762461","year":2024},"citing_paper":{"arxiv_id":"2605.10019","last_updated":"2026-05-11T05:44:18Z","snapshot_observed_at":"2026-07-06T23:22:03.830881Z","submitted_at":"2026-05-11T05:44:18Z","title":"The two clocks and the innovation window: When and how generative models learn rules","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-05-12T03:15:45.257213Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2605.10019"},"observation_digest":"sha256:1a1f9845e841f73b6de89556ac04d39575e10d0dc2a29fb025be69046912a289","observation_id":"fe5c8a4d-2b15-42e1-a141-fe0270cbd952","resolution":{"observed_at":"2026-05-12T03:16:18.110627Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-07-06T19:29:18.770662Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":"2410.05459","doi":"10.48550/arxiv.2410.05459","metadata_source":"arxiv_reference","pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"From sparse dependence to sparse attention: Unveiling how chain-of-thought enhances transformer sample efficiency.ArXiv, abs/2410.05459","venue":"arXiv (Cornell University)","work_id":"12dec1cc-a66a-4ef5-bfe7-061e33762461","year":2024},"citing_paper":{"arxiv_id":"2605.23040","last_updated":"2026-05-21T21:13:14Z","snapshot_observed_at":"2026-08-02T04:36:31.374601Z","submitted_at":"2026-05-21T21:13:14Z","title":"Steered Generation via Gradient-Based Optimization on Sparse Query Features","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-05-25T05:31:29.510639Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2605.23040"},"observation_digest":"sha256:5665e2ee06d51e16b196b4960866ffe5b8fb88f742c07805a2ea74cfeeb4d1fe","observation_id":"4d011ce6-1534-47bc-9fb4-589d17a05c06","resolution":{"observed_at":"2026-05-25T05:36:40.299551Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-07-06T19:29:18.770662Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":"2410.05459","doi":"10.48550/arxiv.2410.05459","metadata_source":"arxiv_reference","pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"From sparse dependence to sparse attention: Unveiling how chain-of-thought enhances transformer sample efficiency.ArXiv, abs/2410.05459","venue":"arXiv (Cornell University)","work_id":"12dec1cc-a66a-4ef5-bfe7-061e33762461","year":2024},"citing_paper":{"arxiv_id":"2605.28600","last_updated":"2026-05-27T15:17:06Z","snapshot_observed_at":"2026-08-01T17:58:43.396208Z","submitted_at":"2026-05-27T15:17:06Z","title":"Transformers Provably Learn to Internalize Chain-of-Thought","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-06-29T14:29:10.010212Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2605.28600"},"observation_digest":"sha256:2479b10684f50b0725a38fcdcbcf3a14d4844fc08afbe28f64229ad869ca1adf","observation_id":"eb03379a-a679-44aa-9a4b-782d9255ef4b","resolution":{"observed_at":"2026-06-29T14:33:30.594670Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","snapshot_observed_at":"2026-07-06T19:29:18.770662Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency","version":2},"cited_work":{"arxiv_id":"2410.05459","doi":"10.48550/arxiv.2410.05459","metadata_source":"arxiv_reference","pith_arxiv_id":"2410.05459","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"From sparse dependence to sparse attention: Unveiling how chain-of-thought enhances transformer sample efficiency.ArXiv, abs/2410.05459","venue":"arXiv (Cornell University)","work_id":"12dec1cc-a66a-4ef5-bfe7-061e33762461","year":2024},"citing_paper":{"arxiv_id":"2606.00183","last_updated":"2026-05-29T14:58:03Z","snapshot_observed_at":"2026-07-06T23:40:56.510371Z","submitted_at":"2026-05-29T14:58:03Z","title":"Agentic Transformers Provably Learn to Search via Reinforcement Learning","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-06-28T23:26:28.158991Z"},"links":{"cited_paper":"/paper/2410.05459","citing_paper":"/paper/2606.00183"},"observation_digest":"sha256:b219bc09b2e4e3d3bd70676d47e437f7cccfdae40bee07deaf16025fe5c6ee96","observation_id":"9d47d62b-7dbe-4f2f-90bf-83035e95c9f1","resolution":{"observed_at":"2026-06-28T23:42:49.906343Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2410.05459/citation-record","integrity":"/paper/2410.05459/integrity","json":"/paper/2410.05459/citation-record.json","paper":"/paper/2410.05459"},"outbound":[],"paper":{"arxiv_id":"2410.05459","last_updated":"2025-03-05T13:57:56Z","latest_version":2,"primary_category":"cs.LG","snapshot_observed_at":"2026-07-06T19:29:18.770662Z","submitted_at":"2024-10-07T19:45:09Z","title":"From Sparse Dependence to Sparse Attention: Unveiling How Chain-of-Thought Enhances Transformer Sample Efficiency"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"thesis":"As of 5 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 8 inbound Pith citation observations for arXiv:2410.05459."}