{"as_of":"2026-08-12T22:20:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:8ebc5b8f97d9d7a0a9417077f4e44f97fbf6c712afced267dd9d07d63b41f478","coverage":[{"denominator":27,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":27,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T00:23:50.523918Z","state":"measured"},{"denominator":27,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":27,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-12T06:34:41.77262+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2608.08168/citation-record","integrity":"/paper/2608.08168/integrity","json":"/paper/2608.08168/citation-record.json","paper":"/paper/2608.08168"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-12T00:23:50.418419Z","title":"Deepseek-r1: Incentivizing reasoning capability in llms via reinforcement learning.arXiv preprint arXiv:2501.12948, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.418419Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:bb45fe983983501c1fa91a77c17dc39ede885db546abf58a438bce8f64875076","observation_id":"9cbcc900-a783-4fb0-9e27-92102fc38583","resolution":{"observed_at":"2026-08-12T00:23:50.418419Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T00:23:50.423415Z","title":"Chain-of-thought prompting elicits reasoning in large language models.Advances in neural information processing systems, 35:24824–24837, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.423415Z"},"links":{"citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:39d05079093b07a53e5d84250d35f7944b90f00e79aec112ca8b1e95756c0305","observation_id":"18ea0f4e-bc63-4f89-859c-f3711b44f138","resolution":{"observed_at":"2026-08-12T00:23:50.423415Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T00:23:50.427827Z","title":"Language models don’t always say what they think: Unfaithful explanations in chain-of-thought prompting.Advances in Neural Information Processing Systems, 36:74952–74965, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.427827Z"},"links":{"citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:673e1c76e670fae53eaed7dbe8aa0b43cb988456222a0522ec1da175990404e1","observation_id":"1b8f143f-e2be-4612-b50f-c6a9e03c2836","resolution":{"observed_at":"2026-08-12T00:23:50.427827Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T00:23:50.431998Z","title":"Let’s verify step by step","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.431998Z"},"links":{"citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:452cd13c211fc49e8f5453bdff4926d49dee9d7078cd0d5fb2df01d8a169513d","observation_id":"290a74bd-ce7a-4316-875e-ad22de85c5d3","resolution":{"observed_at":"2026-08-12T00:23:50.431998Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.22928","last_updated":"2025-07-24T10:25:46Z","snapshot_observed_at":"2026-08-09T08:01:51.012986Z","submitted_at":"2025-07-24T10:25:46Z","title":"How does Chain of Thought Think? Mechanistic Interpretability of Chain-of-Thought Reasoning with Sparse Autoencoding","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.22928","snapshot_observed_at":"2026-08-12T00:23:50.436202Z","title":"How does chain of thought think? mechanistic interpretability of chain-of-thought reasoning with sparse autoencoding.arXiv preprint arXiv:2507.22928, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.436202Z"},"links":{"cited_paper":"/paper/2507.22928","citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:9395a4b5be4f5a47ca1f3f01dd781cdab7195c9a50b81b4be5cc877f33ac2424","observation_id":"33f4f543-12ee-479c-82d3-1435710c04de","resolution":{"observed_at":"2026-08-12T00:23:50.436202Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T00:23:50.887666Z","title":"Finding sparse autoencoder representations of errors in cot prompting","venue":null,"work_id":"2c866910-040d-4336-95c2-28ff989da62f","year":2025},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.441305Z"},"links":{"citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:624a88280d29f654fa794cb957f6388617515e87f53ef1fd9576beaf1ac35676","observation_id":"25f08bb4-d134-458b-8514-0b0de5a37b2b","resolution":{"observed_at":"2026-08-12T00:23:50.891768Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.05217","last_updated":"2023-10-19T21:25:32Z","snapshot_observed_at":"2026-08-02T10:20:00.635719Z","submitted_at":"2023-01-12T18:56:49Z","title":"Progress measures for grokking via mechanistic interpretability","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.05217","snapshot_observed_at":"2026-08-12T00:23:50.445426Z","title":"Progress measures for grokking via mechanistic interpretability.arXiv preprint arXiv:2301.05217, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.445426Z"},"links":{"cited_paper":"/paper/2301.05217","citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:f150ab254a512292b455c1f0b178ecce2334e5a6c46a3ac0bff9c9fbf4c09788","observation_id":"2d4e165a-6dc8-4c43-90b0-a2c62a911c0e","resolution":{"observed_at":"2026-08-12T00:23:50.445426Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2210.13382","last_updated":"2024-06-26T14:27:49Z","snapshot_observed_at":"2026-08-09T03:25:38.596456Z","submitted_at":"2022-10-24T16:29:55Z","title":"Emergent World Representations: Exploring a Sequence Model Trained on a Synthetic Task","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.13382","snapshot_observed_at":"2026-08-12T00:23:50.449277Z","title":"Emergent world representations: Exploring a sequence model trained on a synthetic task.arXiv preprint arXiv:2210.13382, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.449277Z"},"links":{"cited_paper":"/paper/2210.13382","citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:eb57beeb6edf67700aec1760428cd7f94b46d9549fb35c59c15b9b0055d2c0f8","observation_id":"3007f258-dc91-4983-88cd-871ce36b1c59","resolution":{"observed_at":"2026-08-12T00:23:50.449277Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.01405","last_updated":"2025-03-03T06:14:14Z","snapshot_observed_at":"2026-07-06T16:26:38.284922Z","submitted_at":"2023-10-02T17:59:07Z","title":"Representation Engineering: A Top-Down Approach to AI Transparency","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.01405","snapshot_observed_at":"2026-08-12T00:23:50.453594Z","title":"Representation engineering: A top-down approach to ai transparency.arXiv preprint arXiv:2310.01405, 2023","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.453594Z"},"links":{"cited_paper":"/paper/2310.01405","citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:20ebb7a06c0f75a6e30d9fd13d90a91cd19d18e7af976a415d55551131a43e95","observation_id":"87bd8101-a35a-44b4-92c1-0f84e1adcae6","resolution":{"observed_at":"2026-08-12T00:23:50.453594Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.16014","last_updated":"2024-04-30T17:54:04Z","snapshot_observed_at":"2026-08-02T10:59:21.418230Z","submitted_at":"2024-04-24T17:47:22Z","title":"Improving Dictionary Learning with Gated Sparse Autoencoders","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.16014","snapshot_observed_at":"2026-08-12T00:23:50.457846Z","title":"Improving dictionary learning with gated sparse autoencoders.arXiv preprint arXiv:2404.16014, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.457846Z"},"links":{"cited_paper":"/paper/2404.16014","citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:2b7fe810647a886ba1c6da037ec04ac4f02799a27447f6a56bb6e9b39cbbd1e7","observation_id":"5a1bb2e0-1002-4714-95ed-e3b8717a381a","resolution":{"observed_at":"2026-08-12T00:23:50.457846Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.08600","last_updated":"2023-10-04T13:17:38Z","snapshot_observed_at":"2026-08-11T07:21:14.568175Z","submitted_at":"2023-09-15T17:56:55Z","title":"Sparse Autoencoders Find Highly Interpretable Features in Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.08600","snapshot_observed_at":"2026-08-12T00:23:50.462347Z","title":"Sparse autoencoders find highly interpretable features in language models.arXiv preprint arXiv:2309.08600, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.462347Z"},"links":{"cited_paper":"/paper/2309.08600","citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:3b9359d416ba09e7517ddb4b4dca21859dcf33ec880eda3aa2e684fe8d8fc0a9","observation_id":"8328cf95-7b15-438b-b4bb-a9189be04b13","resolution":{"observed_at":"2026-08-12T00:23:50.462347Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T00:23:50.466491Z","title":"Towards monosemanticity: Decomposing language models with dictionary learning.Transformer Circuits Thread, 2, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.466491Z"},"links":{"citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:9a3546ee00126a62e95f32e934818d4a15f7922e6bff97a7a2bd063d888f3ae4","observation_id":"dc9a124f-eb76-40b0-83f5-827bb157a498","resolution":{"observed_at":"2026-08-12T00:23:50.466491Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.15634","last_updated":"2025-07-12T09:42:16Z","snapshot_observed_at":"2026-08-10T11:46:31.963920Z","submitted_at":"2025-05-21T15:17:59Z","title":"Feature Extraction and Steering for Enhanced Chain-of-Thought Reasoning in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.15634","snapshot_observed_at":"2026-08-12T00:23:50.470682Z","title":"Feature extraction and steering for enhanced chain-of-thought reasoning in language models.arXiv preprint arXiv:2505.15634, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.470682Z"},"links":{"cited_paper":"/paper/2505.15634","citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:5a27a497408a3dcac4244ed734fa257e2b36f00c78b97ecdd90841ecaeb9ada9","observation_id":"115a9d7f-3f22-4362-9688-271f013aea15","resolution":{"observed_at":"2026-08-12T00:23:50.470682Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T00:23:50.475049Z","title":"Locating and editing factual associations in gpt","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.475049Z"},"links":{"citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:e8d78322d70695672e542bc2d98b8aa182b5ba6a145c0fbcadbd59610a9660fa","observation_id":"6b0e1e58-301d-4bf1-ae33-d7a88e444086","resolution":{"observed_at":"2026-08-12T00:23:50.475049Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.04093","last_updated":"2024-06-06T14:10:12Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-06T14:10:12Z","title":"Scaling and evaluating sparse autoencoders","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.04093","snapshot_observed_at":"2026-08-12T00:23:50.478491Z","title":"Scaling and evaluating sparse autoencoders.arXiv preprint arXiv:2406.04093, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.478491Z"},"links":{"cited_paper":"/paper/2406.04093","citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:f431a72579ebba2986c758d07929006e05f9159777b5820ac997e5679e38d634","observation_id":"27b113c0-62ab-4821-9d15-3119b15046de","resolution":{"observed_at":"2026-08-12T00:23:50.478491Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.05147","last_updated":"2024-08-19T07:51:05Z","snapshot_observed_at":"2026-08-11T23:41:02.350987Z","submitted_at":"2024-08-09T16:06:42Z","title":"Gemma Scope: Open Sparse Autoencoders Everywhere All At Once on Gemma 2","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.05147","snapshot_observed_at":"2026-08-12T00:23:50.482648Z","title":"Gemma scope: Open sparse autoencoders everywhere all at once on gemma 2.arXiv preprint arXiv:2408.05147, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.482648Z"},"links":{"cited_paper":"/paper/2408.05147","citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:bb79dbe9aaa62865274e22793548574561621eb9dbc1fa8bfc57111f390e4b65","observation_id":"ada82d09-0663-421c-a985-c39ccd67f7e3","resolution":{"observed_at":"2026-08-12T00:23:50.482648Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T00:23:50.860112Z","title":"Sparse autoencoder features for classifications and transferability","venue":null,"work_id":"a672d659-45cb-405c-be48-320ab5d0e5d0","year":2025},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.486695Z"},"links":{"citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:6c271022e9b5ac76e4eb201c9256bfdfbfef43d09aa6bab142bb3c38b6354f65","observation_id":"62dbf1e2-29f5-4e4e-8a34-db0cdfac515b","resolution":{"observed_at":"2026-08-12T00:23:50.864087Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T00:23:50.490578Z","title":"Saes are good for steering–if you select the right features","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.490578Z"},"links":{"citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:5bbc5c86a86e9362b6cacc072a471a341c86234e7959e0b7a0bd973f39983aef","observation_id":"27f7208c-d0eb-4fdf-806f-3f14ac1c762e","resolution":{"observed_at":"2026-08-12T00:23:50.490578Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T00:23:50.845729Z","title":"Lingualens: Towards interpreting linguistic mechanisms of large language models via sparse auto-encoder","venue":null,"work_id":"c257f657-1a2f-4ff5-bab0-833ba519276f","year":2025},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.494278Z"},"links":{"citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:3ef055dc21053e38d460ebf90ea3c45c80ebf86730e631a2f19b993b898c641a","observation_id":"ece60cbd-713e-4dd9-a3dd-e999016a653b","resolution":{"observed_at":"2026-08-12T00:23:50.851200Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.00041","last_updated":"2025-08-27T13:45:53Z","snapshot_observed_at":"2026-08-09T21:03:26.575507Z","submitted_at":"2025-05-28T02:50:17Z","title":"Decoding Dense Embeddings: Sparse Autoencoders for Interpreting and Discretizing Dense Retrieval","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.00041","snapshot_observed_at":"2026-08-12T00:23:50.498053Z","title":"Decoding dense embeddings: Sparse autoencoders for interpreting and discretizing dense retrieval.arXiv preprint arXiv:2506.00041, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.498053Z"},"links":{"cited_paper":"/paper/2506.00041","citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:7d6512612c448f1b943257aa5f199f4319b93caeffa4a97d9c00b679b949cb10","observation_id":"14b9792e-6e85-4ffc-87fe-a920bfaaa251","resolution":{"observed_at":"2026-08-12T00:23:50.498053Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1312.5663","last_updated":"2014-03-22T17:12:07Z","snapshot_observed_at":"2026-08-01T16:20:01.675800Z","submitted_at":"2013-12-19T17:46:46Z","title":"k-Sparse Autoencoders","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1312.5663","snapshot_observed_at":"2026-08-12T00:23:50.501902Z","title":"K-sparse autoencoders.arXiv preprint arXiv:1312.5663, 2013","venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.501902Z"},"links":{"cited_paper":"/paper/1312.5663","citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:93e2aad19e52e4c960b73e775c1dc7a40652dbd64ccf790c2949ccc4643f9448","observation_id":"9a9da866-17ca-4060-9735-6573655149a4","resolution":{"observed_at":"2026-08-12T00:23:50.501902Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.11456","last_updated":"2025-05-22T19:12:14Z","snapshot_observed_at":"2026-08-10T19:38:46.551479Z","submitted_at":"2025-04-15T17:59:51Z","title":"DeepMath-103K: A Large-Scale, Challenging, Decontaminated, and Verifiable Mathematical Dataset for Advancing Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.11456","snapshot_observed_at":"2026-08-12T00:23:50.505227Z","title":"Deepmath-103k: A large-scale, challenging, decontaminated, and verifiable mathematical dataset for advancing reasoning.arXiv preprint arXiv:2504.11456, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.505227Z"},"links":{"cited_paper":"/paper/2504.11456","citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:3fdde5d23889334f3bd99ab262af78328c3f89e4702b5a836d9f19536246cb1f","observation_id":"263e575a-3be9-4270-84bf-193f546a2271","resolution":{"observed_at":"2026-08-12T00:23:50.505227Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.09858","last_updated":"2025-04-14T04:08:16Z","snapshot_observed_at":"2026-08-08T01:06:59.135762Z","submitted_at":"2025-04-14T04:08:16Z","title":"Reasoning Models Can Be Effective Without Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.09858","snapshot_observed_at":"2026-08-12T00:23:50.508557Z","title":"Reasoning models can be effective without thinking.arXiv preprint arXiv:2504.09858, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.508557Z"},"links":{"cited_paper":"/paper/2504.09858","citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:2f37beac665dc9a6b323a22f66321f5d9c4a79c46650cd660e8140b2e6e8b2e0","observation_id":"e503d8e4-528e-487f-9c43-5f1f235040d1","resolution":{"observed_at":"2026-08-12T00:23:50.508557Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T00:23:50.830957Z","title":"Interpreting and steering llm representations with mutual information-based explanations on sparse autoencoders","venue":null,"work_id":"f5d9ad68-3c30-4ae8-b070-12ccf7ab8dd2","year":null},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.512555Z"},"links":{"citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:2bedb24b94be9a0d9245731c48f5873c9874ad17e68bbfbe021578511e2e404d","observation_id":"53ef2980-cb57-4d52-83ca-8bf913379abd","resolution":{"observed_at":"2026-08-12T00:23:50.836653Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1412.6980","last_updated":"2017-01-30T01:27:54Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2014-12-22T13:54:29Z","title":"Adam: A Method for Stochastic Optimization","version":9},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1412.6980","snapshot_observed_at":"2026-08-12T00:23:50.516505Z","title":"A method for stochastic optimization.arXiv preprint arXiv:1412.6980, 1412(6), 2014","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.516505Z"},"links":{"cited_paper":"/paper/1412.6980","citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:5d8296c20028170634f3a19ded298277c22c6b775a86fa762d0d0fd7e67c6fad","observation_id":"b85240b3-6c73-4274-b747-c8c9a8e3cf3d","resolution":{"observed_at":"2026-08-12T00:23:50.516505Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T00:23:50.815528Z","title":"Scaling monosemanticity: Extracting interpretable features from claude 3 sonnet","venue":null,"work_id":"7bab6e44-649e-4b92-9136-ca12e243964c","year":2024},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.520191Z"},"links":{"citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:7cd353ee21b1378c36ccfb02ff58737bbfa300eb2e8a371abf10e53834b3a8ca","observation_id":"eb2de8b2-28f3-4360-a468-ce05979d74fb","resolution":{"observed_at":"2026-08-12T00:23:50.821013Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T00:23:50.800935Z","title":"wait\", \"hmm","venue":null,"work_id":"2b491258-1689-4425-9ab6-f825cadb9b20","year":2016},"citing_paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-12T00:23:50.523918Z"},"links":{"citing_paper":"/paper/2608.08168"},"observation_digest":"sha256:1ad0b0a80ae4dbc6d2a1670ce147a9901e2dc3b388e5dbb7584c4310c57f5397","observation_id":"a8016092-872c-4628-bf30-8c059e3c4502","resolution":{"observed_at":"2026-08-12T00:23:50.806064Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2608.08168","last_updated":"2026-08-08T14:54:13Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-12T22:17:15.464972Z","submitted_at":"2026-08-08T14:54:13Z","title":"Thinking vs. NoThinking: Towards Interpreting Reasoning Mechanisms of Large Language Models via Sparse Autoencoders"},"reference_resolution":{"displayed":27,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":21,"verified_exact":0,"verified_fuzzy":6},"total_outbound_references":27},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"thesis":"As of 12 August 2026, this Paper Citation Record lists 27 of 27 outbound references and 0 inbound Pith citation observations for arXiv:2608.08168."}