{"as_of":"2026-08-07T18:07:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:3728592e78e208ed81b36f95e9b6d619f4ebe9bf384b4b95132bde5f44e6469a","coverage":[{"denominator":60,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":60,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T01:03:43.999176Z","state":"measured"},{"denominator":61,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":61,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-03T00:55:32.686738Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.12217","snapshot_observed_at":"2026-08-03T00:55:32.686738Z","title":"arXiv preprint arXiv:2506.12217 , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28636","last_updated":"2026-05-19T13:56:13Z","snapshot_observed_at":"2026-08-06T00:38:04.006327Z","submitted_at":"2026-05-19T13:56:13Z","title":"Chain-of-Models: Cross-Model Auditing for Bias-Robust LLM Judges","version":1},"reference_index":201,"source":"arxiv_source","source_observed_at":"2026-08-03T00:55:32.686738Z"},"links":{"cited_paper":"/paper/2506.12217","citing_paper":"/paper/2607.28636"},"observation_digest":"sha256:1b6fcd143e8c2be24b427773f3912851e39c9d695e1e7ac5b4d17943349b48eb","observation_id":"09e6d2ee-3503-4e81-8bf6-3281c83c64a6","resolution":{"observed_at":"2026-08-03T00:55:32.686738Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2506.12217/citation-record","integrity":"/paper/2506.12217/integrity","json":"/paper/2506.12217/citation-record.json","paper":"/paper/2506.12217"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2501.09686","last_updated":"2025-01-23T08:44:44Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-16T17:37:58Z","title":"Towards Large Reasoning Models: A Survey of Reinforced Reasoning with Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.09686","snapshot_observed_at":"2026-08-07T01:03:38.561365Z","title":"Towards large reasoning models: A survey of reinforced reasoning with large language models,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:38.561365Z"},"links":{"cited_paper":"/paper/2501.09686","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:7c733425515b8f8f62fa92def8adc10f4d6aa19cf898544c867063f18e05ba4c","observation_id":"1629fedd-c51a-443c-8c86-a9e44716013d","resolution":{"observed_at":"2026-08-07T01:03:38.561365Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:50.006721Z","title":"Self-reasoning language models: Unfold hidden reasoning chains with few reasoning catalyst,","venue":null,"work_id":"27856238-d4e3-4a6d-87a4-556895271826","year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:38.644078Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:b8bdec333e1a0fe76587de4921a17863114d4fa8c3761c04bbc6c3d7aea0987e","observation_id":"c3f47b45-1cf7-4040-b36c-88f44fb55353","resolution":{"observed_at":"2026-08-07T01:03:50.091457Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:38.746258Z","title":"Reinforcement learning with verifiable rewards: Grpo’s effective loss, dynamics, and success amplification,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:38.746258Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:264d01d3bd6c190e9147ad906880b35eadca5b450cd25b898ca15c29ae38819e","observation_id":"d39c0232-4edf-4f1f-93f0-77c3ec0e014d","resolution":{"observed_at":"2026-08-07T01:03:38.746258Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.05379","last_updated":"2025-03-10T07:11:14Z","snapshot_observed_at":"2026-08-07T17:22:15.100571Z","submitted_at":"2025-03-07T12:46:42Z","title":"R1-Omni: Explainable Omni-Multimodal Emotion Recognition with Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.05379","snapshot_observed_at":"2026-08-07T01:03:38.824120Z","title":"R1-omni: Explainable omni-multimodal emotion recognition with rein- forcement learning,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:38.824120Z"},"links":{"cited_paper":"/paper/2503.05379","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:0b588be8148aaaa20dbcb857aa3258f4688f0887953cf77d530d6b41bb2149a0","observation_id":"72ccaf44-fc65-488a-a1cd-10b9b7ca1794","resolution":{"observed_at":"2026-08-07T01:03:38.824120Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:38.946446Z","title":"Reasoning beyond limits: Advances and open problems for llms,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:38.946446Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:9364288a604396fa3e44b265bd9da010c101b8a01f37e94e6e967c67e58cb1c8","observation_id":"95f02528-e180-46f1-b36e-a332307da8d4","resolution":{"observed_at":"2026-08-07T01:03:38.946446Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.23829","last_updated":"2025-04-01T14:48:02Z","snapshot_observed_at":"2026-08-07T16:25:50.798894Z","submitted_at":"2025-03-31T08:22:49Z","title":"Crossing the Reward Bridge: Expanding RL with Verifiable Rewards Across Diverse Domains","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.23829","snapshot_observed_at":"2026-08-07T01:03:39.063087Z","title":"Expanding rl with verifiable rewards across diverse domains,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.063087Z"},"links":{"cited_paper":"/paper/2503.23829","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:f58c2d1c69245731e7904d3845dc09d42aff34f4bfe911b051a900713dbef5a2","observation_id":"45d137f3-ca87-4e2f-b01a-e084297c6e34","resolution":{"observed_at":"2026-08-07T01:03:39.063087Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-07T01:03:39.150691Z","title":"Deepseek-r1: Incentivizing reasoning capability in llms via reinforcement learning,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.150691Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:3055f84761c6be08a8795ea1f9b75f5b5485b0d85aaa906a67ae95b468ce981e","observation_id":"6ce3f4ec-e201-4c21-8cfa-de0702166cd9","resolution":{"observed_at":"2026-08-07T01:03:39.150691Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.20783","last_updated":"2025-10-06T09:30:03Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-03-26T17:59:14Z","title":"Understanding R1-Zero-Like Training: A Critical Perspective","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.20783","snapshot_observed_at":"2026-08-07T01:03:39.248657Z","title":"Understanding r1-zero-like training: A critical perspective,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.248657Z"},"links":{"cited_paper":"/paper/2503.20783","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:2701677162715919e4593349812f7be6a322a757e9f97f665fee65cc0111c1f2","observation_id":"00266b23-2c47-4af4-a0f7-ac7c0164fc40","resolution":{"observed_at":"2026-08-07T01:03:39.248657Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.18892","last_updated":"2025-08-06T08:42:32Z","snapshot_observed_at":"2026-07-06T20:57:57.039376Z","submitted_at":"2025-03-24T17:06:10Z","title":"SimpleRL-Zoo: Investigating and Taming Zero Reinforcement Learning for Open Base Models in the Wild","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.18892","snapshot_observed_at":"2026-08-07T01:03:39.348782Z","title":"Simplerl-zoo: Investigating and taming zero reinforcement learning for open base models in the wild,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.348782Z"},"links":{"cited_paper":"/paper/2503.18892","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:2aca6ec50b65e2a0a4a26100505ab6dfcc960b11dc875d4d4efdcf118e8d06f7","observation_id":"fec69ab1-664a-4a23-b22f-5e9923cbb2d1","resolution":{"observed_at":"2026-08-07T01:03:39.348782Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16084","last_updated":"2025-06-30T15:59:26Z","snapshot_observed_at":"2026-07-06T21:13:13.686703Z","submitted_at":"2025-04-22T17:59:56Z","title":"TTRL: Test-Time Reinforcement Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.16084","snapshot_observed_at":"2026-08-07T01:03:39.439450Z","title":"Ttrl: Test-time reinforcement learning,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.439450Z"},"links":{"cited_paper":"/paper/2504.16084","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:076b7332f1058c580e9677f71fcf96b023e11de69374624ec5ccccec3adca523","observation_id":"39da37bb-30e3-44ac-a74b-d2899675fa2e","resolution":{"observed_at":"2026-08-07T01:03:39.439450Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.13837","last_updated":"2025-11-24T06:11:04Z","snapshot_observed_at":"2026-07-06T21:11:34.701779Z","submitted_at":"2025-04-18T17:59:56Z","title":"Does Reinforcement Learning Really Incentivize Reasoning Capacity in LLMs Beyond the Base Model?","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.13837","snapshot_observed_at":"2026-08-07T01:03:39.555965Z","title":"Does reinforcement learning really incentivize reasoning capacity in llms beyond the base model?,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.555965Z"},"links":{"cited_paper":"/paper/2504.13837","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:50492fca5874d3d758c1e95aeae5c6178b158c48bce12f315f3f6df5510d96f7","observation_id":"6e48c6ec-e596-4e7a-a061-ee7584b866ec","resolution":{"observed_at":"2026-08-07T01:03:39.555965Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10400","last_updated":"2025-02-16T09:33:15Z","snapshot_observed_at":"2026-08-07T04:49:07.615320Z","submitted_at":"2024-06-14T20:07:11Z","title":"Self-Reflection Makes Large Language Models Safer, Less Biased, and Ideologically Neutral","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10400","snapshot_observed_at":"2026-08-07T01:03:39.646395Z","title":"Self-reflection outcome is sensitive to prompt construction,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.646395Z"},"links":{"cited_paper":"/paper/2406.10400","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:43bfc31dbefb7061623481920975fd7298d3fd9d85277e65fc5ce17851772ebf","observation_id":"5012415e-37d7-479a-b840-2af01d93001e","resolution":{"observed_at":"2026-08-07T01:03:39.646395Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:39.763786Z","title":"Dynamic early exit in reasoning models,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.763786Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:0931424920df86b645fb434b69486da0bac81f136b12022e5b9f65c94f7c258f","observation_id":"3a583017-4b0b-4ce9-a51f-739836db59d9","resolution":{"observed_at":"2026-08-07T01:03:39.763786Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.06682","last_updated":"2024-10-16T23:19:46Z","snapshot_observed_at":"2026-07-06T18:12:47.499337Z","submitted_at":"2024-05-05T18:56:46Z","title":"Self-Reflection in LLM Agents: Effects on Problem-Solving Performance","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.06682","snapshot_observed_at":"2026-08-07T01:03:39.875381Z","title":"Self-reflection in llm agents: Effects on problem-solving performance,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.875381Z"},"links":{"cited_paper":"/paper/2405.06682","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:7beeb32a87650d8c829002ee6674e3288ea325fb00cf60ea948db17e10bb75ef","observation_id":"0ff7c1eb-0746-4af3-a970-991237b9ad2b","resolution":{"observed_at":"2026-08-07T01:03:39.875381Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.16419","last_updated":"2025-08-21T19:14:40Z","snapshot_observed_at":"2026-08-07T04:27:23.738927Z","submitted_at":"2025-03-20T17:59:38Z","title":"Stop Overthinking: A Survey on Efficient Reasoning for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.16419","snapshot_observed_at":"2026-08-07T01:03:39.949009Z","title":"Stop overthinking: Asurveyonefficientreasoningforlargelanguagemodels,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.949009Z"},"links":{"cited_paper":"/paper/2503.16419","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:48f0f976e0eaf0d551ebf98e26059d9f50fc54726670e9838a4a902a98c270f5","observation_id":"e75301d6-8fc0-4c2b-8fd8-258fd41755a5","resolution":{"observed_at":"2026-08-07T01:03:39.949009Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:49.808216Z","title":"Steering llama 2 via contrastive activation addition,","venue":null,"work_id":"0745a100-d217-4ccd-87e1-d5d74358f3b1","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.075252Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:75451592405c9162e85d49f157ff70f3d663153783b73c13a1de55dce7a23279","observation_id":"1c3c4d59-7537-49e7-87a4-48675c718d1f","resolution":{"observed_at":"2026-08-07T01:03:49.915796Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1801.10198","last_updated":"2018-01-30T20:07:01Z","snapshot_observed_at":"2026-07-06T06:21:01.297755Z","submitted_at":"2018-01-30T20:07:01Z","title":"Generating Wikipedia by Summarizing Long Sequences","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1801.10198","snapshot_observed_at":"2026-08-07T01:03:40.150110Z","title":"Generating wikipedia by summarizing long sequences,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.150110Z"},"links":{"cited_paper":"/paper/1801.10198","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:dd1005fa56562fde80c71655cef173208ae0cbd5d0d3a59156ac1be9d62a6cc2","observation_id":"dfcd3ad3-4410-464d-bd53-7cafe43242bd","resolution":{"observed_at":"2026-08-07T01:03:40.150110Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:49.607196Z","title":"Beyond accuracy: Evaluating the reasoning behavior of large language models - a survey,","venue":null,"work_id":"cbbc5d79-a6f5-4984-90a3-4b1982b41863","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.258094Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:2581c37467893a7555d3bf6c5696b467185a6f5d9a610c0305c20371989a43f0","observation_id":"d7c5a580-fc9f-4a8e-afd2-9006c9ee36f6","resolution":{"observed_at":"2026-08-07T01:03:49.713626Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:49.409413Z","title":"Oat: A research-friendly framework for llm online alignment,","venue":null,"work_id":"aed4a729-51c1-4061-977d-474fb660b103","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.356585Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:b04a321ad0d3223c400e1b00c85b49677440ee4f7c84e52fc5c0651825f909b3","observation_id":"6edd9d22-89a5-4dc0-9c6f-af10c5a203a7","resolution":{"observed_at":"2026-08-07T01:03:49.490397Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:49.221893Z","title":"Self- consistency improves chain of thought reasoning in language models,","venue":null,"work_id":"54b239e4-7c40-4e94-9625-effc4d5b436e","year":2023},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.427956Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:d29eb1bd602548ace585c820ad685d3bc16b046f98187222fb0a1a6046e76f20","observation_id":"aa276d97-fc20-4b0b-9055-911fb8bdbae9","resolution":{"observed_at":"2026-08-07T01:03:49.324127Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:49.065199Z","title":"X-reasoner: Towards generalizable reasoning across modalities and domains,","venue":null,"work_id":"d7fd41cd-9359-4c55-a606-0653fcea7162","year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.536878Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:c8b85361ce7f779b8d93d08ff328d84c800350d85b347f780d242dba55fbce0d","observation_id":"f3ca3234-ffce-4943-a683-0ebd222aaa98","resolution":{"observed_at":"2026-08-07T01:03:49.145880Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.10400","last_updated":"2025-02-24T08:57:10Z","snapshot_observed_at":"2026-07-06T20:06:43.917795Z","submitted_at":"2024-12-05T16:10:42Z","title":"Reinforcement Learning Enhanced LLMs: A Survey","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.10400","snapshot_observed_at":"2026-08-07T01:03:40.631750Z","title":"Reinforcement learning enhanced llms: A survey,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.631750Z"},"links":{"cited_paper":"/paper/2412.10400","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:43f503812ebdf377e7c2158ed96ed3230ff6fb2b76e66da91c642860a9d0ff18","observation_id":"271c3ef1-143a-4feb-a781-27cf896f307c","resolution":{"observed_at":"2026-08-07T01:03:40.631750Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.03335","last_updated":"2025-10-16T08:23:36Z","snapshot_observed_at":"2026-07-06T21:19:34.329442Z","submitted_at":"2025-05-06T09:08:00Z","title":"Absolute Zero: Reinforced Self-play Reasoning with Zero Data","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.03335","snapshot_observed_at":"2026-08-07T01:03:40.693792Z","title":"Absolute zero: Reinforced self-play reasoning with zero data,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.693792Z"},"links":{"cited_paper":"/paper/2505.03335","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:a5011cf83f439e8eb8b38cff745c7c3560c4688b89c9f4e8a0b518431377b065","observation_id":"a66ec951-743b-4545-99c4-8e4e8a841cbf","resolution":{"observed_at":"2026-08-07T01:03:40.693792Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.09129","last_updated":"2024-04-14T02:47:32Z","snapshot_observed_at":"2026-08-05T16:24:49.303170Z","submitted_at":"2024-04-14T02:47:32Z","title":"When Hindsight is Not 20/20: Testing Limits on Reflective Thinking in Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.09129","snapshot_observed_at":"2026-08-07T01:03:40.835799Z","title":"When hindsight is not 20/20: Testing limits on reflective thinking in large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.835799Z"},"links":{"cited_paper":"/paper/2404.09129","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:6cef46b0dbfda58342937c9b651041be2034bb646fb540d3cda85c43c257e3f7","observation_id":"f9f122bc-638e-4812-94d2-a741f75f391c","resolution":{"observed_at":"2026-08-07T01:03:40.835799Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:48.892164Z","title":"Demystifyinglongchain-of-thoughtreasoninginLLMs,","venue":null,"work_id":"593e40ec-1d05-43bb-b7cf-55f9fd549fd6","year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.918114Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:544b5d8ee10fac769a64fd82fc19a01a5ba44c4d72efce847ec6b47cacc1dff4","observation_id":"d7ed41c2-594d-4bc7-8560-4063ddfa933d","resolution":{"observed_at":"2026-08-07T01:03:48.981885Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.16720","last_updated":"2026-04-30T02:46:40Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-21T18:04:31Z","title":"OpenAI o1 System Card","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.16720","snapshot_observed_at":"2026-08-07T01:03:41.000635Z","title":"Openai o1 system card,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.000635Z"},"links":{"cited_paper":"/paper/2412.16720","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:f75d6d281d0bba02203a3afd4d094101fe238283b98065d4d964bde83a6428a1","observation_id":"9646293b-ee21-4f86-9fc0-d89a56d2de0f","resolution":{"observed_at":"2026-08-07T01:03:41.000635Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.00656","last_updated":"2025-10-08T07:50:45Z","snapshot_observed_at":"2026-07-06T20:15:22.511263Z","submitted_at":"2024-12-31T21:55:10Z","title":"2 OLMo 2 Furious","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.00656","snapshot_observed_at":"2026-08-07T01:03:41.101622Z","title":"2 olmo 2 furious,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.101622Z"},"links":{"cited_paper":"/paper/2501.00656","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:d30a40935d98e2dae03bc516bf3721d64080daadf08a09ef0ee118c36a85d178","observation_id":"469303ad-d1e6-4663-a969-09b9a3aad043","resolution":{"observed_at":"2026-08-07T01:03:41.101622Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:48.708810Z","title":"Qwen3technical report,","venue":null,"work_id":"3e1375b7-3354-43fb-ab92-0eaab029689c","year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.187541Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:b1fd107a5529d7cc323ee321871ddb852a5f6c1bef3bf44ecaa9885e9727c420","observation_id":"a6bebd1d-2878-441b-b47a-317a95e26e40","resolution":{"observed_at":"2026-08-07T01:03:48.803809Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:48.524846Z","title":"Measuring mathematical problem solving with the math dataset,","venue":null,"work_id":"c0c06683-3b2b-4300-ad76-be6b82fc90fc","year":2021},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.281529Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:6c4bdce81b81e22d6cef62a9651e7d6a969238f6ec4754e7684184c56615ed54","observation_id":"7c38a833-c42f-4bdf-b9e3-6d709d13d656","resolution":{"observed_at":"2026-08-07T01:03:48.604954Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.15115","last_updated":"2025-01-03T02:18:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-19T17:56:09Z","title":"Qwen2.5 Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.15115","snapshot_observed_at":"2026-08-07T01:03:41.367515Z","title":"Qwen2.5 technical report,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.367515Z"},"links":{"cited_paper":"/paper/2412.15115","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:98d297500dd2d8165d2f744460e6e876c5c6e360e82e94f3f73afc7bd237041b","observation_id":"c3b0ce6b-1dcd-4094-a847-9ac37c243daf","resolution":{"observed_at":"2026-08-07T01:03:41.367515Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:48.305684Z","title":"Umap: Uniform manifold approximation and projection,","venue":null,"work_id":"4c53cb44-2ed2-4de6-bb60-4224652c4f14","year":2018},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.455519Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:e2a6839323da335539d872e90996085cb6ce339397607312ac0b51dbb9d24a58","observation_id":"5f2a934e-8188-4b61-b082-cfc0a740d4f2","resolution":{"observed_at":"2026-08-07T01:03:48.402203Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:48.105369Z","title":"Refusal in language models is mediated by a single direction,","venue":null,"work_id":"22c5c963-05d1-4048-be22-af5458561e6b","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.581041Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:0cfefd34e577a8281f81ec282f3555fffb90df3daa565e57b7b964fd4270224d","observation_id":"c48b7e33-6880-4a41-9bcc-4e2415ae0261","resolution":{"observed_at":"2026-08-07T01:03:48.223985Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.17148","last_updated":"2025-03-03T21:15:30Z","snapshot_observed_at":"2026-07-06T20:27:32.742145Z","submitted_at":"2025-01-28T18:51:24Z","title":"AxBench: Steering LLMs? Even Simple Baselines Outperform Sparse Autoencoders","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.17148","snapshot_observed_at":"2026-08-07T01:03:41.633267Z","title":"Axbench: Steering llms? even simple baselines outperform sparse autoencoders,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.633267Z"},"links":{"cited_paper":"/paper/2501.17148","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:f326046794f3f6b7ee1b86eaab78ae0ef5ef630faae68a9dc1c65f93d8df5ab4","observation_id":"aa3cfd68-4461-42d7-9c7b-067819b1df5e","resolution":{"observed_at":"2026-08-07T01:03:41.633267Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:47.927657Z","title":"GPQA: A graduate-level google-proof q&a benchmark,","venue":null,"work_id":"b4a52e1e-fe79-414e-b1fd-2d6c6bcb1404","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.688003Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:c8ddbf7b4a43ceda47c73b0f8b964a09ac2135894a476c34838206065f188539","observation_id":"574728a9-8514-40e5-808d-9ebd8bffb441","resolution":{"observed_at":"2026-08-07T01:03:48.007208Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-07-06T18:55:11.576666Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-07T01:03:41.772289Z","title":"The llama 3 herd of models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.772289Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:45b12cc1aa5c5663d311f4f7ae59584525262a63914bc132c9c4057b10a92735","observation_id":"663e3187-2685-436c-8ec6-a625da5ab9ca","resolution":{"observed_at":"2026-08-07T01:03:41.772289Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:47.775398Z","title":"s1: Simple test-time scaling,","venue":null,"work_id":"4c7242c8-6b95-4250-a836-309efa758e88","year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.883084Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:5936c63da2d651be05c599447352406ff5cc75aee6e519376957bab2fb9055cb","observation_id":"176a923f-3bc6-4556-bb92-d9d3e3b7a73e","resolution":{"observed_at":"2026-08-07T01:03:47.843540Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:47.609327Z","title":"Discovering latent knowledge in language models without supervision,","venue":null,"work_id":"d9ecbae2-a49a-4b2c-92b8-cab0d4ac86d1","year":2022},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.999426Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:8f4e0b245071232b7cab317e7add80f408a213420db30e6b5892cbca0351dfcf","observation_id":"455bfbd3-f1f3-4569-9fb0-13256cdc8c65","resolution":{"observed_at":"2026-08-07T01:03:47.687530Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.15043","last_updated":"2023-12-20T20:48:57Z","snapshot_observed_at":"2026-07-06T15:59:23.019044Z","submitted_at":"2023-07-27T17:49:12Z","title":"Universal and Transferable Adversarial Attacks on Aligned Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.15043","snapshot_observed_at":"2026-08-07T01:03:42.097166Z","title":"Universal and transferable adversarial attacks on aligned language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.097166Z"},"links":{"cited_paper":"/paper/2307.15043","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:16c13b67726339afb2188babd409723d1bb87647f295a7af956411a408b71078","observation_id":"e080e39d-6fdc-46bd-b468-be482af58d7d","resolution":{"observed_at":"2026-08-07T01:03:42.097166Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14433","last_updated":"2024-02-22T10:25:14Z","snapshot_observed_at":"2026-08-01T05:10:27.518306Z","submitted_at":"2024-02-22T10:25:14Z","title":"A Language Model's Guide Through Latent Space","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.14433","snapshot_observed_at":"2026-08-07T01:03:42.165517Z","title":"A language model’s guide through latent space,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.165517Z"},"links":{"cited_paper":"/paper/2402.14433","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:bb4ab8b781cee69d0a11606b657a9e9b458577699ac1efc191978d0c652c6f81","observation_id":"066fae06-1e46-4c3e-943f-2ef5c7231f6c","resolution":{"observed_at":"2026-08-07T01:03:42.165517Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.03813","last_updated":"2023-12-06T18:27:07Z","snapshot_observed_at":"2026-07-06T16:57:59.000365Z","submitted_at":"2023-12-06T18:27:07Z","title":"Improving Activation Steering in Language Models with Mean-Centring","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.03813","snapshot_observed_at":"2026-08-07T01:03:42.231617Z","title":"Improvingactivationsteeringinlanguagemodels with mean-centring,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.231617Z"},"links":{"cited_paper":"/paper/2312.03813","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:d426b2450b0da29dc3f68447632107c63a460ddfb6623d871bdbd92519d5a32b","observation_id":"ccda8ba5-f1c8-4351-b338-06c32a8b5839","resolution":{"observed_at":"2026-08-07T01:03:42.231617Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:47.443594Z","title":"Finding alignments between interpretable causal variables and distributed neural representations,","venue":null,"work_id":"704a07e8-8912-4fd7-a283-7781f17fdfa2","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.293333Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:2a610bd651f3cd5be3a68b8326c71cc3d03cde5522f7561a0cc1b2682c67aa48","observation_id":"4b8fa7cc-b175-46cf-a4bc-2e07327b259b","resolution":{"observed_at":"2026-08-07T01:03:47.540585Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:47.231503Z","title":"Generative agents: Interactivesimulacraofhumanbehavior,","venue":null,"work_id":"58604775-c758-4c0d-b3a3-8065e31b59a3","year":2023},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.349684Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:9958e81e39b3a9a54f6c1f264c3fa4be37da00e41ed08fb387e64e053d3e005f","observation_id":"75b62a26-f9ce-45a9-bf6d-25a79750aa6c","resolution":{"observed_at":"2026-08-07T01:03:47.341340Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:47.070162Z","title":"Man is to computer programmer as woman is to homemaker? debiasing word embeddings,","venue":null,"work_id":"951afe25-0bda-4db5-98a1-fa91255a26b8","year":2016},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.442869Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:417442d792df33890e95dbcbf216d23bd531dbebb9ace23a853207147145a662","observation_id":"2e2ee699-326c-43f4-9ce6-e929d40d1b2a","resolution":{"observed_at":"2026-08-07T01:03:47.140600Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:46.909600Z","title":"Sparseautoencodersfindhighlyinter- pretable features in language models,","venue":null,"work_id":"0cf0e31a-b14b-4305-b6eb-9b42fe8f49e7","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.512310Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:0d896e94701500788fb7adef6985fab8ef8e8a233ea0453b6954ba2af4ae6a7b","observation_id":"5646bbed-0a47-48ee-92c0-a8de0ca8edb5","resolution":{"observed_at":"2026-08-07T01:03:46.980720Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:46.760950Z","title":"Gold doesn‘t always glitter: Spectral removal of linear and nonlinear guardedattributeinformation,","venue":null,"work_id":"2e587e80-176b-42ad-9a49-df3b7049bd34","year":2023},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.632456Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:d632e057c4cf4b9a3d346e28558a72304741fc7005a922b0fc9d6e6b59f6be93","observation_id":"f7461889-ce63-41fe-89a9-fbd38f2db7d1","resolution":{"observed_at":"2026-08-07T01:03:46.832967Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:46.595599Z","title":"LEACE: Perfect linear concept erasure in closed form,","venue":null,"work_id":"4df176e9-c6c9-4cca-a5dd-d001179ab7c2","year":2023},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.748817Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:d46c250070d421e38eaa68f24cb0bc8713ab7f0ae21f4588772a8d1f769a5a9f","observation_id":"ee86d914-0601-459b-b18b-37f5071d3aff","resolution":{"observed_at":"2026-08-07T01:03:46.659998Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:46.416338Z","title":"Monitoring latent world states in language models with proposi- tional probes,","venue":null,"work_id":"f193c926-0c04-4ff0-95f2-031e38902e3a","year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.820828Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:e6cb139541cc9e6939f1bc7dcb63fe902e929cc8d17268b5cd7e6f46fbeacede","observation_id":"4876e0ea-6e93-4259-9ab9-3a14a1faeb66","resolution":{"observed_at":"2026-08-07T01:03:46.506000Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:46.276366Z","title":"Let’s verify step by step,","venue":null,"work_id":"6a2867a7-344a-4b77-aaf8-667407e0867b","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.869100Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:c686f5a6c4de876eda182314c558c05bbd6e849b70466b4ae8b5fccfe44cd9d7","observation_id":"14c8a776-3a4d-498f-8938-059650e3aa86","resolution":{"observed_at":"2026-08-07T01:03:46.336643Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:46.141934Z","title":"Self-refine: Iterative refinement with self-feedback,","venue":null,"work_id":"ff0f3286-ad0f-4c29-bc52-e2d444f3d0e8","year":2023},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.995770Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:5dd2a6650b45815ac7159d11e67f7c6039eb1632e52df2ec0f374b2bd2995750","observation_id":"45964bdf-c1b5-4127-be86-22247b70986f","resolution":{"observed_at":"2026-08-07T01:03:46.194754Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:45.954249Z","title":"Fine-tuning with divergent chains of thought boosts reasoning through self-correction in language models,","venue":null,"work_id":"a6694914-fc18-421e-9cdc-66de285902fc","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.095901Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:0998a0625d1967754b76984d3bd9f9786538a43017d8374973035c623abc45ae","observation_id":"4f77234f-d40c-4b67-ac9a-698fee7368e6","resolution":{"observed_at":"2026-08-07T01:03:46.041393Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:45.807342Z","title":"STar: Bootstrapping reasoning with reasoning,","venue":null,"work_id":"f2da2629-3a02-45ed-b049-a19ebdf94b3f","year":2022},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.195392Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:7ceeb9b0fba6fee841d0839c6780acfcb971d71ba2d31bd3a916e9b8ea90824d","observation_id":"b0fac060-c3ff-45bf-be9d-73d9cfe73d47","resolution":{"observed_at":"2026-08-07T01:03:45.858720Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:45.594857Z","title":"Let’s verify step by step,","venue":null,"work_id":"6e0ab76d-2a2b-465f-ba08-d41449af2328","year":2023},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.254472Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:0470e467f4161e23cc8d94518fc95a3e799c9ea82db20c5041ac362aaee0daed","observation_id":"2912de11-1a0d-4b74-84d1-31f3dffcfaae","resolution":{"observed_at":"2026-08-07T01:03:45.699379Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.07374","last_updated":"2025-02-18T05:20:33Z","snapshot_observed_at":"2026-07-06T20:34:39.536540Z","submitted_at":"2025-02-11T08:48:48Z","title":"LLMs Can Easily Learn to Reason from Demonstrations Structure, not content, is what matters!","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.07374","snapshot_observed_at":"2026-08-07T01:03:43.321333Z","title":"Llms can easily learn to reason from demonstrations structure, not content, is what matters!,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.321333Z"},"links":{"cited_paper":"/paper/2502.07374","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:b6b55ce3c3ee351acfe500bc506de121f15073ad780282ca9760f551c6875826","observation_id":"b5856e66-e3bf-4112-b2c0-a838138b03f0","resolution":{"observed_at":"2026-08-07T01:03:43.321333Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:43.439679Z","title":"Shorterbetter: Guidingreasoningmodelstofindoptimalinferencelengthforefficient reasoning,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.439679Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:d50613ed7bfec4f646cca380fab92a009f63c70adccc23a21a037502a6c94b46","observation_id":"1af91da7-bec5-4818-8733-323b5f4b5934","resolution":{"observed_at":"2026-08-07T01:03:43.439679Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:45.332299Z","title":"Unlocking the capabilities of thought: A reasoning boundary framework to quantify and optimize chain-of-thought,","venue":null,"work_id":"08dbf383-56b3-4cdb-bcf3-380ae1905c70","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.510697Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:cbaa05962028324e16287909c83f142d331d6ffce18ae6b4dd29e15de59399fa","observation_id":"ee2f8254-4724-4d15-b0f4-e5a60cf4666e","resolution":{"observed_at":"2026-08-07T01:03:45.472186Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12599","last_updated":"2025-06-03T02:14:54Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-22T02:48:14Z","title":"Kimi k1.5: Scaling Reinforcement Learning with LLMs","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12599","snapshot_observed_at":"2026-08-07T01:03:43.598058Z","title":"Kimi k1. 5: Scaling reinforcement learning with llms,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.598058Z"},"links":{"cited_paper":"/paper/2501.12599","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:84a77bf700ef5b8746ec32c639dce34c8664834a2f5afa5d6393dbc4a7876e3f","observation_id":"1e93143a-2574-4054-be9a-79baf7e88ac1","resolution":{"observed_at":"2026-08-07T01:03:43.598058Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.04022","last_updated":"2025-04-05T02:24:07Z","snapshot_observed_at":"2026-08-07T16:08:28.700767Z","submitted_at":"2025-04-05T02:24:07Z","title":"Rethinking Reflection in Pre-Training","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.04022","snapshot_observed_at":"2026-08-07T01:03:43.678277Z","title":"Rethinking reflection in pre-training,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.678277Z"},"links":{"cited_paper":"/paper/2504.04022","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:1b0ff9bd9f8c4299b2ad414b3d80ac589b835667888e213d5136ba1166589082","observation_id":"c86d77b5-78b6-4952-baaf-a77e3547715d","resolution":{"observed_at":"2026-08-07T01:03:43.678277Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:45.103569Z","title":"Reflexion: Languageagentswithverbal reinforcement learning,","venue":null,"work_id":"c73c2aaf-8ec5-465e-8aaf-39a6e521d6fe","year":2023},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.774554Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:adfd99a7ce3a543c777aedd04221ddf0bd9cb5b7dda3857ca64b1b4b021340e7","observation_id":"ab5573a1-46ee-49f2-b945-5d27a0a59f08","resolution":{"observed_at":"2026-08-07T01:03:45.213263Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.20571","last_updated":"2025-10-24T10:02:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-04-29T09:24:30Z","title":"Reinforcement Learning for Reasoning in Large Language Models with One Training Example","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.20571","snapshot_observed_at":"2026-08-07T01:03:43.838361Z","title":"Rein- forcement learning for reasoning in large language models with one training example,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.838361Z"},"links":{"cited_paper":"/paper/2504.20571","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:4a33a70ec9573ebf09ee56dd011b275e18c0c044e29415da7fd7f7d61c8cff07","observation_id":"3fa9b162-1ea8-45d9-b1a8-faa954e327d7","resolution":{"observed_at":"2026-08-07T01:03:43.838361Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:44.974472Z","title":"There may not be aha moment in r1-zero-like training — a pilot study","venue":null,"work_id":"e5444a1f-56e7-4686-bfe8-778a818f14fd","year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.999176Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:662d071debbbfdfeaf0a279b09e596fb5fb2fde3d7def5d9ac5ac43b9669e9ef","observation_id":"33193c25-3851-45af-84ae-f7a83be9ac32","resolution":{"observed_at":"2026-08-07T01:03:45.036245Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-07T04:49:43.071685Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models"},"reference_resolution":{"displayed":60,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":31,"verified_exact":0,"verified_fuzzy":29},"total_outbound_references":60},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 60 of 60 outbound references and 1 inbound Pith citation observation for arXiv:2506.12217."}