{"as_of":"2026-08-08T09:35:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:7c7654646d5818d6337c37e4cc53496530e8420e100d0e81d6e77050a4916218","coverage":[{"denominator":43,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":43,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T23:10:27.460894Z","state":"measured"},{"denominator":44,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":44,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T17:54:17.954851Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-06T17:54:19.431591Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"cited_work":{"arxiv_id":"2506.19492","doi":null,"metadata_source":"pith","pith_arxiv_id":"2506.19492","snapshot_observed_at":"2026-08-06T17:54:19.431591Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","venue":"cs.CL","work_id":"25f70f04-0a19-4012-851c-ca91ed56f8f4","year":2025},"citing_paper":{"arxiv_id":"2507.09662","last_updated":"2025-07-13T14:51:59Z","snapshot_observed_at":"2026-08-07T01:15:50.475193Z","submitted_at":"2025-07-13T14:51:59Z","title":"Towards Concise and Adaptive Thinking in Large Reasoning Models: A Survey","version":1},"reference_index":220,"source":"arxiv_source","source_observed_at":"2026-08-06T17:54:17.954851Z"},"links":{"cited_paper":"/paper/2506.19492","citing_paper":"/paper/2507.09662"},"observation_digest":"sha256:12b4b4abdb95b718ab2a97324dd0153a58395e844dba0be48353c769bf5462e2","observation_id":"ba371586-8389-4de2-b0ec-df63a0b3ba18","resolution":{"observed_at":"2026-08-06T17:54:19.436865Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2506.19492/citation-record","integrity":"/paper/2506.19492/integrity","json":"/paper/2506.19492/citation-record.json","paper":"/paper/2506.19492"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"1606.06565","last_updated":"2016-07-25T17:23:29Z","snapshot_observed_at":"2026-07-06T05:00:46.434335Z","submitted_at":"2016-06-21T13:37:05Z","title":"Concrete Problems in AI Safety","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1606.06565","snapshot_observed_at":"2026-08-06T23:10:24.360508Z","title":"Concrete problems in ai safety.arXiv preprint arXiv:1606.06565, 2016","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:24.360508Z"},"links":{"cited_paper":"/paper/1606.06565","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:6311937e4172f9f0723d7d209620f48cb97bdfb36710fef79ec6c07618342b59","observation_id":"548782f0-c7a1-4232-8917-7d7783246534","resolution":{"observed_at":"2026-08-06T23:10:24.360508Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:24.429267Z","title":"Chain-of-thought reasoning in the wild is not always faithful","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:24.429267Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:6213ce4e381027b2817d8ca7671a097b1df83f6962e7b6616e763bc4c62af2ed","observation_id":"1d265062-f11e-46eb-b92c-36406210a084","resolution":{"observed_at":"2026-08-06T23:10:24.429267Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.08679","last_updated":"2026-06-16T17:36:22Z","snapshot_observed_at":"2026-08-07T17:12:11.913580Z","submitted_at":"2025-03-11T17:56:30Z","title":"Chain-of-Thought Reasoning In The Wild Is Not Always Faithful","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.08679","snapshot_observed_at":"2026-08-06T23:10:24.498510Z","title":"Chain-of-thought reasoning in the wild is not always faithful, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:24.498510Z"},"links":{"cited_paper":"/paper/2503.08679","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:bf15e7a4e069389c4b3869f3e574b164f82f7e25271250b00f6e0ff73f4ebf43","observation_id":"e4047ab3-1937-40f7-8b9c-f670e5675eb9","resolution":{"observed_at":"2026-08-06T23:10:24.498510Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:24.581791Z","title":"Training language models to reason efficiently.arXiv preprint arXiv:2502.04463, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:24.581791Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:6d973c6f92fa17a189062ef2f15464e68ff9171d5b28a10e924ee5a29b03e9ff","observation_id":"6bc124e3-dd59-4762-b3e3-a751462704b5","resolution":{"observed_at":"2026-08-06T23:10:24.581791Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.11926","last_updated":"2025-03-14T23:50:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-03-14T23:50:34Z","title":"Monitoring Reasoning Models for Misbehavior and the Risks of Promoting Obfuscation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.11926","snapshot_observed_at":"2026-08-06T23:10:24.679284Z","title":"Guan, Aleksander Madry, Wojciech Zaremba, Jakub Pachocki, and David Farhi","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:24.679284Z"},"links":{"cited_paper":"/paper/2503.11926","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:3f2939280c6788719b8a95c0a919a6a0e8110d1511c04d04e15966162ac00441","observation_id":"639c2dd4-23c0-422a-b5cf-7eb820033e8d","resolution":{"observed_at":"2026-08-06T23:10:24.679284Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:31.097257Z","title":"Monitoring reasoning models for misbehavior and the risks of promoting obfuscation","venue":null,"work_id":"c400e572-937f-42ab-a403-ef65c5041465","year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:24.755948Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:ef8ce55bff5094e19f5af7fee9bb64cf84ec8756aabac984f66935fb93e5f089","observation_id":"7c43f894-d01d-4e97-b5bc-acd855f5aac8","resolution":{"observed_at":"2026-08-06T23:10:31.101313Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:31.084242Z","title":"Weak-to-strong generalization: Eliciting strong capabilities with weak supervision, 2023","venue":null,"work_id":"03a7ca2e-780c-4b90-ac24-775a872c10cb","year":2023},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:24.826795Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:b3be6f1f66feb32ed304119377ec6c14f77d369661349ac68abe46452f2a331c","observation_id":"0183fd02-3508-4fa5-8ae8-1d3ce12980d9","resolution":{"observed_at":"2026-08-06T23:10:31.088807Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.21187","last_updated":"2025-02-01T07:57:37Z","snapshot_observed_at":"2026-08-01T16:43:44.704797Z","submitted_at":"2024-12-30T18:55:12Z","title":"Do NOT Think That Much for 2+3=? On the Overthinking of o1-Like LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.21187","snapshot_observed_at":"2026-08-06T23:10:24.922177Z","title":"Do not think that much for 2+ 3=? on the overthinking of o1-like llms.arXiv preprint arXiv:2412.21187, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:24.922177Z"},"links":{"cited_paper":"/paper/2412.21187","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:49bf885c876698b9e918a6ac608fe76ba7a80ba51d3d4b5bd862def05d4796ad","observation_id":"e96086b5-7dee-4f0d-96e2-315d89c06a51","resolution":{"observed_at":"2026-08-06T23:10:24.922177Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:31.071117Z","title":"Reasoning models don’t always say what they think","venue":null,"work_id":"13f4883b-bc8e-40f3-b0f0-5de8260c7c87","year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.009936Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:27d7bf200783af6bb37e81c3b2474388b4aecd6b6ba67dbc03194ab971d40573","observation_id":"60b1c90b-26cc-4740-b5d6-3888a514e631","resolution":{"observed_at":"2026-08-06T23:10:31.074962Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-06T23:10:25.135228Z","title":"Deepseek-r1: Incentivizing reasoning capability in llms via reinforcement learning, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.135228Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:13eae7b3bed6e1d3fbb8049124ac41b78c5eddcb2e5237e835b66f7d9bb3b807","observation_id":"07ca73a6-efda-4613-bfbc-57ee614e6c13","resolution":{"observed_at":"2026-08-06T23:10:25.135228Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10162","last_updated":"2024-06-29T00:28:47Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-14T16:26:20Z","title":"Sycophancy to Subterfuge: Investigating Reward-Tampering in Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10162","snapshot_observed_at":"2026-08-06T23:10:25.217290Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.217290Z"},"links":{"cited_paper":"/paper/2406.10162","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:4ee809fe55577bb3b2a98da4aea10cf17ea02d4401d7f782174ecfce10e06e8d","observation_id":"70f9d6c7-0bb0-43a5-8290-02d958728470","resolution":{"observed_at":"2026-08-06T23:10:25.217290Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:31.058411Z","title":"A wolf in sheep’s clothing: Generalized nested jailbreak prompts can fool large language models easily","venue":null,"work_id":"facc5e9d-99de-4aea-89d7-6ce60aca94a5","year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.285390Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:52e43ec335eb5c5cad0db9e40c7247b8e1146e31119938da031bd4762b105601","observation_id":"697d417d-401a-4482-a561-b06c548eb481","resolution":{"observed_at":"2026-08-06T23:10:31.062985Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:31.044850Z","title":"Reward tampering problems and solutions in reinforcement learning: A causal influence diagram perspective.Synthese, 198(Suppl 27): 6435–6467, 2021","venue":null,"work_id":"7ab4cbbe-ea82-4e3b-99c5-c053f04a9338","year":2021},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.380243Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:3219687562980bcbcb9375a00f3b739f98c6f50cd135b8acc93c0d2f321893b3","observation_id":"a0f602a2-9bc3-4d3f-8c74-1fa6d18c06b4","resolution":{"observed_at":"2026-08-06T23:10:31.049640Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:25.467968Z","title":"Syceval: Evaluating llm sycophancy.arXiv preprint arXiv:2502.08177, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.467968Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:9db49fb5e7b108e9acd3c83546f83dce57c08dc58746472612df72eb486dafad","observation_id":"851d42d2-bced-48be-af82-31191296ee6e","resolution":{"observed_at":"2026-08-06T23:10:25.467968Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:30.910505Z","title":"Who’s asking? user personas and the mechanics of latent misalignment.Advances in Neural Information Processing Systems, 37:125967–126003, 2024","venue":null,"work_id":"0056ff8b-d755-4a71-93d9-e3bf6abb194e","year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.541246Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:58b53479ec428760922bdf2cbe5f6283ddf380852999e907eda625fab448cf77","observation_id":"60f88d12-c1a4-4506-98c6-e08dc9d13f76","resolution":{"observed_at":"2026-08-06T23:10:31.005881Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:30.792289Z","title":"Alignment faking in large language models.CoRR, 2024","venue":null,"work_id":"9a850253-bf6b-4ab3-bc5f-a8ce7c7e9617","year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.613807Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:27e5ef5608c72d5ae06a715630d537fa5d853df9067c16bd074afb95800ba830","observation_id":"8a313f76-fa7e-49ca-9b82-9c58476ae472","resolution":{"observed_at":"2026-08-06T23:10:30.816339Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:25.729917Z","title":"Olympiadbench: A challenging benchmark for promoting agi with olympiad-level bilingual multimodal scientific problems","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.729917Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:bfc10e64ba6ce00a05640bf5b606640d9bebb42bb7cec222fda17e23b1201fa1","observation_id":"5c466036-a97c-4154-87d2-9e9a58e59176","resolution":{"observed_at":"2026-08-06T23:10:25.729917Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:25.812905Z","title":"C3ot: Generating shorter chain-of-thought without compromising effectiveness","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.812905Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:3713fd6f8d8b8ed533a05de01109225a07578c1f6f22949546c0a050647bc0d3","observation_id":"36329b67-aed0-429e-95b3-0c8af85eec41","resolution":{"observed_at":"2026-08-06T23:10:25.812905Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.01141","last_updated":"2025-04-01T00:41:36Z","snapshot_observed_at":"2026-08-08T04:02:43.828729Z","submitted_at":"2025-03-03T03:48:20Z","title":"How Well do LLMs Compress Their Own Chain-of-Thought? A Token Complexity Approach","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.01141","snapshot_observed_at":"2026-08-06T23:10:25.880397Z","title":"How well do llms compress their own chain-of-thought? a token complexity approach, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.880397Z"},"links":{"cited_paper":"/paper/2503.01141","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:3b81f71cbc207257dd1fd801ba886ff1f6cebeb662339f2c7d6c7c91a889093d","observation_id":"91f29c4a-bcc5-4e15-8794-2a82d943533f","resolution":{"observed_at":"2026-08-06T23:10:25.880397Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.13860","last_updated":"2024-03-10T13:58:08Z","snapshot_observed_at":"2026-07-06T15:31:18.144952Z","submitted_at":"2023-05-23T09:33:38Z","title":"Jailbreaking ChatGPT via Prompt Engineering: An Empirical Study","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.13860","snapshot_observed_at":"2026-08-06T23:10:25.963095Z","title":"Jailbreaking chatgpt via prompt engineering: An empirical study.arXiv preprint arXiv:2305.13860, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:25.963095Z"},"links":{"cited_paper":"/paper/2305.13860","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:7f9166a57f30a7ff326c659f5e55944130d3f0f5b75cc9c6c9efab5fc65f68a7","observation_id":"dabfeedd-0961-4400-934d-aab56adf40b8","resolution":{"observed_at":"2026-08-06T23:10:25.963095Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12570","last_updated":"2025-01-29T03:11:03Z","snapshot_observed_at":"2026-08-08T03:22:47.699927Z","submitted_at":"2025-01-22T01:35:11Z","title":"O1-Pruner: Length-Harmonizing Fine-Tuning for O1-Like Reasoning Pruning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12570","snapshot_observed_at":"2026-08-06T23:10:26.017771Z","title":"O1-pruner: Length-harmonizing fine-tuning for o1-like reasoning pruning.arXiv preprint arXiv:2501.12570, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.017771Z"},"links":{"cited_paper":"/paper/2501.12570","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:df7c99a3e670aba507ce50f354c3417892c7b179775ebf18cdbd4c7a2f010ccf","observation_id":"23da9ed6-0f8d-41f2-bffa-d936c0ace837","resolution":{"observed_at":"2026-08-06T23:10:26.017771Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.09858","last_updated":"2025-04-14T04:08:16Z","snapshot_observed_at":"2026-08-08T01:06:59.135762Z","submitted_at":"2025-04-14T04:08:16Z","title":"Reasoning Models Can Be Effective Without Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.09858","snapshot_observed_at":"2026-08-06T23:10:26.064315Z","title":"Reasoning models can be effective without thinking, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.064315Z"},"links":{"cited_paper":"/paper/2504.09858","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:cb6fa328a7f74a1535b53c8915b8eb613c446f93ce7375c356b0465ecf32311b","observation_id":"157ca23f-801b-430f-8ecb-79ca607520ae","resolution":{"observed_at":"2026-08-06T23:10:26.064315Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:30.631840Z","title":"Are self-explanations from large language models faithful? InFindings of the Association for Computational Linguistics ACL 2024, pages 295–337, 2024","venue":null,"work_id":"5fa3de32-a4b2-423a-85cc-2baa44ffd2e1","year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.111750Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:204e9f7afe79a2a239c3cff920f1524f8f149776b1baee90d2bf86c5b144831b","observation_id":"302d724a-88a7-427d-8cea-bfac70a48a75","resolution":{"observed_at":"2026-08-06T23:10:30.727087Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.04984","last_updated":"2025-01-14T20:16:01Z","snapshot_observed_at":"2026-07-29T23:20:20.918596Z","submitted_at":"2024-12-06T12:09:50Z","title":"Frontier Models are Capable of In-context Scheming","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.04984","snapshot_observed_at":"2026-08-06T23:10:26.170863Z","title":"Frontier models are capable of in-context scheming.arXiv preprint arXiv:2412.04984, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.170863Z"},"links":{"cited_paper":"/paper/2412.04984","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:ed19621d9d6ad2c811824261dc1d7db690d3ae0e28c1a1fa65b2928916139862","observation_id":"8d016c3c-1c87-4295-98dd-c8eb93774d26","resolution":{"observed_at":"2026-08-06T23:10:26.170863Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:30.503889Z","title":"Self-training elicits concise reasoning in large language models.CoRR, 2025","venue":null,"work_id":"a87297d3-be05-4573-bf70-bbfb5ca1f649","year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.236575Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:2831378f05ad50d2aef647ceaa4498d8c5ce72ee5b494afddc0754ae444e1ef7","observation_id":"28c01fdc-bcfb-41dd-9cbc-175f53360a1a","resolution":{"observed_at":"2026-08-06T23:10:30.583304Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:30.315614Z","title":"Show your work: Scratchpads for intermediate computation with language models","venue":null,"work_id":"daaf1774-d428-4d62-ac48-f1f3aea741d5","year":null},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.289426Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:5a90712bf124a141c3c54e25f1cb11fcac25070598730ce24b010fba4a6ad7bb","observation_id":"50a0b1e3-21f2-4cd2-91b4-da88d7ff332f","resolution":{"observed_at":"2026-08-06T23:10:30.395731Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:30.061662Z","title":"Learning to reason with llms","venue":null,"work_id":"03964fca-9905-4d2b-b516-e97c1cf1dc02","year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.375978Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:1739c5444aa2027cf90750cc178cdee85d84be7212965fc4d166acba9e8c8510","observation_id":"50b54a60-db45-4499-b960-bc72af29f821","resolution":{"observed_at":"2026-08-06T23:10:30.195407Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:26.439725Z","title":"Discovering language model behaviors with model-written evaluations","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.439725Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:bbb8a144cef5b9905155028f6196bb558317f0f976e2d92c561938a13a053e24","observation_id":"41429ca3-8eaf-4e6b-b9b1-ce061caa908c","resolution":{"observed_at":"2026-08-06T23:10:26.439725Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:29.837598Z","title":"Do Anything Now","venue":null,"work_id":"445b6b54-e098-4fb9-ac06-fb4f024c1f8a","year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.505238Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:7188a9193added65d21e055800ae7d2d099b63ab19fed8f13608d4cd90b5a388","observation_id":"55b4560e-fccb-42ea-9770-3ae2d9078c4c","resolution":{"observed_at":"2026-08-06T23:10:29.918948Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:26.574076Z","title":"Defining and characterizing reward gaming.Advances in Neural Information Processing Systems, 35:9460–9471, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.574076Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:59cdf59184fe10ccfb2dc49d1bc49621c9d0c5970a4a8d5d4b7ac4cb6e041b76","observation_id":"6293ab38-fd35-4d91-aa75-d87d80608643","resolution":{"observed_at":"2026-08-06T23:10:26.574076Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.16419","last_updated":"2025-08-21T19:14:40Z","snapshot_observed_at":"2026-08-07T04:27:23.738927Z","submitted_at":"2025-03-20T17:59:38Z","title":"Stop Overthinking: A Survey on Efficient Reasoning for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.16419","snapshot_observed_at":"2026-08-06T23:10:26.642653Z","title":"Stop overthinking: A survey on efficient reasoning for large language models.arXiv preprint arXiv:2503.16419, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.642653Z"},"links":{"cited_paper":"/paper/2503.16419","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:0cf2d023d4bde6a203ea6736e1b73c47ef9d7b95f3a820434e274f82b56f7f19","observation_id":"a15feb78-3f2c-49ae-835a-4615d4cebdbd","resolution":{"observed_at":"2026-08-06T23:10:26.642653Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12599","last_updated":"2025-06-03T02:14:54Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-22T02:48:14Z","title":"Kimi k1.5: Scaling Reinforcement Learning with LLMs","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12599","snapshot_observed_at":"2026-08-06T23:10:26.694880Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.694880Z"},"links":{"cited_paper":"/paper/2501.12599","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:33370bf688a9c62d6c6856cc11228ad0c47e9c2f2f7e30f7e9a75208cae93145","observation_id":"b04d6c25-3c44-4b35-8bad-021012d32e03","resolution":{"observed_at":"2026-08-06T23:10:26.694880Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:26.755752Z","title":"Qwen3, April 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.755752Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:b21f38b412632419c8897338c814b1ae9add156c57252473058a3f76daf27138","observation_id":"b33198e8-c487-45a2-a93a-dbb6066d5cb3","resolution":{"observed_at":"2026-08-06T23:10:26.755752Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:26.845600Z","title":"Language models don’t always say what they think: Unfaithful explanations in chain-of-thought prompting.Advances in Neural Information Processing Systems, 36:74952–74965, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.845600Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:817f72d7ddd1f7a4b4079be0c92be0dff868d57ad7e4c94c4c758323a02bebb0","observation_id":"838e7af5-614b-4797-ba56-5c4ae7712324","resolution":{"observed_at":"2026-08-06T23:10:26.845600Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:29.495544Z","title":null,"venue":null,"work_id":"166b0865-7a90-4f6f-854a-5092e9dd5bab","year":2024},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.912149Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:c01808311a6c56fe34486539b2638b78e28d5606b92280c5fe0af6e3f8a803c6","observation_id":"f6c5e37f-e127-415b-9a2c-0fb8a1b965af","resolution":{"observed_at":"2026-08-06T23:10:29.642751Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:29.214889Z","title":"Large language models often say one thing and do another","venue":null,"work_id":"7cc95cbc-e60e-4ddd-93f3-15b005a47483","year":2025},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:26.973650Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:347fb652f7a953310d72269053b84baddc678bbc8f86441a29bf6c895a127971","observation_id":"887d404b-d7e8-4e37-9f48-572700bbd90f","resolution":{"observed_at":"2026-08-06T23:10:29.323877Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.15043","last_updated":"2023-12-20T20:48:57Z","snapshot_observed_at":"2026-07-06T15:59:23.019044Z","submitted_at":"2023-07-27T17:49:12Z","title":"Universal and Transferable Adversarial Attacks on Aligned Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.15043","snapshot_observed_at":"2026-08-06T23:10:27.042013Z","title":"no free lunch","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:27.042013Z"},"links":{"cited_paper":"/paper/2307.15043","citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:d2f238fc831ac314f01cc74aaa36ab4f6d0c1142b25d536a535ba7a69986c02e","observation_id":"653ad9f5-5790-4dd2-a177-8a3e4883d017","resolution":{"observed_at":"2026-08-06T23:10:27.042013Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:28.981803Z","title":null,"venue":null,"work_id":"487b18f5-0c96-4652-abff-834448203783","year":null},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:27.100826Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:e75aa9d03334d8b9cdf5b82a3f4b1d306c58b3007e17c0c898a5002eece03436","observation_id":"f57189ff-5941-47fe-9c36-df4163c9c159","resolution":{"observed_at":"2026-08-06T23:10:29.089201Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:28.770920Z","title":"IC” if the two responses show an inconsistency (refusal vs. answer, or different final results). - “CO","venue":null,"work_id":"0e4fa3b7-ef8c-4e29-9481-e000255667e2","year":null},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:27.144731Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:6818adffff674752d49adf655644db4e09b75c87d462ad38b036195b9d083b40","observation_id":"471e33ee-960e-48f7-9c1a-1ec2a0a71a2f","resolution":{"observed_at":"2026-08-06T23:10:28.869126Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:28.559443Z","title":null,"venue":null,"work_id":"25269b88-cc67-47aa-ad54-4cfa020bf629","year":null},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:27.213843Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:5bb99caa44d91008f3eab978a9eb384fbf5d7767f7629a867db6501d0eaf3f00","observation_id":"9b4843ca-055d-464b-bd7b-e6d392f0c876","resolution":{"observed_at":"2026-08-06T23:10:28.670354Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:28.335035Z","title":null,"venue":null,"work_id":"07d205a9-2760-4098-8a80-8f032a8cf968","year":null},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:27.293021Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:88bef6bd5464a95fc6771c48660dec265ad9d34a843155cda0c9819b875661a9","observation_id":"f0f8efbe-7847-4426-a1dd-a6774df564c8","resolution":{"observed_at":"2026-08-06T23:10:28.465354Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:28.119098Z","title":null,"venue":null,"work_id":"29879c81-d06b-4ba8-85ad-31504422820d","year":null},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:27.372720Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:66ba3efa4993bf3cdaa6f4dceabce5afe2126a0075bce32d0096745c99a32cd4","observation_id":"ba0e5744-defa-4b10-aa7e-f23fd86b589b","resolution":{"observed_at":"2026-08-06T23:10:28.200050Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:10:27.929374Z","title":"Deploymode","venue":null,"work_id":"89104294-8e9d-4852-bee4-7fc7971e2152","year":null},"citing_paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:27.460894Z"},"links":{"citing_paper":"/paper/2506.19492"},"observation_digest":"sha256:fda9dad7860b61a25af70a02270329dcc0f1316186c477cd386b752598b12f22","observation_id":"fa8d8b32-6988-4364-92fe-122977d98eef","resolution":{"observed_at":"2026-08-06T23:10:28.026131Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.19492","last_updated":"2025-06-24T10:25:28Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-07T09:55:06.788380Z","submitted_at":"2025-06-24T10:25:28Z","title":"Is Long-to-Short a Free Lunch? Investigating Inconsistency and Reasoning Efficiency in LRMs"},"reference_resolution":{"displayed":43,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":28,"verified_exact":0,"verified_fuzzy":15},"total_outbound_references":43},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 43 of 43 outbound references and 1 inbound Pith citation observation for arXiv:2506.19492."}