{"as_of":"2026-08-05T22:51:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:2f224afbf70b11262f7e38d27dba95c2247a8407662ca605808735ef13548e6d","coverage":[{"denominator":24,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":24,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-26T12:21:42.088795Z","state":"measured"},{"denominator":24,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":24,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-05T06:32:48.257954+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2606.21906/citation-record","integrity":"/paper/2606.21906/integrity","json":"/paper/2606.21906/citation-record.json","paper":"/paper/2606.21906"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2112.00861","last_updated":"2021-12-09T21:40:22Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2021-12-01T22:24:34Z","title":"A General Language Assistant as a Laboratory for Alignment","version":3},"cited_work":{"arxiv_id":"2112.00861","doi":"10.48550/arxiv.2112.00861","metadata_source":"pith","pith_arxiv_id":"2112.00861","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A General Language Assistant as a Laboratory for Alignment","venue":"cs.CL","work_id":"a43f9ea0-01be-47d5-b8ee-a1a9f73381c5","year":2021},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"cited_paper":"/paper/2112.00861","citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:2e8bfc0578087fc2f549ab2f07bb645e6234c746192cc20992ee5b2496b0c7ce","observation_id":"fc86705e-cba9-471b-a09d-89369334a3e5","resolution":{"observed_at":"2026-07-04T07:59:40.105129Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-05-20T18:52:13.802987+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T18:52:13.802987+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08112","last_updated":"2025-11-11T01:13:28Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-03-14T17:47:09Z","title":"Eliciting Latent Predictions from Transformers with the Tuned Lens","version":6},"cited_work":{"arxiv_id":"2303.08112","doi":"10.48550/arxiv.2303.08112","metadata_source":"pith","pith_arxiv_id":"2303.08112","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Eliciting Latent Predictions from Transformers with the Tuned Lens","venue":"cs.LG","work_id":"a127314f-7424-488f-b6d7-8214650c420f","year":2023},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"cited_paper":"/paper/2303.08112","citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:bdbae780cd4fce68ae761c0936d5c7a5983c727de4d9c003c1922ff8240aa26c","observation_id":"56a2b02b-9856-476c-891e-c102a7078b3a","resolution":{"observed_at":"2026-07-04T07:59:40.102550Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-06-01T22:57:41.570848+00:00","source":"crossref_status_cache"},{"observed_at":"2026-06-01T22:57:41.570848+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2110.14168","last_updated":"2021-11-18T00:23:45Z","snapshot_observed_at":"2026-08-04T15:46:25.710484Z","submitted_at":"2021-10-27T04:49:45Z","title":"Training Verifiers to Solve Math Word Problems","version":2},"cited_work":{"arxiv_id":"2110.14168","doi":"10.1002/j.1545-","metadata_source":"pith","pith_arxiv_id":"2110.14168","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Training Verifiers to Solve Math Word Problems","venue":"cs.LG","work_id":"acab1aa8-b4d6-40e0-a3ee-25341701dca2","year":2021},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"cited_paper":"/paper/2110.14168","citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:5d7ce7d997b61a7b88f34117bdf19d3a51529fc1404d8e7d2a1c12e6536740bd","observation_id":"826b59a2-a2ce-4fa0-acb6-7e42761a6d72","resolution":{"observed_at":"2026-07-04T07:59:40.099944Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T12:21:42.088795Z","title":"Transformer feed-forward layers build predictions by promoting concepts in the vocabulary space","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:663ca4c69861586c37c3d3b3119339bdf227a282fa5f7b652583da74b3860f7b","observation_id":"859a0e6c-20e9-48ca-b059-040568c823db","resolution":{"observed_at":"2026-06-26T12:21:42.088795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T12:21:42.088795Z","title":"Dissecting recall of factual associations in auto-regressive language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:cbd39bb28e9a49b24436712e2d03b456b2a39b7ac5fe1b53968b0d653bb13776","observation_id":"d00e2c91-06fb-4487-bfe6-95648d3f4e3a","resolution":{"observed_at":"2026-06-26T12:21:42.088795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2510.18871","doi":"10.48550/arxiv.2510.18871","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"How do llms use their depth?arXiv preprint arXiv:2510.18871","venue":"arXiv (Cornell University)","work_id":"776ea885-519f-4790-8e33-cac642cc672d","year":2026},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:d8f43c9e0902c68890f42dcc29b54c31384d749f2da1149f239c4eee48d061ef","observation_id":"952e87b1-bba6-49f9-9622-bdbbc1374099","resolution":{"observed_at":"2026-07-04T07:59:40.122455Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T12:21:42.088795Z","title":"Inference-time intervention: Eliciting truthful answers from a language model","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:e2aae1a33240cb3e1da35bce08f3978047fca1c4456a734b7d4655b4d0860e48","observation_id":"c3c03b33-93ae-4c1e-9b8e-a3dc965aa4c5","resolution":{"observed_at":"2026-06-26T12:21:42.088795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2601.03542","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T07:59:40.108882Z","title":"Layer-order inversion: Rethinking latent multi-hop reasoning in large language models.arXiv preprint arXiv:2601.03542,","venue":null,"work_id":"983e9b6f-0633-4c08-aa92-57cf0ffb82b0","year":2026},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:7a9a423a7e5460bab8d02b6e086e52193b389488babb89e98d0379c555d8739d","observation_id":"c3e836b1-9439-427f-9b68-4902d744ace3","resolution":{"observed_at":"2026-07-04T07:59:40.110778Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T12:21:42.088795Z","title":"Diffskip: Differential layer skipping in large language models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:36fdb3042933475eb4986a91a8eeff1f98446ca3a988a2efbe7251dc6607f7ea","observation_id":"a5ce01a8-a7d6-474a-b21e-1c559e159c14","resolution":{"observed_at":"2026-06-26T12:21:42.088795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T12:21:42.088795Z","title":"Accessed: 2026-02-16","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:0016ba0f5d0ba70b3d6a0a3e81f45193a908bbc9f9e9c6acb964f1163e319648","observation_id":"cfc81bef-5d02-4bb7-92fe-c6c2037968f8","resolution":{"observed_at":"2026-06-26T12:21:42.088795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T12:21:42.088795Z","title":"Trusting your evidence: Hallucinate less with context-aware decoding","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:806a88b17aee601e1c7142c0ef3719ede44cd3aa1ba19c742417e15998e63667","observation_id":"ab7cb403-a851-4411-a81c-8d902bbfba7b","resolution":{"observed_at":"2026-06-26T12:21:42.088795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2603.23701","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T07:59:40.115162Z","title":"The diminishing returns of early-exit decoding in modern llms.arXiv preprint arXiv:2603.23701,","venue":null,"work_id":"5b9b9669-fa4e-42ed-8ca0-4ff5dce01e17","year":null},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:90d92c99dcb2e1aa25043e1f530cf142f240d79e2353194347e607c48b30d7d2","observation_id":"abbebdb4-c20b-4a27-a55e-1885be4bd1e2","resolution":{"observed_at":"2026-07-04T07:59:40.116684Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2504.15895","doi":"10.48550/arxiv.2504.15895","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Dynamic early exit in reasoning models.arXiv preprint arXiv:2504.15895, 2025a","venue":"ArXiv.org","work_id":"99c0642f-2c3c-4bfd-aa41-08cd3fb032e8","year":2025},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:5670cd44bdb95c26c1b04b99b7b40b8fdaf8caf6fe18bb78cef81a3b713ddeab","observation_id":"4a0a2670-43a2-498c-8f14-7aa5062c04d8","resolution":{"observed_at":"2026-07-04T07:59:40.113991Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T12:21:42.088795Z","title":"Air-bench 2024: A safety benchmark based on regulation and policies specified risk categories","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:91a6348fe1c21cab414b29dfc7bdf099a120bbf177a9ef1b8516154341fc99b5","observation_id":"adb5a8cd-525f-4eb4-ba82-8648338c1a5f","resolution":{"observed_at":"2026-06-26T12:21:42.088795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T12:21:42.088795Z","title":"Active layer-contrastive decoding reduces halluci- nation in large language model generation","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:db57d0a17e18b280a928dda505f608677d5a3dcc83863150ecb0c1d7dfd29db1","observation_id":"ea2ff94a-9bc1-461c-b619-a4c600dbde48","resolution":{"observed_at":"2026-06-26T12:21:42.088795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T12:21:42.088795Z","title":"Generalization or memorization: Dynamic decoding for mode steering","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:9de1f72e72e3601c90b83d3f6851e34a60923cb936f79e99e7e47046dcf00369","observation_id":"79944bff-c872-4e80-9dd5-38ac563f8eec","resolution":{"observed_at":"2026-06-26T12:21:42.088795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T12:21:42.088795Z","title":"Cognition-of-thought elicits social-aligned reasoning in large language models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:931b1e7e746b363034340fe69fa48eea4ea9d4ac5e25307524a06c44ffa5186e","observation_id":"226e3c31-4157-4fa4-b06e-8b18b4e1ae82","resolution":{"observed_at":"2026-06-26T12:21:42.088795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2507.04404","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T07:59:40.118007Z","title":"Please describe the image","venue":null,"work_id":"5780380b-21c8-473f-aee6-2c63c196f66f","year":2025},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:57c3ae4cfeac3ca3580848b31d31504448338a7dbd0e1de58aae70f080e24648","observation_id":"98e93fc6-c7c6-42b3-97e9-521667ff5931","resolution":{"observed_at":"2026-07-04T07:59:40.119866Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T12:21:42.088795Z","title":"Continual pre-training of language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:26708ac1d60fccd3d09c477d18dcb703b9b246185730ff1c7f92d084acb04022","observation_id":"9d1fd99d-8986-45af-ba9e-8d4a8fe9c7d3","resolution":{"observed_at":"2026-06-26T12:21:42.088795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.01405","last_updated":"2025-03-03T06:14:14Z","snapshot_observed_at":"2026-07-06T16:26:38.284922Z","submitted_at":"2023-10-02T17:59:07Z","title":"Representation Engineering: A Top-Down Approach to AI Transparency","version":4},"cited_work":{"arxiv_id":"2310.01405","doi":"10.1609/aaai.v36i9.21196","metadata_source":"pith","pith_arxiv_id":"2310.01405","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Representation Engineering: A Top-Down Approach to AI Transparency","venue":"cs.LG","work_id":"45b326e2-e962-41a5-a542-2559e103a19b","year":2023},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"cited_paper":"/paper/2310.01405","citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:28b2613f2bffbd3cfe19014cb96263c52495968057555b6083124d0b4380f4fd","observation_id":"b026f487-944f-4156-a230-b616cc5b548b","resolution":{"observed_at":"2026-07-04T07:59:40.107743Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T12:21:42.088795Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:25b82557f2ee48dbace4e1249a81b22935f599cc60dfaa3f93500cf56c287a49","observation_id":"f9842565-f8c5-4a8d-bbec-e797a694ef28","resolution":{"observed_at":"2026-06-26T12:21:42.088795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T12:21:42.088795Z","title":"You should think step-by-step and put your final answer within \\boxed{}","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:4dc3c5a977dc077c962058961900d1087b90e7a5f02fe25cc98cf969fcaaedbb","observation_id":"f0ec6ebf-09d1-417b-9619-88a3e4d27942","resolution":{"observed_at":"2026-06-26T12:21:42.088795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T12:21:42.088795Z","title":"You are an expert evaluator with extensive experience in evaluating response of given query","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:b9d14deda6ee6823f8ded1857ad1820997a282a587ced4872cb3c2f8583ccd98","observation_id":"2999d430-0b19-48e5-bd0d-2b75c1082b5e","resolution":{"observed_at":"2026-06-26T12:21:42.088795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T12:21:42.088795Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-06-26T12:21:42.088795Z"},"links":{"citing_paper":"/paper/2606.21906"},"observation_digest":"sha256:df8e7cebc2cd1abd4ce8898e931c62f59f7f4ff8072f74c4d8397d600d5e568f","observation_id":"930d7d20-bda1-495e-a295-ba8190150f63","resolution":{"observed_at":"2026-06-26T12:21:42.088795Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2606.21906","last_updated":"2026-06-20T07:03:26Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-03T21:45:33.488775Z","submitted_at":"2026-06-20T07:03:26Z","title":"Deeper is Not Always Better: Mitigating the Alignment Tax via Confident Layer Decoding"},"reference_resolution":{"displayed":24,"state_counts":{"malformed_identifier":1,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":14,"verified_exact":8,"verified_fuzzy":0},"total_outbound_references":24},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"thesis":"As of 5 August 2026, this Paper Citation Record lists 24 of 24 outbound references and 0 inbound Pith citation observations for arXiv:2606.21906."}