{"as_of":"2026-08-17T17:08:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:8f5cdb7095cc6c0655996464d31998a169a16c1000271e20509211632ab02605","coverage":[{"denominator":51,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":51,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-14T04:31:15.712697Z","state":"measured"},{"denominator":51,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":51,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-17T06:30:58.91139+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2608.08829/citation-record","integrity":"/paper/2608.08829/integrity","json":"/paper/2608.08829/citation-record.json","paper":"/paper/2608.08829"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:15.513741Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.513741Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:297c627f66e0c5bc7f98c493da610e26394f4137c5645621ee40d57140706ccb","observation_id":"3a328d52-7d1c-4cb9-8145-4a467ce677f8","resolution":{"observed_at":"2026-08-14T04:31:15.513741Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2601.21505","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.254829Z","title":null,"venue":null,"work_id":"9bb5a31e-f4c8-42c1-b5cf-91c01f352ea5","year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.522859Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:cae04fb80001c32f56e81474b1b009993f3adc3701a1289b0404883c10dd9eed","observation_id":"29b639c1-6335-47c6-b0db-1bd3914264b8","resolution":{"observed_at":"2026-08-14T04:31:16.261038Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2604.03867","last_updated":"2026-04-04T21:16:47Z","snapshot_observed_at":"2026-08-15T03:27:56.018408Z","submitted_at":"2026-04-04T21:16:47Z","title":"Where to Steer: Input-Dependent Layer Selection for Steering Improves LLM Alignment","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2604.03867","snapshot_observed_at":"2026-08-14T04:31:15.527593Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.527593Z"},"links":{"cited_paper":"/paper/2604.03867","citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:7dc2adcd984115c9e8f3c271d3ca0750200e1aefd987df81322d217cec6268fc","observation_id":"23a4b0c4-4adb-4c6c-985e-36a823a07543","resolution":{"observed_at":"2026-08-14T04:31:15.527593Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:15.536458Z","title":"Amoukou, Tom Bewley, Saumitra Mishra, and Manuela Veloso","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.536458Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:2555c366b9a8f765107e56b505072b0cd60ba17d3d9b3dc636a717fd574524cb","observation_id":"99f4b790-ca0e-45e7-be4d-a22e474b0fac","resolution":{"observed_at":"2026-08-14T04:31:15.536458Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.701158Z","title":"a rvelin and Jaana Kek \\","venue":null,"work_id":"430ac867-0588-4998-8279-c4c01ec4bfcf","year":2002},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.540407Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:89ec78f70d17d2960d3cbea8dd3ae0cd71d622e900d8b061a1285bc17a0d7ce3","observation_id":"8380b5bd-f92d-47e3-ae7f-163f8a100f40","resolution":{"observed_at":"2026-08-14T04:31:16.705448Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.05907","last_updated":"2025-02-17T20:23:19Z","snapshot_observed_at":"2026-08-16T13:20:42.975625Z","submitted_at":"2024-09-06T15:47:40Z","title":"Programming Refusal with Conditional Activation Steering","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.05907","snapshot_observed_at":"2026-08-14T04:31:15.544135Z","title":"Lee, Inkit Padhi, Karthikeyan Natesan Ramamurthy, Erik Miehling, Pierre Dognin, Manish Nagireddy, and Amit Dhurandhar","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.544135Z"},"links":{"cited_paper":"/paper/2409.05907","citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:e31d7fcee5caa01242a296ac48d61c61e9bab1a7d8d2a164569740bb0f9c499a","observation_id":"401c321e-44b5-4132-a258-8b30ad40b35a","resolution":{"observed_at":"2026-08-14T04:31:15.544135Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.03341","last_updated":"2024-06-26T14:11:53Z","snapshot_observed_at":"2026-08-13T19:18:46.730073Z","submitted_at":"2023-06-06T01:26:53Z","title":"Inference-Time Intervention: Eliciting Truthful Answers from a Language Model","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.03341","snapshot_observed_at":"2026-08-14T04:31:15.548550Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.548550Z"},"links":{"cited_paper":"/paper/2306.03341","citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:d180c8973f300f98e56a44ceea8924978fae170eecdaef610803abad7e6a37e1","observation_id":"5ee537b1-17de-4b68-a925-170ee5431af7","resolution":{"observed_at":"2026-08-14T04:31:15.548550Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.12446","last_updated":"2026-07-10T01:07:20Z","snapshot_observed_at":"2026-08-16T12:57:30.396986Z","submitted_at":"2025-02-18T02:27:23Z","title":"Multi-Attribute Steering of Language Models via Targeted Intervention","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.12446","snapshot_observed_at":"2026-08-14T04:31:15.553780Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.553780Z"},"links":{"cited_paper":"/paper/2502.12446","citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:29cc139152286ad6b1ffd53cb000036129ba32d1e934719e4f5a32ea5d8698cf","observation_id":"41c0fca4-3cc3-4e06-9cec-6ea3764f0209","resolution":{"observed_at":"2026-08-14T04:31:15.553780Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.24535","last_updated":"2026-04-04T17:56:48Z","snapshot_observed_at":"2026-08-16T13:14:10.143261Z","submitted_at":"2025-05-30T12:41:19Z","title":"Beyond Linear Steering: Unified Multi-Attribute Control for Language Models","version":3},"cited_work":{"arxiv_id":"2505.24535","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.24535","snapshot_observed_at":"2026-08-14T04:31:16.077945Z","title":"Beyond Linear Steering: Unified Multi-Attribute Control for Language Models","venue":"cs.LG","work_id":"43aff66d-a237-4647-ae88-f600b21eaae1","year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.557756Z"},"links":{"cited_paper":"/paper/2505.24535","citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:2d68c5a373312ed747a6ac7311a163eb1d137d572604becb090fdad308bfbe05","observation_id":"2e53a852-784a-4e55-bd11-f5bf1633323e","resolution":{"observed_at":"2026-08-14T04:31:16.083075Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2508.12815","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.060639Z","title":null,"venue":null,"work_id":"a9654fc2-8660-429d-ad56-005d8eeaf0bc","year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.561567Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:667ec0ab169a12d2d890f8e9be19e5c0eb7aa8fb19521a58ef4af9701ae1b2fb","observation_id":"12c3a859-4411-430e-b96a-81575a2249d6","resolution":{"observed_at":"2026-08-14T04:31:16.066459Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2212.09251","last_updated":"2022-12-19T05:13:52Z","snapshot_observed_at":"2026-08-14T22:08:45.647927Z","submitted_at":"2022-12-19T05:13:52Z","title":"Discovering Language Model Behaviors with Model-Written Evaluations","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2212.09251","snapshot_observed_at":"2026-08-14T04:31:15.565020Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.565020Z"},"links":{"cited_paper":"/paper/2212.09251","citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:769ca90c1ecee0ef26a4820fe350949cdda45b147c7bcaa6977685a50eafa79d","observation_id":"c5f6e279-5088-4b1f-a9d3-ed7c968dc54e","resolution":{"observed_at":"2026-08-14T04:31:15.565020Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.06681","last_updated":"2024-07-05T15:30:45Z","snapshot_observed_at":"2026-08-15T08:04:06.283165Z","submitted_at":"2023-12-09T04:40:46Z","title":"Steering Llama 2 via Contrastive Activation Addition","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.06681","snapshot_observed_at":"2026-08-14T04:31:15.568778Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.568778Z"},"links":{"cited_paper":"/paper/2312.06681","citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:25cf33b428fb0b8eddde6528a8b73d5e2755b2637d82b9170c9af59e683073a9","observation_id":"3fad16b0-bd9b-4aa4-953e-b56195579704","resolution":{"observed_at":"2026-08-14T04:31:15.568778Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.23054","last_updated":"2024-11-22T16:04:44Z","snapshot_observed_at":"2026-08-16T13:04:32.014923Z","submitted_at":"2024-10-30T14:21:33Z","title":"Controlling Language and Diffusion Models by Transporting Activations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.23054","snapshot_observed_at":"2026-08-14T04:31:15.572573Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.572573Z"},"links":{"cited_paper":"/paper/2410.23054","citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:69ea4d4f86ff6c4f5d4997e8049d7343540c63460408ef93261d19fffc4887aa","observation_id":"d4313d60-9e4b-49cc-b6fb-7a1bcc1ce14a","resolution":{"observed_at":"2026-08-14T04:31:15.572573Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.17563","last_updated":"2024-06-25T14:00:42Z","snapshot_observed_at":"2026-08-16T13:39:45.805437Z","submitted_at":"2024-06-25T14:00:42Z","title":"Multi-property Steering of Large Language Models with Dynamic Activation Composition","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.17563","snapshot_observed_at":"2026-08-14T04:31:15.576336Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.576336Z"},"links":{"cited_paper":"/paper/2406.17563","citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:a55beb39e0127fada68600c542c6499b2053bc690af0a9728f6144c5c1cabdcc","observation_id":"0752f5b6-2cae-4141-bcef-cfe8b8357741","resolution":{"observed_at":"2026-08-14T04:31:15.576336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.688492Z","title":null,"venue":null,"work_id":"0d8afd6b-9f8a-40a9-ac98-2c33c87999ae","year":1953},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.580188Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:5472d9e22e3edb98c4815c86b7b83b93d29c6301d00a739cb5eb1de21aad4926","observation_id":"e103ba08-32cd-4c34-8d47-295dcc60f336","resolution":{"observed_at":"2026-08-14T04:31:16.692327Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09929","last_updated":"2025-04-02T13:20:40Z","snapshot_observed_at":"2026-08-15T14:36:36.921298Z","submitted_at":"2025-01-17T02:55:23Z","title":"Interpretable Steering of Large Language Models with Feature Guided Activation Additions","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.09929","snapshot_observed_at":"2026-08-14T04:31:15.583725Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.583725Z"},"links":{"cited_paper":"/paper/2501.09929","citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:d7b79ffc57b6388c1904dfd94828d900bd576a9821988816f14c2cbddb30ecc3","observation_id":"fcf235cf-9bf0-4e30-9026-1235a3239dac","resolution":{"observed_at":"2026-08-14T04:31:15.583725Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.675711Z","title":null,"venue":null,"work_id":"cd13565b-c876-4fbf-9d9f-8a0991e0a4c3","year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.587637Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:d2211031168cc5c043365ba2816f26f7f2944eb27511a55d080397111f9c96ed","observation_id":"a6c10ae6-841b-428a-949f-4042b87a9e58","resolution":{"observed_at":"2026-08-14T04:31:16.679688Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.12404","last_updated":"2025-05-04T23:07:56Z","snapshot_observed_at":"2026-08-16T13:33:30.744227Z","submitted_at":"2024-07-17T08:32:03Z","title":"Analyzing the Generalization and Reliability of Steering Vectors","version":8},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.12404","snapshot_observed_at":"2026-08-14T04:31:15.591259Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.591259Z"},"links":{"cited_paper":"/paper/2407.12404","citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:1af2ceeb5053e8eef3a29700fd60fa23b1977808d9a08daf821f3663b21ee515","observation_id":"bb08709f-17b0-4bcb-b189-f2393e9c9de0","resolution":{"observed_at":"2026-08-14T04:31:15.591259Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:15.598713Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.598713Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:917e40705c479f54f0534d88f653d33b735451830205098cdf5109b64b95859b","observation_id":"57fcd832-8b10-4b1d-b7fe-0479db593a95","resolution":{"observed_at":"2026-08-14T04:31:15.598713Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.00034","last_updated":"2025-02-26T14:07:05Z","snapshot_observed_at":"2026-08-16T13:49:15.535352Z","submitted_at":"2024-05-26T21:39:53Z","title":"Adaptive Activation Steering: A Tuning-Free LLM Truthfulness Improvement Method for Diverse Hallucinations Categories","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.00034","snapshot_observed_at":"2026-08-14T04:31:15.602397Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.602397Z"},"links":{"cited_paper":"/paper/2406.00034","citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:dca2f7ef7c0f70dfe936f11221cb3a1cfcd6772945b9832b161518cbe8fe3df7","observation_id":"d02a6e65-a461-4bfe-af80-1f35aef22f1c","resolution":{"observed_at":"2026-08-14T04:31:15.602397Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2507.13255","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:15.833960Z","title":null,"venue":null,"work_id":"bd401e79-e16f-4f05-9ab5-3e796881a2cc","year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.606282Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:7a4daf7523faa60f8dce32134f5c2e24040ce524f740e50f9fa48a4d6824ee0b","observation_id":"5df269de-b46d-4893-a927-da2b3773583f","resolution":{"observed_at":"2026-08-14T04:31:15.842404Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.05176","last_updated":"2025-06-11T02:54:49Z","snapshot_observed_at":"2026-08-17T07:57:12.232097Z","submitted_at":"2025-06-05T15:49:48Z","title":"Qwen3 Embedding: Advancing Text Embedding and Reranking Through Foundation Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.05176","snapshot_observed_at":"2026-08-14T04:31:15.610027Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.610027Z"},"links":{"cited_paper":"/paper/2506.05176","citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:cb44472e44a706971adc3c6bc92e2c824423cb6e380f5d4b160358fe1642e673","observation_id":"a991e96a-601b-4ef2-be7d-4a89f8cab3ac","resolution":{"observed_at":"2026-08-14T04:31:15.610027Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.01405","last_updated":"2025-03-03T06:14:14Z","snapshot_observed_at":"2026-07-06T16:26:38.284922Z","submitted_at":"2023-10-02T17:59:07Z","title":"Representation Engineering: A Top-Down Approach to AI Transparency","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.01405","snapshot_observed_at":"2026-08-14T04:31:15.613830Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.613830Z"},"links":{"cited_paper":"/paper/2310.01405","citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:e1d03118c58691a341142adb053910867482d1007a0bc5430bd80ead37e25b65","observation_id":"1a014a55-c730-45d6-9d2c-f6b834ef89f7","resolution":{"observed_at":"2026-08-14T04:31:15.613830Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.663304Z","title":null,"venue":null,"work_id":"0ed787ba-35cb-44e2-9c5c-16779710a6e9","year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.617597Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:925a599505a8467372b2089d23cfca9dd08c5969fc4deeb06030bc0e14ebcc6a","observation_id":"f8cb505d-690d-4398-8806-f38492f866d7","resolution":{"observed_at":"2026-08-14T04:31:16.667463Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-08-13T17:20:44.002518Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-14T04:31:15.621185Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.621185Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:4a144b7c1d042e9fcc0dd627068e2db917411dcdd805a8f7e9a6266dac8dc4bb","observation_id":"c995ecc9-cd80-4072-86f1-ac98985f1aff","resolution":{"observed_at":"2026-08-14T04:31:15.621185Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.04261","last_updated":"2024-12-05T15:41:06Z","snapshot_observed_at":"2026-08-13T14:49:22.470549Z","submitted_at":"2024-12-05T15:41:06Z","title":"Aya Expanse: Combining Research Breakthroughs for a New Multilingual Frontier","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.04261","snapshot_observed_at":"2026-08-14T04:31:15.624835Z","title":"arXiv preprint arXiv:2412.04261 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.624835Z"},"links":{"cited_paper":"/paper/2412.04261","citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:360ceb2c61cbd6660b27253305a6c397c7af2c365807f18a6370e19dc5a3e479","observation_id":"f227f6ae-ac66-452b-832b-5783aadc042f","resolution":{"observed_at":"2026-08-14T04:31:15.624835Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.650680Z","title":"2025 , note=","venue":null,"work_id":"0895ecda-c25b-4e7d-a2de-9eaad0b5e693","year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.628251Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:41f7a5acd7917fccc0ad6c36c35a5f855f430d84205ab434fa71e67fe09e77bc","observation_id":"bbd5c563-8c5f-490b-9758-aca758764c1f","resolution":{"observed_at":"2026-08-14T04:31:16.655315Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.10248","last_updated":"2024-10-10T13:20:13Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-08-20T12:21:05Z","title":"Steering Language Models With Activation Engineering","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.10248","snapshot_observed_at":"2026-08-14T04:31:15.632218Z","title":"arXiv preprint arXiv:2308.10248 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.632218Z"},"links":{"cited_paper":"/paper/2308.10248","citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:58703650e8ee39e7c971792edbe7a4acadb8be7f71fc4f3724c9f604b555e86f","observation_id":"9b892a46-342f-4839-a239-05d59577845c","resolution":{"observed_at":"2026-08-14T04:31:15.632218Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.638711Z","title":"Steering","venue":null,"work_id":"6c3d78f9-6e44-4d0f-818d-c5ac9275c458","year":2024},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.635612Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:4477dd64fb5643e8a2c871f1806cb6ebd5da407f19664b8391a7a4c5f76304bc","observation_id":"d172221a-bd6b-48cc-8822-9a735c17bd26","resolution":{"observed_at":"2026-08-14T04:31:16.642641Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.625313Z","title":"Findings of the Association for Computational Linguistics: ACL 2023 , pages=","venue":null,"work_id":"96b3756f-52b3-4ef3-85b9-f59ad9998099","year":2023},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.639515Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:9aab5241679e9122e012802c6a9f2fdb9cf3095b32b4399ba05adacd517be3f9","observation_id":"762b0273-6442-4435-ac04-024b4f1b6312","resolution":{"observed_at":"2026-08-14T04:31:16.629555Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.611443Z","title":"Where to Steer: Input-Dependent Layer Selection for Steering Improves","venue":null,"work_id":"304f71cd-b73e-4f51-9bd6-8aa25a87d70b","year":2026},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.642819Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:e2de99c9555035f8a4f626b528cf1b0394fbf5c1f032f01488dc35ecad33167a","observation_id":"3391caa3-6c31-438a-b9e2-97e0c5b58515","resolution":{"observed_at":"2026-08-14T04:31:16.616274Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.597628Z","title":"Proceedings of the International Conference on Learning Representations (ICLR) , year=","venue":null,"work_id":"b07d22c5-705e-4b5e-b019-7df84a340adf","year":null},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.646245Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:fa07bbf6d08685f3cd13acdb3b89f08296e545789f7fc694dc228a622d225cf9","observation_id":"850a3bf1-f0cb-4fa3-9c35-d62e441b282f","resolution":{"observed_at":"2026-08-14T04:31:16.602569Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.584623Z","title":"Activation-Space Personality Steering: Hybrid Layer Selection for Stable Trait Control in","venue":null,"work_id":"2d1d7e02-b52c-48dc-8874-48405ef09851","year":2026},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.649885Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:d04e37372d2c33b59d3a8082b264cccb36c402c7e987a16e563b6881133b7c45","observation_id":"6bedf5b3-07b2-4e2b-a16b-6aa05feac8c6","resolution":{"observed_at":"2026-08-14T04:31:16.589301Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.570140Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":"9bc911a4-5d0b-4ebc-a2bd-a351e7ab0760","year":2023},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.653495Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:5551b9bb945712bdaaa1bf70984d925a392740c7816b17cc797cbbb8f0b9e301","observation_id":"ac5d36c7-efac-4250-ac28-a3bdceffd84a","resolution":{"observed_at":"2026-08-14T04:31:16.575616Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.555773Z","title":"Proceedings of the 2025 Conference on Empirical Methods in Natural Language Processing (EMNLP) , pages=","venue":null,"work_id":"0b155505-b8b8-4dfa-9db7-b1ff1d2073a9","year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.656791Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:be3ee9681225e6c289172f90fb3569d4e0ac3cedab1922c5fe4c9a86a4998ce7","observation_id":"ee06a78d-ba64-4c59-8891-59c70233cc60","resolution":{"observed_at":"2026-08-14T04:31:16.560403Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.542209Z","title":"Learning to Steer: Input-dependent Steering for Multimodal","venue":null,"work_id":"675226b5-3ae3-44a7-b9a3-e6a9ed2b5623","year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.660098Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:b1b86e0b1b868d546446aee21041b5c9e4d28ae7272edc60c7ee86cce850ab96","observation_id":"6b990954-a555-4861-8292-ea7fdaec0db9","resolution":{"observed_at":"2026-08-14T04:31:16.546648Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.528043Z","title":"To Steer or Not to Steer?","venue":null,"work_id":"b6c936cf-4324-45c1-b48e-c7a2f8bc6c16","year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.663437Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:acd407ecfba24506b893c582b3a7729d755be4eb6d7c00062917e54d854c3363","observation_id":"2b95bd0e-f3a7-4ff6-bee1-c905c1386bd9","resolution":{"observed_at":"2026-08-14T04:31:16.532856Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.514310Z","title":"Proceedings of the International Conference on Learning Representations (ICLR) , year=","venue":null,"work_id":"a662f400-e2cd-439f-8de5-204973119831","year":null},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.666880Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:d573bd5c0f98f89dc45ba97231217194ca23b43ee77f2e00b83965230c46e92d","observation_id":"d6da67ac-d6cf-4ef7-888c-20bfa227ec29","resolution":{"observed_at":"2026-08-14T04:31:16.519003Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.500101Z","title":"Proceedings of the 63rd Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers) , pages=","venue":null,"work_id":"599fb0af-5b44-4b47-9290-f05b2d1d2233","year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.670647Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:677de776b698bb5cf3147415e5de800bf086c38369eca20a7ce281b8d2247664","observation_id":"8945d3c9-1a97-4389-b935-05593fc8b4f4","resolution":{"observed_at":"2026-08-14T04:31:16.505404Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:15.673956Z","title":"Mechanistic Interpretability Workshop at NeurIPS 2025 , year=","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.673956Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:1218c310c0206df9453c7fb19e054e823dc2bdfedd63d53a36c610036b0231b8","observation_id":"5e4d165f-a597-4f70-849f-16dfcce23033","resolution":{"observed_at":"2026-08-14T04:31:15.673956Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.478104Z","title":"Zico and Hendrycks, Dan , journal=","venue":null,"work_id":"7649f36c-d8dd-4397-8718-b4f50d064b1e","year":null},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.677409Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:ea7a51ec5f9b9d25b065e13b1098a5b3cca9e39cf309499027c162ce3d55595e","observation_id":"ca238f9b-27db-42de-9a01-b89ba7ffaa5b","resolution":{"observed_at":"2026-08-14T04:31:16.482887Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.464181Z","title":"Advances in Neural Information Processing Systems (NeurIPS) , year=","venue":null,"work_id":"c870dfce-09d9-4328-806c-9b3b58570e9f","year":null},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.680932Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:300caa6fe73616669a240eac0f981789941dfa0284932ee249d6843272c14454","observation_id":"677a6eb4-e898-4fb9-905c-d0efc6990435","resolution":{"observed_at":"2026-08-14T04:31:16.468379Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.450818Z","title":"Advances in Neural Information Processing Systems (NeurIPS) , year=","venue":null,"work_id":"f1490557-751f-4d91-bb14-a1f8d7c2fd0a","year":null},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.684147Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:9b5deba663d8b5f97ab4f7657295dc3a8030a0ec1307ca9f5a6cc523c438a73b","observation_id":"dd200205-079e-4e54-b055-ae15a8d42e89","resolution":{"observed_at":"2026-08-14T04:31:16.455387Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.434816Z","title":"Findings of the Association for Computational Linguistics: EMNLP 2025 , pages=","venue":null,"work_id":"1493e758-14a5-4756-bc94-0bdf8f6cb847","year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.687645Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:c42e060f644fa2dd93872e13950da2524fdd97b77e0d0ad879870aca56bd8b43","observation_id":"4ca6793e-271a-4b47-96d3-8c629ad39d32","resolution":{"observed_at":"2026-08-14T04:31:16.440351Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.421394Z","title":"ICLR Workshop on Building Trust in","venue":null,"work_id":"3ca6a36e-e1dd-4f6c-b484-8da5a11077a3","year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.691095Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:00cb081837898a6d18700f495a4aacb633f1cfadb6ea195463433e5ffda813bf","observation_id":"1dfa1963-31f0-46d9-aa64-f5a91b5e9bbe","resolution":{"observed_at":"2026-08-14T04:31:16.425675Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.406837Z","title":"IEEE Access , volume=","venue":null,"work_id":"7330688f-743e-4603-bbe1-f11375fa43dd","year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.694700Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:6289bbbfcaed3032aab3e822b3606d0f1499912d62ea688ecaeb02265084c00f","observation_id":"1c938743-b8d1-49c8-9068-0432889a97cd","resolution":{"observed_at":"2026-08-14T04:31:16.411323Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.391547Z","title":"Proceedings of the 7th BlackboxNLP Workshop: Analyzing and Interpreting Neural Networks for NLP , pages=","venue":null,"work_id":"9ae3d92b-e2cd-4aeb-a437-6991f8124a41","year":2024},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.698273Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:d584956e6bb815de2808dd9347d2a5dbd463bf9b49d751810af6c0a50fb87035","observation_id":"e58d983d-b676-423d-898e-decb9e19ac05","resolution":{"observed_at":"2026-08-14T04:31:16.395852Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.376843Z","title":"Adaptive Activation Steering: A Tuning-Free","venue":null,"work_id":"577e0a36-db41-41da-9a13-3ca24e7ad3f3","year":2025},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.701771Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:c1dc5ecb398a6beab483adf00bb73404e1d7671d3a58ed77fa3107e8e48d5257","observation_id":"81471620-200f-465f-9a28-4183555c391d","resolution":{"observed_at":"2026-08-14T04:31:16.382008Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.363901Z","title":"Contributions to the Theory of Games II , editor=","venue":null,"work_id":"6b620462-1542-40e5-9229-0e40baaca1ca","year":1953},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.705247Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:28a090c34a35becc17f1614f889aff3b95f6259656d45c836ba4857dc7b09d6b","observation_id":"b3d40be1-1e69-4b3d-a63e-f98957a282d2","resolution":{"observed_at":"2026-08-14T04:31:16.368045Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:15.708930Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.708930Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:4e6c09bae81a233ed187bcdee2017445539d1ec397b66fddda24f1930bb1a3c9","observation_id":"45f33878-6155-469d-9625-5934a97c2a55","resolution":{"observed_at":"2026-08-14T04:31:15.708930Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-14T04:31:16.341286Z","title":"Cumulated Gain-Based Evaluation of","venue":null,"work_id":"0457dbba-8804-4e3b-a9f3-5137a46a5be9","year":null},"citing_paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-14T04:31:15.712697Z"},"links":{"citing_paper":"/paper/2608.08829"},"observation_digest":"sha256:8b969af79fc8391ea31c35d3bc2b47dceae4e4ef479a4fcaa14de55495a05dcc","observation_id":"d9ff5e7a-57e5-491e-ad00-b7b72521dbe5","resolution":{"observed_at":"2026-08-14T04:31:16.345395Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2608.08829","last_updated":"2026-08-09T17:23:00Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-14T05:40:32.837436Z","submitted_at":"2026-08-09T17:23:00Z","title":"Deployable Per-Instance Multi-Layer Activation Steering for Large Language Models"},"reference_resolution":{"displayed":51,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":24,"verified_exact":4,"verified_fuzzy":23},"total_outbound_references":51},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"thesis":"As of 17 August 2026, this Paper Citation Record lists 51 of 51 outbound references and 0 inbound Pith citation observations for arXiv:2608.08829."}