{"as_of":"2026-08-06T05:24:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:46f5747aa3821bdb99a9bca92a779f70a9a8db649975242e64e82fa3c4a3a191","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":16,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":16,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-05T06:32:48.257954+00:00","state":"measured"},{"denominator":16,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":16,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T22:51:28.064772Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2503.00177","last_updated":"2025-02-28T20:43:45Z","snapshot_observed_at":"2026-07-06T20:44:41.083168Z","submitted_at":"2025-02-28T20:43:45Z","title":"Steering Large Language Model Activations in Sparse Spaces","version":1},"cited_work":{"arxiv_id":"2503.00177","doi":"10.48550/arxiv.2503.00177","metadata_source":"arxiv_reference","pith_arxiv_id":"2503.00177","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Steering large language model activations in sparse spaces.arXiv preprint arXiv:2503.00177","venue":"ArXiv.org","work_id":"e10f6b0d-fd6f-4116-a4a5-4f88c1e94d92","year":2014},"citing_paper":{"arxiv_id":"2506.01247","last_updated":"2026-05-26T21:34:46Z","snapshot_observed_at":"2026-07-06T21:34:43.519872Z","submitted_at":"2025-06-02T01:51:20Z","title":"Beyond Interpretability: When, Why, and How Sparse Autoencoders Enable Label-Free Visual Steering","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-05-19T11:35:01.836096Z"},"links":{"cited_paper":"/paper/2503.00177","citing_paper":"/paper/2506.01247"},"observation_digest":"sha256:2a2f83cbc80a08c95f60696344e2b180712f5309aa896f2c9b3aaa4f8e2c3e9e","observation_id":"a9a7c9c2-6377-43b2-ae6c-c27d0133efa0","resolution":{"observed_at":"2026-05-19T11:37:15.808881Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.00177","last_updated":"2025-02-28T20:43:45Z","snapshot_observed_at":"2026-07-06T20:44:41.083168Z","submitted_at":"2025-02-28T20:43:45Z","title":"Steering Large Language Model Activations in Sparse Spaces","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.00177","snapshot_observed_at":"2026-08-05T22:51:28.064772Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.06418","last_updated":"2025-08-08T16:05:27Z","snapshot_observed_at":"2026-08-06T02:11:05.197113Z","submitted_at":"2025-08-08T16:05:27Z","title":"Quantifying Conversation Drift in MCP via Latent Polytope","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-05T22:51:28.064772Z"},"links":{"cited_paper":"/paper/2503.00177","citing_paper":"/paper/2508.06418"},"observation_digest":"sha256:0b12892a4a565c10cdea24bf87cdd35737aa69f222559d7e940b034be09581ee","observation_id":"6ff86b81-6e43-4bf1-9a83-10780ba5bdf7","resolution":{"observed_at":"2026-08-05T22:51:28.064772Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.00177","last_updated":"2025-02-28T20:43:45Z","snapshot_observed_at":"2026-07-06T20:44:41.083168Z","submitted_at":"2025-02-28T20:43:45Z","title":"Steering Large Language Model Activations in Sparse Spaces","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.00177","snapshot_observed_at":"2026-08-05T10:55:31.191794Z","title":"Bayat, A","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.03518","last_updated":"2025-09-03T17:59:45Z","snapshot_observed_at":"2026-08-05T10:55:25.765282Z","submitted_at":"2025-09-03T17:59:45Z","title":"Can LLMs Lie? Investigation beyond Hallucination","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-05T10:55:31.191794Z"},"links":{"cited_paper":"/paper/2503.00177","citing_paper":"/paper/2509.03518"},"observation_digest":"sha256:828a42c6f66cbf94b616606cc26b8964590ae2e32341e0b1ee7e9abc447c237b","observation_id":"1c478744-4bce-4874-a330-e8481626d9f6","resolution":{"observed_at":"2026-08-05T10:55:31.191794Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.00177","last_updated":"2025-02-28T20:43:45Z","snapshot_observed_at":"2026-07-06T20:44:41.083168Z","submitted_at":"2025-02-28T20:43:45Z","title":"Steering Large Language Model Activations in Sparse Spaces","version":1},"cited_work":{"arxiv_id":"2503.00177","doi":"10.48550/arxiv.2503.00177","metadata_source":"arxiv_reference","pith_arxiv_id":"2503.00177","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Steering large language model activations in sparse spaces.arXiv preprint arXiv:2503.00177","venue":"ArXiv.org","work_id":"e10f6b0d-fd6f-4116-a4a5-4f88c1e94d92","year":2014},"citing_paper":{"arxiv_id":"2509.22739","last_updated":"2026-05-15T15:16:23Z","snapshot_observed_at":"2026-07-06T22:30:55.313733Z","submitted_at":"2025-09-25T23:25:47Z","title":"Painless Activation Steering: An Automated, Lightweight Approach for Post-Training Large Language Models","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-21T21:44:36.351517Z"},"links":{"cited_paper":"/paper/2503.00177","citing_paper":"/paper/2509.22739"},"observation_digest":"sha256:a58c48ce95e77eb096ccbe8cc09712a81c6f97312da33d03a4a0a00caac332b1","observation_id":"012dc5a2-03c8-46ba-a6f2-0072f61787d7","resolution":{"observed_at":"2026-05-21T21:45:40.572737Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.00177","last_updated":"2025-02-28T20:43:45Z","snapshot_observed_at":"2026-07-06T20:44:41.083168Z","submitted_at":"2025-02-28T20:43:45Z","title":"Steering Large Language Model Activations in Sparse Spaces","version":1},"cited_work":{"arxiv_id":"2503.00177","doi":"10.48550/arxiv.2503.00177","metadata_source":"arxiv_reference","pith_arxiv_id":"2503.00177","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Steering large language model activations in sparse spaces.arXiv preprint arXiv:2503.00177","venue":"ArXiv.org","work_id":"e10f6b0d-fd6f-4116-a4a5-4f88c1e94d92","year":2014},"citing_paper":{"arxiv_id":"2601.14004","last_updated":"2026-04-14T03:49:06Z","snapshot_observed_at":"2026-07-31T09:24:49.481740Z","submitted_at":"2026-01-20T14:23:23Z","title":"Locate, Steer, and Improve: A Practical Survey of Actionable Mechanistic Interpretability in Large Language Models","version":4},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-16T12:39:57.398423Z"},"links":{"cited_paper":"/paper/2503.00177","citing_paper":"/paper/2601.14004"},"observation_digest":"sha256:ed910982a6e5466dc9ffed59d0283a830af549151b4e7c64561ccd975237a7e9","observation_id":"f643abdf-37bf-4643-92a3-e221791edef7","resolution":{"observed_at":"2026-05-16T12:40:54.808987Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.00177","last_updated":"2025-02-28T20:43:45Z","snapshot_observed_at":"2026-07-06T20:44:41.083168Z","submitted_at":"2025-02-28T20:43:45Z","title":"Steering Large Language Model Activations in Sparse Spaces","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.00177","snapshot_observed_at":"2026-08-02T23:52:40.237099Z","title":"misaligned persona","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.12418","last_updated":"2026-06-27T01:07:47Z","snapshot_observed_at":"2026-08-03T15:50:26.557526Z","submitted_at":"2026-02-12T21:17:32Z","title":"Sparse Autoencoders are Capable LLM Jailbreak Mitigators","version":2},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-02T23:52:40.237099Z"},"links":{"cited_paper":"/paper/2503.00177","citing_paper":"/paper/2602.12418"},"observation_digest":"sha256:df0f9c990f56ebd7061b1a1bc44d4aa1a5e5154c9580ed11935626d0b6d3cdb5","observation_id":"1a677b74-7079-44f6-a909-3ad05bc53d35","resolution":{"observed_at":"2026-08-02T23:52:40.237099Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.00177","last_updated":"2025-02-28T20:43:45Z","snapshot_observed_at":"2026-07-06T20:44:41.083168Z","submitted_at":"2025-02-28T20:43:45Z","title":"Steering Large Language Model Activations in Sparse Spaces","version":1},"cited_work":{"arxiv_id":"2503.00177","doi":"10.48550/arxiv.2503.00177","metadata_source":"arxiv_reference","pith_arxiv_id":"2503.00177","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Steering large language model activations in sparse spaces.arXiv preprint arXiv:2503.00177","venue":"ArXiv.org","work_id":"e10f6b0d-fd6f-4116-a4a5-4f88c1e94d92","year":2014},"citing_paper":{"arxiv_id":"2604.15789","last_updated":"2026-04-17T07:50:42Z","snapshot_observed_at":"2026-07-06T23:03:16.345488Z","submitted_at":"2026-04-17T07:50:42Z","title":"A Systematic Study of Training-Free Methods for Trustworthy Large Language Models","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-10T08:19:42.671690Z"},"links":{"cited_paper":"/paper/2503.00177","citing_paper":"/paper/2604.15789"},"observation_digest":"sha256:b30d0bccfadda89f8325b08abfb3423aa9730f80a00f8f219ff8689da65d3b96","observation_id":"f56f0c2d-1217-4559-99f8-9b2a41b49403","resolution":{"observed_at":"2026-05-10T08:22:37.358896Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.00177","last_updated":"2025-02-28T20:43:45Z","snapshot_observed_at":"2026-07-06T20:44:41.083168Z","submitted_at":"2025-02-28T20:43:45Z","title":"Steering Large Language Model Activations in Sparse Spaces","version":1},"cited_work":{"arxiv_id":"2503.00177","doi":"10.48550/arxiv.2503.00177","metadata_source":"arxiv_reference","pith_arxiv_id":"2503.00177","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Steering large language model activations in sparse spaces.arXiv preprint arXiv:2503.00177","venue":"ArXiv.org","work_id":"e10f6b0d-fd6f-4116-a4a5-4f88c1e94d92","year":2014},"citing_paper":{"arxiv_id":"2605.08426","last_updated":"2026-06-02T11:37:31Z","snapshot_observed_at":"2026-07-06T23:20:38.385189Z","submitted_at":"2026-05-08T19:37:07Z","title":"Mechanism Design Is Not Enough: Prosocial Agents for Cooperative AI","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-12T00:58:59.587622Z"},"links":{"cited_paper":"/paper/2503.00177","citing_paper":"/paper/2605.08426"},"observation_digest":"sha256:aed0da041d20d42fc35bb6297b88efda360e79e52df29f7aefe22c0661e02977","observation_id":"8a0bf699-6824-4266-8ef0-b48d1788f6d5","resolution":{"observed_at":"2026-05-12T08:36:24.206288Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.00177","last_updated":"2025-02-28T20:43:45Z","snapshot_observed_at":"2026-07-06T20:44:41.083168Z","submitted_at":"2025-02-28T20:43:45Z","title":"Steering Large Language Model Activations in Sparse Spaces","version":1},"cited_work":{"arxiv_id":"2503.00177","doi":"10.48550/arxiv.2503.00177","metadata_source":"arxiv_reference","pith_arxiv_id":"2503.00177","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Steering large language model activations in sparse spaces.arXiv preprint arXiv:2503.00177","venue":"ArXiv.org","work_id":"e10f6b0d-fd6f-4116-a4a5-4f88c1e94d92","year":2014},"citing_paper":{"arxiv_id":"2605.08426","last_updated":"2026-06-02T11:37:31Z","snapshot_observed_at":"2026-07-06T23:20:38.385189Z","submitted_at":"2026-05-08T19:37:07Z","title":"Mechanism Design Is Not Enough: Prosocial Agents for Cooperative AI","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-06-30T23:18:57.280524Z"},"links":{"cited_paper":"/paper/2503.00177","citing_paper":"/paper/2605.08426"},"observation_digest":"sha256:bcca90fc9bf75a0c6474086ac213286afda268891c7608d65dce43a67683e557","observation_id":"dbd43405-b993-4fcd-a586-968b2d8deedb","resolution":{"observed_at":"2026-07-01T13:25:45.211495Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.00177","last_updated":"2025-02-28T20:43:45Z","snapshot_observed_at":"2026-07-06T20:44:41.083168Z","submitted_at":"2025-02-28T20:43:45Z","title":"Steering Large Language Model Activations in Sparse Spaces","version":1},"cited_work":{"arxiv_id":"2503.00177","doi":"10.48550/arxiv.2503.00177","metadata_source":"arxiv_reference","pith_arxiv_id":"2503.00177","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Steering large language model activations in sparse spaces.arXiv preprint arXiv:2503.00177","venue":"ArXiv.org","work_id":"e10f6b0d-fd6f-4116-a4a5-4f88c1e94d92","year":2014},"citing_paper":{"arxiv_id":"2605.11887","last_updated":"2026-05-12T10:01:06Z","snapshot_observed_at":"2026-07-29T16:09:50.456033Z","submitted_at":"2026-05-12T10:01:06Z","title":"Qwen-Scope: Turning Sparse Features into Development Tools for Large Language Models","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-05-13T05:38:31.094834Z"},"links":{"cited_paper":"/paper/2503.00177","citing_paper":"/paper/2605.11887"},"observation_digest":"sha256:5eece50985a71f6d4be19ba47e686b2691fb9fc95e4539cd7600aee487fd6d62","observation_id":"faa2a36e-cccb-4985-b630-8bb727c0e3b9","resolution":{"observed_at":"2026-05-13T05:42:21.056231Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.00177","last_updated":"2025-02-28T20:43:45Z","snapshot_observed_at":"2026-07-06T20:44:41.083168Z","submitted_at":"2025-02-28T20:43:45Z","title":"Steering Large Language Model Activations in Sparse Spaces","version":1},"cited_work":{"arxiv_id":"2503.00177","doi":"10.48550/arxiv.2503.00177","metadata_source":"arxiv_reference","pith_arxiv_id":"2503.00177","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Steering large language model activations in sparse spaces.arXiv preprint arXiv:2503.00177","venue":"ArXiv.org","work_id":"e10f6b0d-fd6f-4116-a4a5-4f88c1e94d92","year":2014},"citing_paper":{"arxiv_id":"2605.23040","last_updated":"2026-05-21T21:13:14Z","snapshot_observed_at":"2026-08-02T04:36:31.374601Z","submitted_at":"2026-05-21T21:13:14Z","title":"Steered Generation via Gradient-Based Optimization on Sparse Query Features","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-25T05:31:29.510639Z"},"links":{"cited_paper":"/paper/2503.00177","citing_paper":"/paper/2605.23040"},"observation_digest":"sha256:d9ff3204ad5a3d33238fc3d35fa6d28d78c607d3064d20d6d778fb79df232ac9","observation_id":"e2d545c9-8b9c-48e3-bff0-18fe56995128","resolution":{"observed_at":"2026-05-25T05:36:40.256994Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.00177","last_updated":"2025-02-28T20:43:45Z","snapshot_observed_at":"2026-07-06T20:44:41.083168Z","submitted_at":"2025-02-28T20:43:45Z","title":"Steering Large Language Model Activations in Sparse Spaces","version":1},"cited_work":{"arxiv_id":"2503.00177","doi":"10.48550/arxiv.2503.00177","metadata_source":"arxiv_reference","pith_arxiv_id":"2503.00177","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Steering large language model activations in sparse spaces.arXiv preprint arXiv:2503.00177","venue":"ArXiv.org","work_id":"e10f6b0d-fd6f-4116-a4a5-4f88c1e94d92","year":2014},"citing_paper":{"arxiv_id":"2605.28664","last_updated":"2026-05-27T15:59:45Z","snapshot_observed_at":"2026-07-06T23:38:14.056361Z","submitted_at":"2026-05-27T15:59:45Z","title":"Activation Steering for Synthetic Data Generation: The Role of Diversity in Downstream Safety Detection","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-29T13:53:27.306664Z"},"links":{"cited_paper":"/paper/2503.00177","citing_paper":"/paper/2605.28664"},"observation_digest":"sha256:b0647d55ff4eef4d8252614e5e13246d0949c66e08311ad21037d3c739a6b4d1","observation_id":"095ecfb9-1be0-4759-ace8-4ee5190672b3","resolution":{"observed_at":"2026-06-29T14:03:29.936584Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.00177","last_updated":"2025-02-28T20:43:45Z","snapshot_observed_at":"2026-07-06T20:44:41.083168Z","submitted_at":"2025-02-28T20:43:45Z","title":"Steering Large Language Model Activations in Sparse Spaces","version":1},"cited_work":{"arxiv_id":"2503.00177","doi":"10.48550/arxiv.2503.00177","metadata_source":"arxiv_reference","pith_arxiv_id":"2503.00177","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Steering large language model activations in sparse spaces.arXiv preprint arXiv:2503.00177","venue":"ArXiv.org","work_id":"e10f6b0d-fd6f-4116-a4a5-4f88c1e94d92","year":2014},"citing_paper":{"arxiv_id":"2606.00726","last_updated":"2026-07-10T01:15:34Z","snapshot_observed_at":"2026-08-02T05:37:27.284595Z","submitted_at":"2026-05-30T13:38:06Z","title":"Latent Reward Steering: An Adaptive Inference-Time Framework that Implicitly Promotes Cognitive Behaviors in Reasoning LLMs","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-06-28T18:49:56.917505Z"},"links":{"cited_paper":"/paper/2503.00177","citing_paper":"/paper/2606.00726"},"observation_digest":"sha256:22ea021c03689e00abd8d03e60042bacd6a92e5b4d46567e4dedc925b85b2977","observation_id":"e73107f4-1d3e-4921-89f9-3be0aa915c1f","resolution":{"observed_at":"2026-06-28T19:52:35.549158Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.00177","last_updated":"2025-02-28T20:43:45Z","snapshot_observed_at":"2026-07-06T20:44:41.083168Z","submitted_at":"2025-02-28T20:43:45Z","title":"Steering Large Language Model Activations in Sparse Spaces","version":1},"cited_work":{"arxiv_id":"2503.00177","doi":"10.48550/arxiv.2503.00177","metadata_source":"arxiv_reference","pith_arxiv_id":"2503.00177","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Steering large language model activations in sparse spaces.arXiv preprint arXiv:2503.00177","venue":"ArXiv.org","work_id":"e10f6b0d-fd6f-4116-a4a5-4f88c1e94d92","year":2014},"citing_paper":{"arxiv_id":"2606.03002","last_updated":"2026-06-04T19:39:22Z","snapshot_observed_at":"2026-08-04T10:42:06.850785Z","submitted_at":"2026-06-02T01:17:05Z","title":"Perplexity Can Miss SAE Feature Damage Under Quantization","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-06-28T11:41:18.460538Z"},"links":{"cited_paper":"/paper/2503.00177","citing_paper":"/paper/2606.03002"},"observation_digest":"sha256:85d8caf61941f889fac65645df53b094f0e78707de8af9d4b1081cdfb8707af2","observation_id":"17e05bde-bc72-4c61-a1c7-78c18fbbd42c","resolution":{"observed_at":"2026-07-02T01:36:25.800030Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.00177","last_updated":"2025-02-28T20:43:45Z","snapshot_observed_at":"2026-07-06T20:44:41.083168Z","submitted_at":"2025-02-28T20:43:45Z","title":"Steering Large Language Model Activations in Sparse Spaces","version":1},"cited_work":{"arxiv_id":"2503.00177","doi":"10.48550/arxiv.2503.00177","metadata_source":"arxiv_reference","pith_arxiv_id":"2503.00177","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Steering large language model activations in sparse spaces.arXiv preprint arXiv:2503.00177","venue":"ArXiv.org","work_id":"e10f6b0d-fd6f-4116-a4a5-4f88c1e94d92","year":2014},"citing_paper":{"arxiv_id":"2606.06315","last_updated":"2026-06-04T15:54:34Z","snapshot_observed_at":"2026-07-06T23:46:10.612537Z","submitted_at":"2026-06-04T15:54:34Z","title":"LLM Self-Recognition: Steering and Retrieving Activation Signatures","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-06-28T01:41:03.518190Z"},"links":{"cited_paper":"/paper/2503.00177","citing_paper":"/paper/2606.06315"},"observation_digest":"sha256:fe7c77f3460dd17a0161905927da559569ce25686134ed9880f6a6cb838842fe","observation_id":"5de0e7ba-4a1d-41f6-9e97-026363e213ad","resolution":{"observed_at":"2026-06-28T01:41:29.197320Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.00177","last_updated":"2025-02-28T20:43:45Z","snapshot_observed_at":"2026-07-06T20:44:41.083168Z","submitted_at":"2025-02-28T20:43:45Z","title":"Steering Large Language Model Activations in Sparse Spaces","version":1},"cited_work":{"arxiv_id":"2503.00177","doi":"10.48550/arxiv.2503.00177","metadata_source":"arxiv_reference","pith_arxiv_id":"2503.00177","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Steering large language model activations in sparse spaces.arXiv preprint arXiv:2503.00177","venue":"ArXiv.org","work_id":"e10f6b0d-fd6f-4116-a4a5-4f88c1e94d92","year":2014},"citing_paper":{"arxiv_id":"2606.23670","last_updated":"2026-06-22T17:56:25Z","snapshot_observed_at":"2026-07-06T23:58:20.975630Z","submitted_at":"2026-06-22T17:56:25Z","title":"Tapered Language Models","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-06-26T09:11:20.341634Z"},"links":{"cited_paper":"/paper/2503.00177","citing_paper":"/paper/2606.23670"},"observation_digest":"sha256:9c7f42925be096aa95f462fcc4379236a116d04a5139eea153c1db32c3edaa38","observation_id":"046e59be-ba9f-46fb-b961-813ea7e5fc52","resolution":{"observed_at":"2026-07-04T10:09:44.002989Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2503.00177/citation-record","integrity":"/paper/2503.00177/integrity","json":"/paper/2503.00177/citation-record.json","paper":"/paper/2503.00177"},"outbound":[],"paper":{"arxiv_id":"2503.00177","last_updated":"2025-02-28T20:43:45Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-07-06T20:44:41.083168Z","submitted_at":"2025-02-28T20:43:45Z","title":"Steering Large Language Model Activations in Sparse Spaces"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"thesis":"As of 6 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 16 inbound Pith citation observations for arXiv:2503.00177."}