{"as_of":"2026-08-07T10:35:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:e80acf6ec036bb35b3a44ee0c29534c85a3a17310c6ef3c401c021756f5f78fb","coverage":[{"denominator":13,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":13,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T14:27:25.454156Z","state":"measured"},{"denominator":13,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":13,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2508.21324/citation-record","integrity":"/paper/2508.21324/integrity","json":"/paper/2508.21324/citation-record.json","paper":"/paper/2508.21324"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2503.09532","last_updated":"2025-06-04T13:27:28Z","snapshot_observed_at":"2026-08-03T07:32:09.915873Z","submitted_at":"2025-03-12T16:49:02Z","title":"SAEBench: A Comprehensive Benchmark for Sparse Autoencoders in Language Model Interpretability","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.09532","snapshot_observed_at":"2026-08-05T14:27:25.439860Z","title":"Saebench: A comprehensive benchmark for sparse autoencoders in language model interpretability","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.21324","last_updated":"2025-08-29T04:42:17Z","snapshot_observed_at":"2026-08-05T14:27:24.567504Z","submitted_at":"2025-08-29T04:42:17Z","title":"Distribution-Aware Feature Selection for SAEs","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-05T14:27:25.439860Z"},"links":{"cited_paper":"/paper/2503.09532","citing_paper":"/paper/2508.21324"},"observation_digest":"sha256:4050039e8d30227ff75d3ebae7c6041dae2fd591fe08cba93a84ea1f479eb5b1","observation_id":"5c41650a-07f1-43c6-bec8-c9c5004543f8","resolution":{"observed_at":"2026-08-05T14:27:25.439860Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.16615","last_updated":"2025-01-29T19:18:45Z","snapshot_observed_at":"2026-07-06T20:27:09.296190Z","submitted_at":"2025-01-28T01:24:16Z","title":"Sparse Autoencoders Trained on the Same Data Learn Different Features","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.16615","snapshot_observed_at":"2026-08-05T14:27:25.446156Z","title":"Sparse autoencoders trained on the same data learn different features","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.21324","last_updated":"2025-08-29T04:42:17Z","snapshot_observed_at":"2026-08-05T14:27:24.567504Z","submitted_at":"2025-08-29T04:42:17Z","title":"Distribution-Aware Feature Selection for SAEs","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-05T14:27:25.446156Z"},"links":{"cited_paper":"/paper/2501.16615","citing_paper":"/paper/2508.21324"},"observation_digest":"sha256:e41594a6cf251365ffdd1d8d07f317b6779531287e93cfb3fbf9c77722094a6b","observation_id":"96cfe9ce-a37d-41f3-9def-752e842cb6ed","resolution":{"observed_at":"2026-08-05T14:27:25.446156Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:27:25.595762Z","title":"activation lottery","venue":null,"work_id":"0801ddd3-af2b-4a74-8299-9a4a1cafb467","year":2021},"citing_paper":{"arxiv_id":"2508.21324","last_updated":"2025-08-29T04:42:17Z","snapshot_observed_at":"2026-08-05T14:27:24.567504Z","submitted_at":"2025-08-29T04:42:17Z","title":"Distribution-Aware Feature Selection for SAEs","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-05T14:27:25.454156Z"},"links":{"citing_paper":"/paper/2508.21324"},"observation_digest":"sha256:151141717c174b589b79e443fc5de05444f63ac874a0c543ddaf4da75396a5f6","observation_id":"c04e6590-8c83-4c82-909d-740fa0f5ac6a","resolution":{"observed_at":"2026-08-05T14:27:25.599812Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:27:25.604973Z","title":"High frequency latents are features, not bugs","venue":null,"work_id":"162d3dcf-0e05-4c7e-a7d5-85c8a158c46e","year":2025},"citing_paper":{"arxiv_id":"2508.21324","last_updated":"2025-08-29T04:42:17Z","snapshot_observed_at":"2026-08-05T14:27:24.567504Z","submitted_at":"2025-08-29T04:42:17Z","title":"Distribution-Aware Feature Selection for SAEs","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-05T14:27:25.451622Z"},"links":{"citing_paper":"/paper/2508.21324"},"observation_digest":"sha256:2e803afd31e6426e1021a5c694573114fe634b2c8b383a819a255c6cb765fa72","observation_id":"f4ba490c-26f2-40af-afd2-8114fd7b3b5f","resolution":{"observed_at":"2026-08-05T14:27:25.609772Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2101.00027","last_updated":"2020-12-31T19:00:10Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-12-31T19:00:10Z","title":"The Pile: An 800GB Dataset of Diverse Text for Language Modeling","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2101.00027","snapshot_observed_at":"2026-08-05T14:27:25.430598Z","title":"The pile: An 800gb dataset of diverse text for language modeling","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.21324","last_updated":"2025-08-29T04:42:17Z","snapshot_observed_at":"2026-08-05T14:27:24.567504Z","submitted_at":"2025-08-29T04:42:17Z","title":"Distribution-Aware Feature Selection for SAEs","version":1},"reference_index":2004,"source":"pdf_text","source_observed_at":"2026-08-05T14:27:25.430598Z"},"links":{"cited_paper":"/paper/2101.00027","citing_paper":"/paper/2508.21324"},"observation_digest":"sha256:c9ecbd2eb6ca45d893c32a579e5cb5754b182a0ad0a6b9c410982505bb318c22","observation_id":"de1f2a05-692c-493d-acd3-caa96f5d87da","resolution":{"observed_at":"2026-08-05T14:27:25.430598Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.14926","last_updated":"2025-02-07T19:22:32Z","snapshot_observed_at":"2026-08-02T04:19:16.503331Z","submitted_at":"2025-01-24T21:31:12Z","title":"Interpretability in Parameter Space: Minimizing Mechanistic Description Length with Attribution-based Parameter Decomposition","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.14926","snapshot_observed_at":"2026-08-05T14:27:25.414379Z","title":"Interpretability in parameter space: Minimizing mechanistic description length with attribution-based parameter decomposition","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.21324","last_updated":"2025-08-29T04:42:17Z","snapshot_observed_at":"2026-08-05T14:27:24.567504Z","submitted_at":"2025-08-29T04:42:17Z","title":"Distribution-Aware Feature Selection for SAEs","version":1},"reference_index":2009,"source":"pdf_text","source_observed_at":"2026-08-05T14:27:25.414379Z"},"links":{"cited_paper":"/paper/2501.14926","citing_paper":"/paper/2508.21324"},"observation_digest":"sha256:c000ed47cde4643cb6f109bc837b256d0dc23a86dda40d0800317b7138f1cc67","observation_id":"faf94b8e-5da5-4fb6-b848-1c0faa36b836","resolution":{"observed_at":"2026-08-05T14:27:25.414379Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:27:25.614740Z","title":"Online row sampling","venue":null,"work_id":"b4eb1181-dc86-4da1-880c-9587d3cdafa4","year":2016},"citing_paper":{"arxiv_id":"2508.21324","last_updated":"2025-08-29T04:42:17Z","snapshot_observed_at":"2026-08-05T14:27:24.567504Z","submitted_at":"2025-08-29T04:42:17Z","title":"Distribution-Aware Feature Selection for SAEs","version":1},"reference_index":2015,"source":"pdf_text","source_observed_at":"2026-08-05T14:27:25.424101Z"},"links":{"citing_paper":"/paper/2508.21324"},"observation_digest":"sha256:aa6a5604915535304ea5ce0d6e74a076b842672382bc6b1d08e99075d3294c70","observation_id":"d186e98d-a9c8-4a13-a445-c6d7c0511051","resolution":{"observed_at":"2026-08-05T14:27:25.617473Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.01610","last_updated":"2023-06-02T21:52:17Z","snapshot_observed_at":"2026-08-02T17:50:59.894886Z","submitted_at":"2023-05-02T17:13:55Z","title":"Finding Neurons in a Haystack: Case Studies with Sparse Probing","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.01610","snapshot_observed_at":"2026-08-05T14:27:25.436740Z","title":"Finding neurons in a haystack: Case studies with sparse probing","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.21324","last_updated":"2025-08-29T04:42:17Z","snapshot_observed_at":"2026-08-05T14:27:24.567504Z","submitted_at":"2025-08-29T04:42:17Z","title":"Distribution-Aware Feature Selection for SAEs","version":1},"reference_index":2016,"source":"pdf_text","source_observed_at":"2026-08-05T14:27:25.436740Z"},"links":{"cited_paper":"/paper/2305.01610","citing_paper":"/paper/2508.21324"},"observation_digest":"sha256:9bdcc923d8a785d71a26f92feb7f334b7433b2f2e9f8b731b32cc7cbf997cfac","observation_id":"c5b71cd4-73d5-4923-b094-3d5c2d3cb41a","resolution":{"observed_at":"2026-08-05T14:27:25.436740Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.08600","last_updated":"2023-10-04T13:17:38Z","snapshot_observed_at":"2026-07-06T16:19:05.495349Z","submitted_at":"2023-09-15T17:56:55Z","title":"Sparse Autoencoders Find Highly Interpretable Features in Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.08600","snapshot_observed_at":"2026-08-05T14:27:25.427171Z","title":"Sparse autoen- coders find highly interpretable features in language models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.21324","last_updated":"2025-08-29T04:42:17Z","snapshot_observed_at":"2026-08-05T14:27:24.567504Z","submitted_at":"2025-08-29T04:42:17Z","title":"Distribution-Aware Feature Selection for SAEs","version":1},"reference_index":2017,"source":"pdf_text","source_observed_at":"2026-08-05T14:27:25.427171Z"},"links":{"cited_paper":"/paper/2309.08600","citing_paper":"/paper/2508.21324"},"observation_digest":"sha256:9d8adb0259d0fad5218ccc00df6a875e0b020a0ff1ea0e408a9122c514b8df0e","observation_id":"ff08ff48-df8c-4a13-b271-604b726fe0b5","resolution":{"observed_at":"2026-08-05T14:27:25.427171Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.16014","last_updated":"2024-04-30T17:54:04Z","snapshot_observed_at":"2026-08-02T10:59:21.418230Z","submitted_at":"2024-04-24T17:47:22Z","title":"Improving Dictionary Learning with Gated Sparse Autoencoders","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.16014","snapshot_observed_at":"2026-08-05T14:27:25.433612Z","title":"Improving dictionary learning with gated sparse autoencoders","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.21324","last_updated":"2025-08-29T04:42:17Z","snapshot_observed_at":"2026-08-05T14:27:24.567504Z","submitted_at":"2025-08-29T04:42:17Z","title":"Distribution-Aware Feature Selection for SAEs","version":1},"reference_index":2020,"source":"pdf_text","source_observed_at":"2026-08-05T14:27:25.433612Z"},"links":{"cited_paper":"/paper/2404.16014","citing_paper":"/paper/2508.21324"},"observation_digest":"sha256:d264232f3917e224a67e8b372e6bd332bb4677dde51864b0a4880f099f15e025","observation_id":"83ed672f-fb79-42bb-a6ad-a2f7a6671e5b","resolution":{"observed_at":"2026-08-05T14:27:25.433612Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:27:25.622515Z","title":"Bart Bussmann, Joseph Jermyn, and Nix Robertson","venue":null,"work_id":"d63e0497-380a-4088-b81d-f8d2801e30ec","year":2023},"citing_paper":{"arxiv_id":"2508.21324","last_updated":"2025-08-29T04:42:17Z","snapshot_observed_at":"2026-08-05T14:27:24.567504Z","submitted_at":"2025-08-29T04:42:17Z","title":"Distribution-Aware Feature Selection for SAEs","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-05T14:27:25.418043Z"},"links":{"citing_paper":"/paper/2508.21324"},"observation_digest":"sha256:30dfcf062ede74e50fb62a93a02e02e5da1feff96c44e2b607bf008630bdd89f","observation_id":"9715d8c0-cb74-4c08-8466-e1f0af7674b3","resolution":{"observed_at":"2026-08-05T14:27:25.625688Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:27:25.420785Z","title":"David Chanin, James Wilken-Smith, Tomáš Dulka, Hardik Bhatnagar, Satvik Golechha, and Joseph Bloom","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.21324","last_updated":"2025-08-29T04:42:17Z","snapshot_observed_at":"2026-08-05T14:27:24.567504Z","submitted_at":"2025-08-29T04:42:17Z","title":"Distribution-Aware Feature Selection for SAEs","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-05T14:27:25.420785Z"},"links":{"citing_paper":"/paper/2508.21324"},"observation_digest":"sha256:731948d63fc60ab5d9cd26605a1bcb41ee94b5781f421e93ac1c0a9bc715d125","observation_id":"8712a6ef-8836-44d3-9563-67d00474fb30","resolution":{"observed_at":"2026-08-05T14:27:25.420785Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.13928","last_updated":"2025-08-06T13:47:10Z","snapshot_observed_at":"2026-07-06T19:35:35.229351Z","submitted_at":"2024-10-17T17:56:01Z","title":"Automatically Interpreting Millions of Features in Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.13928","snapshot_observed_at":"2026-08-05T14:27:25.448907Z","title":"Automatically interpreting millions of features in large language models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.21324","last_updated":"2025-08-29T04:42:17Z","snapshot_observed_at":"2026-08-05T14:27:24.567504Z","submitted_at":"2025-08-29T04:42:17Z","title":"Distribution-Aware Feature Selection for SAEs","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-05T14:27:25.448907Z"},"links":{"cited_paper":"/paper/2410.13928","citing_paper":"/paper/2508.21324"},"observation_digest":"sha256:4b51fc863e1e1d60f47dc584f6bc5c81d8a6ccd5dbdd69f2177a361fb0bf024c","observation_id":"38cada16-d398-4e83-9f60-d806293f9950","resolution":{"observed_at":"2026-08-05T14:27:25.448907Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2508.21324","last_updated":"2025-08-29T04:42:17Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-05T14:27:24.567504Z","submitted_at":"2025-08-29T04:42:17Z","title":"Distribution-Aware Feature Selection for SAEs"},"reference_resolution":{"displayed":13,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":9,"verified_exact":0,"verified_fuzzy":4},"total_outbound_references":13},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 13 of 13 outbound references and 0 inbound Pith citation observations for arXiv:2508.21324."}