{"as_of":"2026-08-20T00:27:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:a5b719512fa74ffccc224d763164ec1278bd26c5580a9cbf0d8187ca17c1683c","coverage":[{"denominator":173,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":100,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-15T16:30:36.783894Z","state":"measured"},{"denominator":174,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":174,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-19T06:32:44.657259+00:00","state":"measured"},{"denominator":74,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":74,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-16T00:21:22.489512Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"pith","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":2,"observed_at":"2026-08-05T02:28:24.338817Z","source":"pith"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-09T21:23:48.689713Z","title":"R., Zhao, D., Patel, N., Naghiyev, J., LeCun, Y., and Shwartz-Ziv, R","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2501.19099","last_updated":"2025-05-23T12:41:11Z","snapshot_observed_at":"2026-08-18T09:27:52.410768Z","submitted_at":"2025-01-31T12:46:04Z","title":"Elucidating Subspace Perturbation in Zeroth-Order Optimization: Theory and Practice at Scale","version":2},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-08-09T21:23:48.689713Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2501.19099"},"observation_digest":"sha256:7af2c5707ac3642b97791b99c4f0dc6296a46ad416d6fc6bd136a4c89a44cb45","observation_id":"a0447577-4b80-4fd5-9157-9650a6c037bf","resolution":{"observed_at":"2026-08-09T21:23:48.689713Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-15T21:21:22.487821Z","title":"Layer by layer: Uncovering hidden representations in language mod- els.arXiv:2502.02013, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.10046","last_updated":"2025-05-15T07:43:23Z","snapshot_observed_at":"2026-08-16T10:46:43.017879Z","submitted_at":"2025-05-15T07:43:23Z","title":"Exploring the Deep Fusion of Large Language Models and Diffusion Transformers for Text-to-Image Synthesis","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-15T21:21:22.487821Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2505.10046"},"observation_digest":"sha256:219f76279ba8c806b76fa80a49f6ea4988d067ab5c7e1e59e8232cff564b8ec5","observation_id":"32912298-69e7-4213-9abf-079c51bc419b","resolution":{"observed_at":"2026-08-15T21:21:22.487821Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-07T15:28:09.716568Z","title":"Layer by layer: Uncovering hidden representations in language models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.17099","last_updated":"2025-05-21T05:29:55Z","snapshot_observed_at":"2026-08-19T08:07:00.043951Z","submitted_at":"2025-05-21T05:29:55Z","title":"Learning Interpretable Representations Leads to Semantically Faithful EEG-to-Text Generation","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T15:28:09.716568Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2505.17099"},"observation_digest":"sha256:ca7625d58818c42292d1a1d2ee7c29e84f5a3dbe385850a526932d3a972c66e1","observation_id":"64c46a49-80e5-4931-ba5b-32eeb903e014","resolution":{"observed_at":"2026-08-07T15:28:09.716568Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-07T14:42:13.180944Z","title":"Layer by layer: Uncovering hidden representations in language models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.17998","last_updated":"2025-05-23T15:03:51Z","snapshot_observed_at":"2026-08-14T18:24:13.420208Z","submitted_at":"2025-05-23T15:03:51Z","title":"TRACE for Tracking the Emergence of Semantic Representations in Transformers","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T14:42:13.180944Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2505.17998"},"observation_digest":"sha256:0d7d5a5e00990fb548942fdf7d4e234c822124479d101135b40bcdc3b0df5cac","observation_id":"d324e018-af54-4058-b8c0-3f9035691025","resolution":{"observed_at":"2026-08-07T14:42:13.180944Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-07T14:40:52.012569Z","title":"Skean, M","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.18132","last_updated":"2025-06-17T09:32:43Z","snapshot_observed_at":"2026-08-14T05:10:56.832153Z","submitted_at":"2025-05-23T17:41:54Z","title":"BiggerGait: Unlocking Gait Recognition with Layer-wise Representations from Large Vision Models","version":3},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T14:40:52.012569Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2505.18132"},"observation_digest":"sha256:0d839415c529b20ee7f22939ad8eaf9c5a8ca6ab8233a2e2716e6ef592868f03","observation_id":"704b3598-2ac1-4503-8393-456cf31e0c5b","resolution":{"observed_at":"2026-08-07T14:40:52.012569Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-07T14:02:50.561305Z","title":"Layer by layer: Uncovering hidden representations in language models.arXiv preprint arXiv:2502.02013, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.20225","last_updated":"2025-05-26T17:06:25Z","snapshot_observed_at":"2026-08-14T18:24:52.393233Z","submitted_at":"2025-05-26T17:06:25Z","title":"FLAME-MoE: A Transparent End-to-End Research Platform for Mixture-of-Experts Language Models","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T14:02:50.561305Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2505.20225"},"observation_digest":"sha256:f7ee376a840d375aaa9419ba46725ad77d9166216e43fc44fc3c979c4b77c785","observation_id":"3b62554e-adc6-4893-8488-a990ba9a1b9b","resolution":{"observed_at":"2026-08-07T14:02:50.561305Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-07T14:21:38.302301Z","title":"Layer by layer: Uncovering hidden representations in language models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.23790","last_updated":"2025-05-25T22:31:24Z","snapshot_observed_at":"2026-08-16T23:09:15.764496Z","submitted_at":"2025-05-25T22:31:24Z","title":"Rethinking the Understanding Ability across LLMs through Mutual Information","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T14:21:38.302301Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2505.23790"},"observation_digest":"sha256:031a8b13b8b81caaa9cc46c921c32fa959319abba26a3e222cfadc012d7b8981","observation_id":"e78e6351-2daa-4618-9c63-9105da7e6fec","resolution":{"observed_at":"2026-08-07T14:21:38.302301Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-07T12:24:23.145178Z","title":"Layer by layer: Uncovering hidden representations in language models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24597","last_updated":"2025-05-30T13:45:19Z","snapshot_observed_at":"2026-08-14T18:24:48.983843Z","submitted_at":"2025-05-30T13:45:19Z","title":"Mixture-of-Experts for Personalized and Semantic-Aware Next Location Prediction","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T12:24:23.145178Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2505.24597"},"observation_digest":"sha256:7a38907cbf1412f72ada26c730e2739f3203257bd08faaf1e2f2d3d018f11df2","observation_id":"e11a1a2a-9afa-46e8-90a1-2263d1f5e603","resolution":{"observed_at":"2026-08-07T12:24:23.145178Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2507.05387","last_updated":"2026-04-03T19:07:42Z","snapshot_observed_at":"2026-08-15T02:17:51.224778Z","submitted_at":"2025-07-07T18:18:51Z","title":"The Generalization Ridge: Information Flow in Natural Language Generation","version":5},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-19T05:30:38.612759Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2507.05387"},"observation_digest":"sha256:9254ce9d8f6120a5c8876ec9d93ed12d8ad6eb88ead4957284adf88d3b30ebd7","observation_id":"643f1f6a-8032-42c9-8164-66cf0a1df9bf","resolution":{"observed_at":"2026-05-19T05:32:05.678765Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-06T19:14:30.983300Z","title":"Layer by layer: Uncovering hidden representations in language models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.06203","last_updated":"2025-07-10T16:43:36Z","snapshot_observed_at":"2026-08-07T04:57:37.201438Z","submitted_at":"2025-07-08T17:29:07Z","title":"A Survey on Latent Reasoning","version":2},"reference_index":96,"source":"arxiv_source","source_observed_at":"2026-08-06T19:14:30.983300Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2507.06203"},"observation_digest":"sha256:fa40c936a6987bfdc1fa78070c5e710da5f61aec9dbc4ee838a9b34faac5d1dd","observation_id":"e63889f9-8b2f-4859-ba78-abafe48faaa9","resolution":{"observed_at":"2026-08-06T19:14:30.983300Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-06T17:48:07.190751Z","title":"R., Zhao, D., Patel, N., Naghiyev, J., LeCun, Y., and Shwartz-Ziv, R","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.09973","last_updated":"2025-07-14T06:43:00Z","snapshot_observed_at":"2026-08-08T02:31:23.778902Z","submitted_at":"2025-07-14T06:43:00Z","title":"Tiny Reward Models","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-06T17:48:07.190751Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2507.09973"},"observation_digest":"sha256:0d3e92562095a8136559bdab84d4f459fdd428c397ba9cc5b39602e1ea0af66d","observation_id":"986bf3e7-c475-4689-8ffb-a0a7508a53d1","resolution":{"observed_at":"2026-08-06T17:48:07.190751Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-06T16:31:51.367025Z","title":"Layer by layer: Uncovering hidden representations in language models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.13236","last_updated":"2025-07-17T15:47:22Z","snapshot_observed_at":"2026-08-07T18:24:30.449370Z","submitted_at":"2025-07-17T15:47:22Z","title":"Enhancing Cross-task Transfer of Large Language Models via Activation Steering","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T16:31:51.367025Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2507.13236"},"observation_digest":"sha256:564723ae24fa3511b5e828d0284dfc6b7340e12d522a16946b39450a81137e37","observation_id":"c26df17d-df75-4dfd-9e5b-f09f6607941d","resolution":{"observed_at":"2026-08-06T16:31:51.367025Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T21:10:35.752852Z","title":"R., Zhao, D., Patel, N., Naghiyev, J., LeCun, Y., and Shwartz-Ziv, R","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.10057","last_updated":"2025-08-12T21:38:46Z","snapshot_observed_at":"2026-08-13T13:10:53.109705Z","submitted_at":"2025-08-12T21:38:46Z","title":"Large Language Models Show Signs of Alignment with Human Neurocognition During Abstract Reasoning","version":1},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-08-05T21:10:35.752852Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2508.10057"},"observation_digest":"sha256:a7c17017b4391bb27b26d8fabab2ebec337bc20eebd20c80ca400e85de4700ae","observation_id":"75c65697-76fd-41a9-8427-7b2a413cf990","resolution":{"observed_at":"2026-08-05T21:10:35.752852Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T16:47:20.290385Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models , 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.17734","last_updated":"2025-08-25T07:21:49Z","snapshot_observed_at":"2026-08-11T03:00:13.674300Z","submitted_at":"2025-08-25T07:21:49Z","title":"Layerwise Importance Analysis of Feed-Forward Networks in Transformer-based Language Models","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-05T16:47:20.290385Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2508.17734"},"observation_digest":"sha256:6c98df37e9927987e3e8993e00d8331c56eb3ab7ea581b75ba760d87264d4f53","observation_id":"8fc0d314-0d4d-48bc-9497-53fbcc73747a","resolution":{"observed_at":"2026-08-05T16:47:20.290385Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-04T23:34:30.729282Z","title":"Layer by layer: Uncovering hidden representations in language models.arXiv preprint arXiv:2502.02013, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.06472","last_updated":"2025-09-09T08:54:11Z","snapshot_observed_at":"2026-08-16T22:28:12.226820Z","submitted_at":"2025-09-08T09:37:20Z","title":"Rethinking LLM Parametric Knowledge as Post-retrieval Confidence for Dynamic Retrieval and Reranking","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-04T23:34:30.729282Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2509.06472"},"observation_digest":"sha256:6bb27251f4276bfc8fde7ee260fd0f1f8e115daab40dd0e310daff84cb121cb0","observation_id":"9c91241a-debf-4e80-a783-da2bdf65a134","resolution":{"observed_at":"2026-08-04T23:34:30.729282Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-04T23:33:34.841718Z","title":"R., Zhao, D., Patel, N., Naghiyev, J., LeCun, Y., and Shwartz-Ziv, R","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.06518","last_updated":"2025-09-08T10:24:19Z","snapshot_observed_at":"2026-08-14T18:24:10.709024Z","submitted_at":"2025-09-08T10:24:19Z","title":"Crown, Frame, Reverse: Layer-Wise Scaling Variants for LLM Pre-Training","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-04T23:33:34.841718Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2509.06518"},"observation_digest":"sha256:b6402cc0ce7d56946f7a630c0811a5a2b2a8f83556cf43abebbb049aa4ce5d08","observation_id":"6119f402-35ea-4626-abc0-ad606cedb02f","resolution":{"observed_at":"2026-08-04T23:33:34.841718Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-04T20:53:30.156275Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.08315","last_updated":"2025-09-10T06:32:49Z","snapshot_observed_at":"2026-08-15T18:17:20.123085Z","submitted_at":"2025-09-10T06:32:49Z","title":"EvolKV: Evolutionary KV Cache Compression for LLM Inference","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-04T20:53:30.156275Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2509.08315"},"observation_digest":"sha256:8537c333015e13e141d12896a6953de72f72c91f3d05908e85825d1d6135b48f","observation_id":"d8738f86-aa29-4df8-809c-ad338e73fd5c","resolution":{"observed_at":"2026-08-04T20:53:30.156275Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-03T23:27:37.302772Z","title":"Layer by layer: Uncovering hidden representations in language models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.05933","last_updated":"2026-06-24T14:20:52Z","snapshot_observed_at":"2026-08-07T19:34:44.238702Z","submitted_at":"2025-11-08T08:56:29Z","title":"Reinforcement Learning Improves Traversal of Parametric Knowledge in LLMs","version":2},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-03T23:27:37.302772Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2511.05933"},"observation_digest":"sha256:f1155c791d40661f5332c6e40ddfd64bcd5a2d9d5bed1930208ca37303b90932","observation_id":"00ab11fc-7e79-42a3-b60f-f2c392629488","resolution":{"observed_at":"2026-08-03T23:27:37.302772Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2511.06516","last_updated":"2026-06-28T18:20:13Z","snapshot_observed_at":"2026-08-15T08:59:56.241757Z","submitted_at":"2025-11-09T19:58:24Z","title":"You Had One Job: Per-Task Quantization Using LLMs' Hidden Representations","version":3},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-05-21T18:46:04.926179Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2511.06516"},"observation_digest":"sha256:16ea38446d208992dce9ad16d60bae22748b5b4ed847fa87129552de90648be5","observation_id":"16697031-5ad6-4c31-8449-ea9662681579","resolution":{"observed_at":"2026-05-21T18:50:30.251328Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-03T23:21:41.258490Z","title":"R., Zhao, D., Patel, N., Naghiyev, J., LeCun, Y., and Shwartz-Ziv, R","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.06516","last_updated":"2026-06-28T18:20:13Z","snapshot_observed_at":"2026-08-15T08:59:56.241757Z","submitted_at":"2025-11-09T19:58:24Z","title":"You Had One Job: Per-Task Quantization Using LLMs' Hidden Representations","version":4},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-03T23:21:41.258490Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2511.06516"},"observation_digest":"sha256:ac2116a40e83400aaff0fa4158c606fef7c11aec464c77a1f0eae0298250d864","observation_id":"117f3b57-dbf6-4ea1-8824-773c880e3bef","resolution":{"observed_at":"2026-08-03T23:21:41.258490Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2601.03233","last_updated":"2026-01-06T18:24:41Z","snapshot_observed_at":"2026-08-18T20:29:58.857701Z","submitted_at":"2026-01-06T18:24:41Z","title":"LTX-2: Efficient Joint Audio-Visual Foundation Model","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-13T07:06:20.470686Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2601.03233"},"observation_digest":"sha256:ed07405b667b0191cfbfaee891eb451671666366ab8e260fac98c3e45cb9cdf0","observation_id":"37c9b1ad-76a6-4247-a8a7-b074a68ec74c","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2601.21619","last_updated":"2026-05-09T01:57:30Z","snapshot_observed_at":"2026-08-11T01:08:48.247252Z","submitted_at":"2026-01-29T12:22:45Z","title":"On the Overscaling Curse of Parallel Thinking: System Efficacy Contradicts Sample Efficiency","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-16T10:38:09.875786Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2601.21619"},"observation_digest":"sha256:1b300f40b11f42fb2f6fa265c863b570dfddd74b3448bc06fe09982dab4f8861","observation_id":"12e4354c-2f55-497a-bb8f-2c392140af50","resolution":{"observed_at":"2026-05-16T10:40:51.366510Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2602.21750","last_updated":"2026-04-24T15:04:27Z","snapshot_observed_at":"2026-08-15T00:56:06.072424Z","submitted_at":"2026-02-25T10:06:12Z","title":"From Words to Amino Acids: Does the Curse of Depth Persist?","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-05-15T19:16:01.481502Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2602.21750"},"observation_digest":"sha256:58d429ce0c77a22b09964da727a456effc43c0a03b49dd401cf98e4a42812110","observation_id":"192ce5e2-8039-4af1-b63f-51d22e88d3f1","resolution":{"observed_at":"2026-05-15T19:16:30.957867Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2603.07475","last_updated":"2026-08-02T20:40:52Z","snapshot_observed_at":"2026-08-17T12:15:25.151167Z","submitted_at":"2026-03-08T05:31:52Z","title":"A Comparative analysis of Layer-wise Representational Capacity in AR and Diffusion LLMs","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-15T15:03:23.792608Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2603.07475"},"observation_digest":"sha256:3f4cb5caf9cc24036c5e25cd2196ca0869a3ee8dc04fb2958dbeb248d95c8829","observation_id":"6d3364af-3104-4d76-bdfd-3177e8274d77","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-04T05:57:08.053888Z","title":"Layer by layer: Uncovering hidden representations in language models.arXiv preprint arXiv:2502.02013,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2603.07475","last_updated":"2026-08-02T20:40:52Z","snapshot_observed_at":"2026-08-17T12:15:25.151167Z","submitted_at":"2026-03-08T05:31:52Z","title":"A Comparative analysis of Layer-wise Representational Capacity in AR and Diffusion LLMs","version":4},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-04T05:57:08.053888Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2603.07475"},"observation_digest":"sha256:f4c11bcc2f1f60258c1f4ebaf8c09f4bf86d573c80247eb140751835d3af4691","observation_id":"f495359f-8576-4650-8afd-74a105c47da5","resolution":{"observed_at":"2026-08-04T05:57:08.053888Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2603.12451","last_updated":"2026-07-15T13:22:58Z","snapshot_observed_at":"2026-08-07T07:24:11.757671Z","submitted_at":"2026-03-12T21:05:33Z","title":"Overcoming the Modality Gap in Context-Aided Forecasting","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-15T12:29:43.025344Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2603.12451"},"observation_digest":"sha256:d31b3c165678dd3d1dbe5e413c29fa5b27fa9f8024c76760dd5236eedc701e3f","observation_id":"638a2585-c086-470b-8259-643c3d5e8fa5","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-02T18:22:23.210467Z","title":"Layer by Layer: Uncovering hidden representations in language models.arXiv preprint arXiv:2502.02013,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2603.12451","last_updated":"2026-07-15T13:22:58Z","snapshot_observed_at":"2026-08-07T07:24:11.757671Z","submitted_at":"2026-03-12T21:05:33Z","title":"Overcoming the Modality Gap in Context-Aided Forecasting","version":4},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-02T18:22:23.210467Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2603.12451"},"observation_digest":"sha256:ae99c4dca7de86681ded45b145a3e36692b4d3adf213a2f9f876718024a0bdd8","observation_id":"957c7492-96a2-4e34-bc96-c2f07eef6d15","resolution":{"observed_at":"2026-08-02T18:22:23.210467Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-07-13T12:56:04.388200Z","title":"arXiv preprint arXiv:2502.02013 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2604.03589","last_updated":"2026-04-04T04:35:20Z","snapshot_observed_at":"2026-08-14T16:45:48.473541Z","submitted_at":"2026-04-04T04:35:20Z","title":"Entropy and Attention Dynamics in Small Language Models: A Trace-Level Structural Analysis on the TruthfulQA Benchmark","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-07-13T12:56:04.388200Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2604.03589"},"observation_digest":"sha256:d25d958a716ff822bf1d6aa78873595b551885021bf0257490a19b1528b9be05","observation_id":"3d2e27c0-0185-4349-8c9e-074d358fb090","resolution":{"observed_at":"2026-07-13T12:56:04.388200Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2604.06377","last_updated":"2026-05-05T18:37:43Z","snapshot_observed_at":"2026-08-11T13:29:25.505319Z","submitted_at":"2026-04-07T19:02:10Z","title":"The Master Key Hypothesis: Unlocking Cross-Model Capability Transfer via Linear Subspace Alignment","version":3},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-05-10T18:36:44.401045Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2604.06377"},"observation_digest":"sha256:ecca055626217ad09143b75bb99769b80e72b9ef926e3fb1f69c084e96a22334","observation_id":"86d5bae5-eaf0-4d57-a865-d7c846bb7149","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2604.09425","last_updated":"2026-04-10T15:38:00Z","snapshot_observed_at":"2026-08-17T05:59:10.434500Z","submitted_at":"2026-04-10T15:38:00Z","title":"Do Vision Language Models Need to Process Image Tokens?","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-10T18:26:37.371488Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2604.09425"},"observation_digest":"sha256:28c1de0501662112731764db387afcedc258e83d20f5c6b096985f91c36fa1b6","observation_id":"05ea5b76-5c40-473a-9f5b-f2edf56ccfbc","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2604.10448","last_updated":"2026-04-19T03:58:00Z","snapshot_observed_at":"2026-08-15T05:17:22.597707Z","submitted_at":"2026-04-12T04:11:12Z","title":"Instruction Data Selection via Answer Divergence","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-10T16:15:59.297896Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2604.10448"},"observation_digest":"sha256:5867297edffc5b19cdbaf035d3b0d242bc7dcc5c2d8ee3c0d460b7d1d5a801d5","observation_id":"6fa136fe-6294-481e-95b1-efd45dff5e81","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2604.14838","last_updated":"2026-04-16T10:16:11Z","snapshot_observed_at":"2026-07-06T23:02:31.717911Z","submitted_at":"2026-04-16T10:16:11Z","title":"Intermediate Layers Encode Optimal Biological Representations in Single-Cell Foundation Models","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-10T11:48:46.157380Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2604.14838"},"observation_digest":"sha256:970a6b53be0e69c08c7507c80382259d1983d1fd64f81d9f86bb8a1c0dd77145","observation_id":"a1428d06-fd34-41e6-b9f1-ba564b5932ad","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-07-12T19:12:16.091691Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2604.16878","last_updated":"2026-08-13T02:38:42Z","snapshot_observed_at":"2026-08-19T21:52:53.836763Z","submitted_at":"2026-04-18T07:01:28Z","title":"OC-Distill: Ontology-aware Contrastive Learning with Cross-Modal Distillation for ICU Risk Prediction","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-07-12T19:12:16.091691Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2604.16878"},"observation_digest":"sha256:2049ef14af2b4c4641c97ed8b9d78785f39ccf9da0fd0eb34c1fc8aa47590395","observation_id":"a70ab8c0-4922-4f38-8042-3a31b02f79ff","resolution":{"observed_at":"2026-07-12T19:12:16.091691Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2604.16879","last_updated":"2026-04-22T07:00:49Z","snapshot_observed_at":"2026-08-11T13:26:32.867543Z","submitted_at":"2026-04-18T07:07:53Z","title":"Adaptive Forensic Feature Refinement via Intrinsic Importance Perception","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-05-10T07:57:28.950520Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2604.16879"},"observation_digest":"sha256:d17f0224df33bb46be33a0f6f498d535ca37c1f9449218309a68d9fdc49c825c","observation_id":"a1e52f2f-783e-4fb6-9789-9ce9ea585a54","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2604.16902","last_updated":"2026-04-29T05:22:54Z","snapshot_observed_at":"2026-08-10T23:36:15.762831Z","submitted_at":"2026-04-18T08:25:52Z","title":"Beyond Text-Dominance: Understanding Modality Preference of Omni-modal Large Language Models","version":3},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-05-10T07:02:02.752466Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2604.16902"},"observation_digest":"sha256:967c16548e1f2d77c17866c3e5f8cd7b1c464eed8480ec88df5577ca7df628d8","observation_id":"b26e66c1-9b55-4b1b-a1fe-557aa2971347","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2604.18519","last_updated":"2026-04-20T17:17:07Z","snapshot_observed_at":"2026-08-17T04:02:17.790778Z","submitted_at":"2026-04-20T17:17:07Z","title":"LLM Safety From Within: Detecting Harmful Content with Internal Representations","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-05-10T04:33:54.058475Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2604.18519"},"observation_digest":"sha256:eb3ca90f7252c4530c00bd6bf7218a524548f43ae95483a4f63e815c82ad987c","observation_id":"a4b800ca-a3da-43b0-a7b1-a1fda654f7fb","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2604.27169","last_updated":"2026-04-29T20:17:43Z","snapshot_observed_at":"2026-08-16T09:10:57.313128Z","submitted_at":"2026-04-29T20:17:43Z","title":"Semantic Structure of Feature Space in Large Language Models","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-07T09:45:00.255245Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2604.27169"},"observation_digest":"sha256:489e652b6dafa43ba7a7b95305bdde5ff1107937ba210c050e5bb319d7a1970a","observation_id":"f7f73460-4eea-43f4-831a-a40f8a250b64","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2605.00226","last_updated":"2026-04-30T21:04:38Z","snapshot_observed_at":"2026-08-16T18:04:22.107795Z","submitted_at":"2026-04-30T21:04:38Z","title":"Why Do LLMs Struggle in Strategic Play? Broken Links Between Observations, Beliefs, and Actions","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-09T20:12:36.456731Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.00226"},"observation_digest":"sha256:ba3e88591f4480dc0fbdb98bd22b47bfb8743da7619cbc48603593f15ef37dd2","observation_id":"66cd463d-9cce-4dd3-a969-3fda975097e5","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2605.05668","last_updated":"2026-05-07T04:45:52Z","snapshot_observed_at":"2026-08-13T23:26:24.177992Z","submitted_at":"2026-05-07T04:45:52Z","title":"Large Vision-Language Models Get Lost in Attention","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-05-08T11:54:01.224588Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.05668"},"observation_digest":"sha256:29a4c191d8b97f3f684616eb5ad38e2b43b8d5c12006157fca85cf4d3a3bfba1","observation_id":"059b7852-da1f-4841-986f-df280b06a196","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2605.09430","last_updated":"2026-05-12T03:20:13Z","snapshot_observed_at":"2026-08-12T22:45:19.762221Z","submitted_at":"2026-05-10T09:07:20Z","title":"FlashAR: Efficient Post-Training Acceleration for Autoregressive Image Generation","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-12T02:20:25.071428Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.09430"},"observation_digest":"sha256:5fefd5409040247206370d86de86742586db899a79df676f0fb8ef4b1113dab1","observation_id":"fa66e777-f9b3-4610-b8f5-54146f3c7af4","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2605.09430","last_updated":"2026-05-12T03:20:13Z","snapshot_observed_at":"2026-08-12T22:45:19.762221Z","submitted_at":"2026-05-10T09:07:20Z","title":"FlashAR: Efficient Post-Training Acceleration for Autoregressive Image Generation","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-13T05:54:24.248910Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.09430"},"observation_digest":"sha256:b1dc8059f7ec5bc21452c4f68d2b706f101c34b2995f6123e41f9a0cb2b7cac8","observation_id":"216d043b-7da5-4d20-9463-6e1dae92f61c","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2605.11739","last_updated":"2026-05-21T08:50:47Z","snapshot_observed_at":"2026-08-18T16:51:41.400204Z","submitted_at":"2026-05-12T08:19:15Z","title":"Learning to Foresee: Unveiling the Unlocking Efficiency of On-Policy Distillation","version":1},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-05-13T07:01:46.325498Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.11739"},"observation_digest":"sha256:8c3aca9ab19c355f6887d6083ef7bfb297105aff738900de0c9cc075f716e714","observation_id":"7587f3a6-0db6-483d-bcf9-80e09e5c6538","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2605.11739","last_updated":"2026-05-21T08:50:47Z","snapshot_observed_at":"2026-08-18T16:51:41.400204Z","submitted_at":"2026-05-12T08:19:15Z","title":"Learning to Foresee: Unveiling the Unlocking Efficiency of On-Policy Distillation","version":2},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-05-14T21:07:43.168136Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.11739"},"observation_digest":"sha256:9f64f5fe841a715da0aedbcc23c5b6682196c4aeaded4f274525b0d4345c8cdf","observation_id":"2c7d2fbb-10e7-4364-98f0-c8efa3d627c8","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2605.11808","last_updated":"2026-05-12T09:03:19Z","snapshot_observed_at":"2026-08-17T19:57:35.135626Z","submitted_at":"2026-05-12T09:03:19Z","title":"Mitigating Action-Relation Hallucinations in LVLMs via Relation-aware Visual Enhancement","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-05-13T05:44:05.093926Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.11808"},"observation_digest":"sha256:39b7ba61313773502104fcd91d3f82837a7bc9316de2170a5d71ac010e17c197","observation_id":"f5eba219-5514-426a-b3e2-24fb9a78b747","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2605.11856","last_updated":"2026-05-12T09:40:03Z","snapshot_observed_at":"2026-07-06T23:23:39.723268Z","submitted_at":"2026-05-12T09:40:03Z","title":"UniVLR: Unifying Text and Vision in Visual Latent Reasoning for Multimodal LLMs","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-13T07:32:03.466222Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.11856"},"observation_digest":"sha256:4b9e00960db2f0a1f6a82f2382b6576a3572bc6a0a7ddabf110b1d7b6ec0cebd","observation_id":"2d9ec589-cfb8-4756-bf8c-b3f019891e01","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2605.12714","last_updated":"2026-05-12T20:22:45Z","snapshot_observed_at":"2026-08-16T06:18:32.460798Z","submitted_at":"2026-05-12T20:22:45Z","title":"Layer-wise Representation Dynamics: An Empirical Investigation Across Embedders and Base LLMs","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-05-14T21:50:10.564922Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.12714"},"observation_digest":"sha256:0bb4ad3c1f3f289ec26d0bc9de548f892926161825f509ced5cdfd8bbaa7ccaa","observation_id":"1ac69b58-fa81-4e4a-83c0-e0db3887e24e","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2605.12765","last_updated":"2026-07-14T01:51:38Z","snapshot_observed_at":"2026-08-15T00:45:07.253204Z","submitted_at":"2026-05-12T21:26:25Z","title":"Inference-Time Machine Unlearning via Gated Activation Redirection","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-05-14T20:58:47.559501Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.12765"},"observation_digest":"sha256:e1e715489b09200f5f4b2662b33708baaed763ee3541eaf5820ae4484fed7b82","observation_id":"ab1d68ba-7668-402c-a8b7-de2badcacc9a","resolution":{"observed_at":"2026-05-15T16:30:37.615120Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2605.12765","last_updated":"2026-07-14T01:51:38Z","snapshot_observed_at":"2026-08-15T00:45:07.253204Z","submitted_at":"2026-05-12T21:26:25Z","title":"Inference-Time Machine Unlearning via Gated Activation Redirection","version":2},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-05-20T21:57:53.833175Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.12765"},"observation_digest":"sha256:e45e167d650de2f575857d98f0f6ed747fea3a115ed9e4af182c2689c49aa7fd","observation_id":"a0ab0069-b266-4531-96c3-b9cb8b0f73c5","resolution":{"observed_at":"2026-05-20T21:59:06.030088Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2605.12890","last_updated":"2026-05-13T02:14:21Z","snapshot_observed_at":"2026-07-06T23:24:32.412966Z","submitted_at":"2026-05-13T02:14:21Z","title":"Steer-to-Detect: Probing Hidden Representations for Detection of LLM-Generated Texts","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-06-30T21:53:20.873904Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.12890"},"observation_digest":"sha256:b10e369d9d27057ac32bcf756ca8ebef194299802fef0026c42ae20bf3bc4397","observation_id":"09754284-151c-4362-9b70-e18ee34a924c","resolution":{"observed_at":"2026-06-30T21:55:05.801380Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2605.14738","last_updated":"2026-06-10T08:53:16Z","snapshot_observed_at":"2026-08-12T21:07:08.281089Z","submitted_at":"2026-05-14T12:01:05Z","title":"TAPIOCA: Why Task- Aware Pruning Improves OOD model Capability","version":2},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-05-22T10:09:46.259358Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.14738"},"observation_digest":"sha256:2859e5710632c6c4b19f3a2ad3cc0df02f66afa61d4f2e9ac75e9d08de31becc","observation_id":"afc4691a-c24e-4cd5-8a17-12f162a38a2b","resolution":{"observed_at":"2026-05-22T10:11:23.173954Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2605.17084","last_updated":"2026-05-16T17:01:17Z","snapshot_observed_at":"2026-08-08T22:04:33.865459Z","submitted_at":"2026-05-16T17:01:17Z","title":"Scale Determines Whether Language Models Organize Representation Geometry for Prediction","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-20T15:29:49.147997Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.17084"},"observation_digest":"sha256:6da016bad007d82d155b3aabd1d1187559e59a94dd92abc46f7d3d4b4d29c212","observation_id":"aba8df6f-183a-4e0e-81e6-b5e580d37ab2","resolution":{"observed_at":"2026-05-20T15:33:25.671690Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2605.23033","last_updated":"2026-05-21T20:58:42Z","snapshot_observed_at":"2026-08-19T00:05:14.528322Z","submitted_at":"2026-05-21T20:58:42Z","title":"Uncovering the Latent Potential of Deep Intermediate Representations","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-05-25T05:36:24.743558Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.23033"},"observation_digest":"sha256:87eb9ea3c3f474e651c62f9c66515ef97a00b622cfcefe8318610c0d4c4f12ec","observation_id":"567c69c1-f9d6-452e-8563-9a178b1478cd","resolution":{"observed_at":"2026-05-25T05:36:39.306766Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2605.24956","last_updated":"2026-07-12T09:24:31Z","snapshot_observed_at":"2026-08-07T14:43:53.645603Z","submitted_at":"2026-05-24T09:13:12Z","title":"NITP: Next Implicit Token Prediction for LLM Pre-training","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-06-30T12:23:42.587689Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.24956"},"observation_digest":"sha256:9be1f60a6a1238586b36709eb24dbe4aa4844afbaa5ef358433991663d994384","observation_id":"b5319dda-ab9b-4ebc-8370-c6657b1d371f","resolution":{"observed_at":"2026-06-30T12:24:39.165800Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2605.24956","last_updated":"2026-07-12T09:24:31Z","snapshot_observed_at":"2026-08-07T14:43:53.645603Z","submitted_at":"2026-05-24T09:13:12Z","title":"NITP: Next Implicit Token Prediction for LLM Pre-training","version":2},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-07-04T00:38:33.708836Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.24956"},"observation_digest":"sha256:2ed9b37a44e5ed65880047d66a13911b985a5c2a397840e51426d9092ecf92c9","observation_id":"4e39a470-2357-42d4-8e2b-3bdb59bd1bba","resolution":{"observed_at":"2026-07-04T00:39:16.287745Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-07-14T18:45:28.635910Z","title":"R., Zhao, D., Patel, N., Naghiyev, J., LeCun, Y ., and Shwartz-Ziv, R","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2605.24956","last_updated":"2026-07-12T09:24:31Z","snapshot_observed_at":"2026-08-07T14:43:53.645603Z","submitted_at":"2026-05-24T09:13:12Z","title":"NITP: Next Implicit Token Prediction for LLM Pre-training","version":3},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-07-14T18:45:28.635910Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.24956"},"observation_digest":"sha256:af11f0123c4c53c77ceacbad01691d77ee017d156cdf209665e17e18109c9161","observation_id":"25fbb8c4-dd70-4ea6-a349-250c75b9e57b","resolution":{"observed_at":"2026-07-14T18:45:28.635910Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2605.26366","last_updated":"2026-06-02T00:26:44Z","snapshot_observed_at":"2026-08-14T18:24:53.334746Z","submitted_at":"2026-05-25T22:28:23Z","title":"Automatic Layer Selection for Hallucination Detection","version":3},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-06-29T21:15:37.369334Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.26366"},"observation_digest":"sha256:79f811ccbd6aaa94cc83628dbf46149a084b2eb2b7ef6b4926774a2e03b6ca33","observation_id":"541ca24b-dfc0-40a5-aefd-c12c873b6d0b","resolution":{"observed_at":"2026-06-29T22:04:00.849175Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2605.27970","last_updated":"2026-05-27T05:04:29Z","snapshot_observed_at":"2026-08-16T08:39:56.820314Z","submitted_at":"2026-05-27T05:04:29Z","title":"Geometry of Human Perceptual Domains Emerges Transiently in LLM Representations","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-06-29T12:31:21.660647Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2605.27970"},"observation_digest":"sha256:cfebdc330348c21d17bb4466dd9405e4381a6b2585127c2fe7ee02163a682b37","observation_id":"cc842d94-a698-496a-97a6-71162001cd37","resolution":{"observed_at":"2026-06-29T12:33:24.299685Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2606.00535","last_updated":"2026-05-30T05:05:24Z","snapshot_observed_at":"2026-08-16T00:43:44.943332Z","submitted_at":"2026-05-30T05:05:24Z","title":"DREAM-S: Speculative Decoding with Searchable Drafting and Target-Aware Refinement for Multimodal Generation","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-06-28T18:55:51.474956Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2606.00535"},"observation_digest":"sha256:fc29b805339bee1ba35088741561ed4fa55703cf50d4d9bcc32715ba95259e83","observation_id":"57cacfe0-845e-4b47-90e2-ace280556329","resolution":{"observed_at":"2026-06-28T19:02:34.576108Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2606.05875","last_updated":"2026-06-04T08:47:46Z","snapshot_observed_at":"2026-08-14T03:15:05.778194Z","submitted_at":"2026-06-04T08:47:46Z","title":"QCFuse: Query-Aware Cache Fusion via Compressed View for Efficient RAG Serving","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-06-28T01:47:43.240850Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2606.05875"},"observation_digest":"sha256:e054a041727e9e4f86060aaac0f16e2dbd8de9d26d909d7128e37c2ddf3f8169","observation_id":"381184eb-eb0e-42a0-a616-12d7dae7f8b2","resolution":{"observed_at":"2026-07-02T12:46:57.480466Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2606.06906","last_updated":"2026-06-05T04:49:37Z","snapshot_observed_at":"2026-08-02T10:42:01.638389Z","submitted_at":"2026-06-05T04:49:37Z","title":"EASE-TTT: Evidence-Aligned Selective Test-Time Training for Long-Context Question Answering","version":1},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-06-27T22:05:00.537690Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2606.06906"},"observation_digest":"sha256:acc7f54c844f19bf6f457fc716f4895025821fd9325f21b9dfa2bdb4f76f6b11","observation_id":"1f9c7fbb-4257-415a-9807-d692378722b9","resolution":{"observed_at":"2026-07-02T17:17:15.127594Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2606.07604","last_updated":"2026-05-29T09:40:38Z","snapshot_observed_at":"2026-08-03T04:00:50.272221Z","submitted_at":"2026-05-29T09:40:38Z","title":"Contribution Weights: A Geometrical Analysis of Self-Attention Transformers","version":1},"reference_index":74,"source":"arxiv_source","source_observed_at":"2026-06-28T23:29:02.457697Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2606.07604"},"observation_digest":"sha256:4cde0260bc528cd1d6d3068b8d0ad5515427c1cd26bd79a011855b9de06cba01","observation_id":"e354138e-2f39-4995-bb07-f471f1f6f56a","resolution":{"observed_at":"2026-06-28T23:32:46.726027Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2606.19398","last_updated":"2026-06-17T08:39:09Z","snapshot_observed_at":"2026-08-15T13:00:40.457818Z","submitted_at":"2026-06-17T08:39:09Z","title":"S-JEPA : Soft Clustering Anchors for Self-Supervised Speech Representation Learning","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-06-26T19:46:47.653439Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2606.19398"},"observation_digest":"sha256:44fbf0614e88c7c25554f878dd6c49385fd96af0d794cf1ed1fd4be6c2ba6953","observation_id":"bbe4f69e-34a3-46a2-9b2d-fccb541f1388","resolution":{"observed_at":"2026-07-04T02:19:24.233642Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2606.23670","last_updated":"2026-06-22T17:56:25Z","snapshot_observed_at":"2026-08-12T14:08:49.266214Z","submitted_at":"2026-06-22T17:56:25Z","title":"Tapered Language Models","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-06-26T09:11:20.341634Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2606.23670"},"observation_digest":"sha256:8120910f45c622d93eb84c89b9d87445cd33a8f1ff6ed5a7496377aac30cd63c","observation_id":"eea63741-de8d-4b39-87db-b886eeed383e","resolution":{"observed_at":"2026-07-04T10:09:44.022860Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2606.29640","last_updated":"2026-06-28T23:15:02Z","snapshot_observed_at":"2026-08-13T00:45:21.155031Z","submitted_at":"2026-06-28T23:15:02Z","title":"Fast Wireless Foundation Models with Early-Exits","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-06-30T01:44:30.763412Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2606.29640"},"observation_digest":"sha256:8cff6fa6416cc6d016aef40eda180f44178b0c09877bc84bd633547d1e37a986","observation_id":"b5dddb1c-5b71-4699-91c6-299d0ce48e54","resolution":{"observed_at":"2026-06-30T03:24:12.827632Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2606.30813","last_updated":"2026-06-29T18:37:34Z","snapshot_observed_at":"2026-08-14T12:09:47.392599Z","submitted_at":"2026-06-29T18:37:34Z","title":"Gradient Smoothing: Coupling Layer-wise Updates for Improved Optimization","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-07-01T06:36:48.524846Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2606.30813"},"observation_digest":"sha256:34d4060e3d322910d8841931a34b0ac23336fccbdad9dcea7105e5399d7308f5","observation_id":"fb747d1e-de14-4470-b655-9bedaf80887c","resolution":{"observed_at":"2026-07-01T09:25:41.253949Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2606.30961","last_updated":"2026-06-29T22:38:09Z","snapshot_observed_at":"2026-08-17T15:56:28.873092Z","submitted_at":"2026-06-29T22:38:09Z","title":"ElemeNet: Multiscale Molecular Machine Learning with Uncertainty Quantification Across the Periodic Table","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-07-01T00:49:10.411961Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2606.30961"},"observation_digest":"sha256:42c1d66e8b29fb26f6c680fb82b4e709809d465729a4e4d821951453bbb60543","observation_id":"31b1d2d6-8bbc-47de-8a5d-177f51135a2b","resolution":{"observed_at":"2026-07-01T00:55:11.460664Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2607.00434","last_updated":"2026-07-01T04:45:41Z","snapshot_observed_at":"2026-08-08T21:41:10.025868Z","submitted_at":"2026-07-01T04:45:41Z","title":"Information-Regularized Attention for Visual-Centric Reasoning","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-07-02T15:12:43.475802Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2607.00434"},"observation_digest":"sha256:8a97c5cc22c2cb44e2343547c2fd24cb84de8ce0850495c8a2f5b50db77f6635","observation_id":"867a91c2-2f74-4ad7-9047-f14110be65fe","resolution":{"observed_at":"2026-07-02T15:17:07.251203Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-07-11T22:57:36.974470Z","title":"Layer by layer: Uncovering hidden representations in language models,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.03929","last_updated":"2026-07-04T15:47:38Z","snapshot_observed_at":"2026-08-18T16:46:04.158748Z","submitted_at":"2026-07-04T15:47:38Z","title":"Probe, Don't Prompt: A Hidden-State Probe for Metadata Filtering in Multi-Meta-RAG","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-07-11T22:57:36.974470Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2607.03929"},"observation_digest":"sha256:9047d79b2a0de03870d0b43eaa152b4ef76390c7eccdc24a0b48395ef1ef9940","observation_id":"97e58f36-40c5-4e7e-8218-a134ff679e59","resolution":{"observed_at":"2026-07-11T22:57:36.974470Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":"2502.02013","doi":"10.48550/arxiv.2502.02013","metadata_source":"pith","pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","venue":"cs.LG","work_id":"7b4ac06a-e804-4f0a-8305-c45f2735afb5","year":2025},"citing_paper":{"arxiv_id":"2607.06392","last_updated":"2026-07-07T15:25:47Z","snapshot_observed_at":"2026-08-17T20:01:50.771212Z","submitted_at":"2026-07-07T15:25:47Z","title":"InsideSSL: Understanding Self-Supervised Speech Representations using a Model-Centric Perspective","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-07-08T07:33:22.898615Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2607.06392"},"observation_digest":"sha256:b441eff8d5be3b3add4f8733eca29681823f7530c7d59d3f898dcb51fdd1c25a","observation_id":"75a8d211-9522-437a-96ad-26166bbd215e","resolution":{"observed_at":"2026-07-08T07:34:43.095089Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-01T23:15:30.797805Z","title":"Layer by layer: Uncovering hidden representations in language models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.15495","last_updated":"2026-07-16T22:54:30Z","snapshot_observed_at":"2026-08-19T09:54:53.193795Z","submitted_at":"2026-07-16T22:54:30Z","title":"Verbalizable Representations Form a Global Workspace in Language Models","version":1},"reference_index":152,"source":"pdf_text","source_observed_at":"2026-08-01T23:15:30.797805Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2607.15495"},"observation_digest":"sha256:aa5ce9e1e1ef07674774998afe3ea31de584c296ba7ca7b3cfe7b06fa3e3568b","observation_id":"c5fd4910-575a-4f6e-b630-6feb93611ff0","resolution":{"observed_at":"2026-08-01T23:15:30.797805Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-01T13:20:37.546194Z","title":"arXiv preprint arXiv:2502.02013 , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.26565","last_updated":"2026-07-30T02:21:33Z","snapshot_observed_at":"2026-08-17T05:03:09.558782Z","submitted_at":"2026-07-29T07:35:59Z","title":"Representation Trajectories Matters: Complementary Evidence for OOD Detection and Image Classification","version":2},"reference_index":76,"source":"arxiv_source","source_observed_at":"2026-08-01T13:20:37.546194Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2607.26565"},"observation_digest":"sha256:f9bba68720c7bc08a977d771ab573e07a667ba956b54c5f30c7dd74f9c08180e","observation_id":"c23ece97-f158-4978-9fba-ed8aa731dc5b","resolution":{"observed_at":"2026-08-01T13:20:37.546194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-03T16:32:05.577174Z","title":"Layer by layer: Uncovering hidden representations in language models.arXiv preprint arXiv:2502.02013,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28959","last_updated":"2026-07-31T02:26:04Z","snapshot_observed_at":"2026-08-19T02:46:52.788835Z","submitted_at":"2026-07-31T02:26:04Z","title":"Efficient LLM Adversarial Training via Low-Rank Defense and Circuit-Guided Surrogates","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-03T16:32:05.577174Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2607.28959"},"observation_digest":"sha256:f3991dae737b9282fa94c51919357149b8b6eef4b40c882bd02e9b2ef7154e1c","observation_id":"fb6ffc55-2fdb-4cc1-b425-57bc77fd86ea","resolution":{"observed_at":"2026-08-03T16:32:05.577174Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-04T20:22:23.021410Z","title":"ArXiv:2502.02013.https://doi.org/10.48550/arXiv.2502.02013","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.01816","last_updated":"2026-08-03T07:26:45Z","snapshot_observed_at":"2026-08-17T11:42:01.519075Z","submitted_at":"2026-08-03T07:26:45Z","title":"Divergent large language model predictions from convergent representations in ambiguous word pairs","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-04T20:22:23.021410Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2608.01816"},"observation_digest":"sha256:e0d6e72cb6ac31be8d0c186ff1cdebafc5bec79572e4052293ae5ca5b754e165","observation_id":"c78f88d9-d7b6-4b1d-8153-6867186725db","resolution":{"observed_at":"2026-08-04T20:22:23.021410Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.02013","snapshot_observed_at":"2026-08-16T00:21:22.489512Z","title":"Layer by layer: Uncovering hid- den representations in language models.arXiv preprint arXiv:2502.02013, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.12090","last_updated":"2026-08-13T07:47:02Z","snapshot_observed_at":"2026-08-19T11:37:43.941012Z","submitted_at":"2026-08-12T14:13:48Z","title":"Task- and dataset-specific information in protein language models","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-16T00:21:22.489512Z"},"links":{"cited_paper":"/paper/2502.02013","citing_paper":"/paper/2608.12090"},"observation_digest":"sha256:d8e24383684a4712f9e3199bdc31aadb081b076234c3495a5c3f90de08031349","observation_id":"62f72406-c10e-48c5-811f-6bbf8a4b46d6","resolution":{"observed_at":"2026-08-16T00:21:22.489512Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2502.02013/citation-record","integrity":"/paper/2502.02013/integrity","json":"/paper/2502.02013/citation-record.json","paper":"/paper/2502.02013"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"K., Mondal, A","venue":null,"work_id":"8fe0fb88-86eb-414b-b26d-362bdcf0cf9e","year":2022},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:1e81669c196ab6eebe5d53b3da501beb917ec2ec82d85ac86daa54dd0cb5ee14","observation_id":"31cf633f-c9f9-4cc1-a1a0-35f9a0b9dcbb","resolution":{"observed_at":"2026-05-15T16:30:36.891700Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"and Bengio, Y","venue":null,"work_id":"a916bb7d-1cf2-481f-88c9-5d8cc9240c10","year":2017},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:d014d02d20b6fa68673a77f4733c63afab41dc53505b6bacad107fc6197182ab","observation_id":"8bc9f2da-8675-43f1-84a4-a4e459c2299a","resolution":{"observed_at":"2026-05-15T16:30:36.900464Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"R., Subbaraj, G., Gontier, N., LeCun, Y., Rish, I., Shwartz-Ziv, R., and Pal, C","venue":null,"work_id":"b4a8c8bf-2a6e-433d-8438-d1898900a6ac","year":2025},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:8b611b6b46a404aade3f76ceae3914d73208d8e36849dbfa656a99b0bc220c42","observation_id":"3051e983-3955-4369-a128-371c1609b98a","resolution":{"observed_at":"2026-05-15T16:30:36.906242Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Information theory with kernel methods","venue":null,"work_id":"cd38bb25-b39e-4543-8a17-5c6e2db8d88b","year":2022},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:5ade034b8c9bc342d19ec32be888c19b9a017d9978bed31719e0c773e1b2124d","observation_id":"2738776c-7371-49f8-b95f-51e628c84f36","resolution":{"observed_at":"2026-05-15T16:30:36.912641Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"BeIT : Bert pre-training of image transformers","venue":null,"work_id":"c29e2625-5fb6-47c4-b15e-ad4380212f00","year":2022},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:6a2b4d62d8d5557761543b7f5346485c972472d500ac42f1ff94dd014d29089f","observation_id":"a466ba98-cf03-4ef2-8fe6-169754c35537","resolution":{"observed_at":"2026-05-15T16:30:36.919252Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Why do LLMs attend to the first token? arXiv","venue":null,"work_id":"db68e00c-04a3-4f61-ba5f-2399eecf78dd","year":2025},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:4ef3981ff9890e23e1e77d9231e87342502d1929f3665682ca7caaf0a2172a99","observation_id":"c2f9a502-91a6-4cad-9928-28b055e27067","resolution":{"observed_at":"2026-05-15T16:30:36.925019Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"LLM2Vec : Large language models are secretly powerful text encoders","venue":null,"work_id":"976e55ca-33cc-478f-a92c-775122e5a7d9","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:f936db2b37ec56f65620b755c27d505e387e8b07d4ed22619131ab1849d99f4f","observation_id":"0e92d3c8-c143-4bf3-8f9b-94a31f48ad08","resolution":{"observed_at":"2026-05-15T16:30:36.931197Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"G., Bradley, H., O’Brien, K., Hallahan, E., Khan, M","venue":null,"work_id":"3448a721-54f4-42b6-af81-20ad9be4efcc","year":2023},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:d19a002ffb9ec135dc50b25a9deebdd899af8793d957d985fc12a1a3ed7ede86","observation_id":"b17f7ca2-3aa0-4f8c-a5ab-41b8fe171919","resolution":{"observed_at":"2026-05-15T16:30:36.936345Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"P., and Wilming, H","venue":null,"work_id":"c4c2d36c-aa20-45c5-9bd8-1813d2b67783","year":2019},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:d09bf964b016e7d31b44a9d302296434f97dbd6f636268b167ff6e0cd13446df","observation_id":"9f850dfe-85bd-41d0-9cb4-e68da5c58a27","resolution":{"observed_at":"2026-05-15T16:30:36.941430Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Guillotine regularization: Why removing layers is needed to improve generalization in self-supervised learning","venue":null,"work_id":"0a995bde-15da-4a5f-9de6-ba2314b89a08","year":2023},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:61ea1269d4479ebfa27e5df8e18c0030779827c72616fa388fd4a27920164f71","observation_id":"98496c56-071d-4904-a401-55bbeda3969e","resolution":{"observed_at":"2026-05-15T16:30:36.946917Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"6e5d3343-e030-48bc-8d46-fede3fcba765","year":2020},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:c8ed36c599c54d6ae24f152f929261e956ea9721fc2db2e6d0c65afb39751b6a","observation_id":"5ebcd6e6-75f3-4f87-b93e-b920f5dcd7b6","resolution":{"observed_at":"2026-05-15T16:30:36.952551Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"On identifiability in transformers","venue":null,"work_id":"97467427-01b0-4434-af64-bbe459351987","year":2020},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:e478bfeb2ea6287b044fc6e2889955d6165f1e9ab400190fb088b4db4cdb5e22","observation_id":"ec9094ae-588f-4d5d-9e87-2a2e90920a44","resolution":{"observed_at":"2026-05-15T16:30:36.958420Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Discovering latent knowledge in language models without supervision","venue":null,"work_id":"60d61a44-516e-4121-90d6-1d0b11e16d47","year":2023},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:74ca142671f158b6c2c8316e0c4045eba05f7d53489dc5aa15e9cef67105457a","observation_id":"eeab8e8e-9c59-48df-9eef-47609c2a9edc","resolution":{"observed_at":"2026-05-15T16:30:36.963785Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Generative pretraining from pixels","venue":null,"work_id":"3dfd7651-ad20-48e9-aaf5-5931c82f61a7","year":2020},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:747b7a9941c754be8e524739085dfdfe74d1e9558d5d26459f4af6e7ac28b99f","observation_id":"141a0824-ff1c-426c-b68d-a01a49330425","resolution":{"observed_at":"2026-05-15T16:30:36.968493Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Emergence of a high-dimensional abstraction phase in language transformers","venue":null,"work_id":"3d79a865-e3bd-4339-ae8a-edaa06d6e939","year":2025},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:f444f923a51650b8bdcadf4542c0cc9e985d21e1c93c876c287512122e2d3d58","observation_id":"b837f4f1-7d9d-470c-8ce8-af64225313b9","resolution":{"observed_at":"2026-05-15T16:30:36.973636Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"D., and Potts, C","venue":null,"work_id":"db255408-18e0-4513-921e-256c0bf363f4","year":2025},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:173bb618a4916bcad7d6c7666410c05be0dbb8fbbbd78aa8291543ddd7024cc1","observation_id":"414c3fc5-5c62-4788-a85a-1e605feab38c","resolution":{"observed_at":"2026-05-15T16:30:36.977669Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Deepseek- R1 : Incentivizing reasoning capability in LLMs via reinforcement learning","venue":null,"work_id":"f1760c5b-ea65-4952-a347-60ac6d0f73a7","year":2025},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:a547a89710fae0428dc65f4defb91814b0fc7da9899425236acdb15d76c0ad91","observation_id":"c513e41b-540c-43c6-8fc4-e57d26e04c36","resolution":{"observed_at":"2026-05-15T16:30:36.981920Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"K., Aitchison, M., Orseau, L., et al","venue":null,"work_id":"f532436c-4026-4c16-8679-41e7adc485d8","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:4ad95da7cdf0abb446a5a761499b8cb90b908b403e0f8ed6898a144b69f82395","observation_id":"d242bd67-898b-4459-a010-d5d4e1749052","resolution":{"observed_at":"2026-05-15T16:30:36.987343Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"BERT : Pre-training of deep bidirectional transformers for language understanding","venue":null,"work_id":"56f57e1e-6e8d-4917-b3dc-eff6b5de4761","year":2019},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:bcdcd6eed800559063971f32afe5163ec98cb48c618512eafcf5d953bb031a79","observation_id":"facf7fc6-a961-48ef-85f9-b5a87cb25890","resolution":{"observed_at":"2026-05-15T16:30:36.993611Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"An image is worth 16x16 words: Transformers for image recognition at scale","venue":null,"work_id":"86275509-cedb-48a9-ae57-7b6087df18f8","year":2021},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:f31e5c2954490a9954346e0bf9acf8f446f5873c4ce9179193aef01ce8ec663b","observation_id":"94213d65-4a75-4cd4-8aaf-ac49036c31a2","resolution":{"observed_at":"2026-05-15T16:30:36.998822Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The Llama 3 herd of models","venue":null,"work_id":"383498b0-5e3a-4474-a777-c0398c56928b","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:0a45bbcd3f086e22b04284d5bbd9bfd92be8b97dd3b1d546d59a7abac9d7b479","observation_id":"92d9c984-0e6c-48e9-a38a-c406f225495d","resolution":{"observed_at":"2026-05-15T16:30:37.002687Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"A., Toshev, A., Shankar, V., Susskind, J","venue":null,"work_id":"5d5d9118-6ee6-4cbd-85c0-5705b1b87730","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:9ef4c55d81581b20d16f729bf01391f1f339b29b6c82302cf70cdc74387e626c","observation_id":"fbbbd822-2c03-488b-aea6-c9b666292412","resolution":{"observed_at":"2026-05-15T16:30:37.007250Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Not all layers of LLMs are necessary during inference","venue":null,"work_id":"db421106-e0a1-42e5-ae73-6c681e0a29a0","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:86a84b7c2817f97160705aa9589b292c8bc1f83185c6ae902afd70852ca1dfb0","observation_id":"6b30f9a3-f465-4868-a48c-64c539927c73","resolution":{"observed_at":"2026-05-15T16:30:37.011612Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"6e56d872-1ae5-4b91-ab5a-6048c6ba5310","year":2025},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:6265b4f233872a18b482e2e86c7f99db7dd0b58204dfbc1d92aab4386d6d0669","observation_id":"07560d33-e355-46f1-9c14-60607b74d140","resolution":{"observed_at":"2026-05-15T16:30:37.015178Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"RankMe : Assessing the downstream performance of pretrained self-supervised representations by their rank","venue":null,"work_id":"dda2c3a2-39fb-479d-bfac-b8f5ac46b36b","year":2023},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:46718777031549f17e5972a1d3d5870d0bbeb571613aa7a970f485aa40ea2be6","observation_id":"4c455f60-38ff-49f8-aee6-14545da0d511","resolution":{"observed_at":"2026-05-15T16:30:37.019388Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"9237ad75-b36a-4352-8f8c-e121291970b6","year":2014},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:5d6c858fb89a10a0a709cf405e894e1782a3bfcdb649d5e9341a3b1543225680","observation_id":"d7cf5f06-c4b9-47c0-8f09-2d5080c56d6b","resolution":{"observed_at":"2026-05-15T16:30:37.023216Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"and Dao, T","venue":null,"work_id":"8e1bb26d-8fa3-4a10-8449-7bd2ef57e2fc","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:5558aceea8c5541defdaccb36eef29aef529fc67e7789103f1a91514f537265e","observation_id":"f8ac9a06-ef17-43fd-a749-70e0308addee","resolution":{"observed_at":"2026-05-15T16:30:37.026839Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"When attention sink emerges in language models: An empirical view","venue":null,"work_id":"e4df4052-6c15-4e84-873f-7566723bf4e2","year":2025},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:399138a168fac86c1d2bedaf84ebc2ee534f2e9d3dbdfe1fef4250e72d91ccff","observation_id":"177e3ccc-b2f8-4fd2-bab2-de8783762839","resolution":{"observed_at":"2026-05-15T16:30:37.030408Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"and Tegmark, M","venue":null,"work_id":"82e58c79-0793-4a75-92b8-a9756f4958b9","year":2023},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:1ce35eb3908c3a062a4ac496b946dbe3d63105e57c6d405e9f45825830eb3d6c","observation_id":"c1e06328-7b9e-42c9-af01-0f2abff51626","resolution":{"observed_at":"2026-05-15T16:30:37.033646Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Training large language models to reason in a continuous latent space","venue":null,"work_id":"d744a31b-457e-44a4-9b91-fc88e9471c19","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:bfc0417ddf34105c71f5c17081915d00bfa74cb14fde84bb498a500e844f1170","observation_id":"deb9c79b-8ae5-4d77-be80-7c29fcef812e","resolution":{"observed_at":"2026-05-15T16:30:37.037461Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Masked autoencoders are scalable vision learners","venue":null,"work_id":"cc5da116-277b-4e0e-b5b4-d8c9951a7e7c","year":2022},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:def394da2fe0cfee90959c8a98fdacb1f03e0452bb0d253900a38b404557a24e","observation_id":"0a4cc76b-09e8-4d00-9a57-7b8d6d1064f8","resolution":{"observed_at":"2026-05-15T16:30:37.043952Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"and Fedorenko, E","venue":null,"work_id":"c7fffa6d-b2ea-4824-932b-2367ac04d90e","year":2023},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:cb6d262e4fb1374b2b2f8aaa726250b20d7591a520feff097dba7c6ee95cf1e2","observation_id":"ec1d638b-0758-4556-8c85-07fa7e24494a","resolution":{"observed_at":"2026-05-15T16:30:37.047393Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Exploring concept depth: How large language models acquire knowledge at different layers? arXiv","venue":null,"work_id":"99bb7c00-5123-4f52-9fbb-f1bb2ac70b33","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:f825ecf7a75df49732fc8396f433fd183499dc0260f095c033618936c51136c3","observation_id":"cee485c9-3fd3-4211-a9da-16700ad40b26","resolution":{"observed_at":"2026-05-15T16:30:37.050733Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The remarkable robustness of LLMs : Stages of inference? arXiv","venue":null,"work_id":"8b510075-64ed-4791-9447-846760f3d5a9","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:7253f79a22c25b63cc5b5fae15a9b1a5e7415700957d14a6149957f7101af485","observation_id":"e60bba0a-3537-4b3f-972e-95223c355eb5","resolution":{"observed_at":"2026-05-15T16:30:37.054339Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Competition-level code generation with alphacode","venue":null,"work_id":"01efd198-238c-42a0-8a2d-77b2ada3f64a","year":2022},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:0de66d13d58c1e4997f2d80f5e70bb82ee290f480ae2a4d0f1a58e0489367b46","observation_id":"cdf2df38-37f7-4505-bfbe-904b8769cc03","resolution":{"observed_at":"2026-05-15T16:30:37.058155Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"F., Gardner, M., Belinkov, Y., Peters, M","venue":null,"work_id":"fff3c9fb-b519-4ffb-ace1-078aefbc7020","year":2019},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:8433085bbcf152b5928373885736a6db6d3ba05f36f980263483bb68d3451d97","observation_id":"85f6fca0-b958-4937-b2ad-6160991b1c17","resolution":{"observed_at":"2026-05-15T16:30:37.062186Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"NLP Augmentation","venue":null,"work_id":"3bdf562c-9cc5-4dca-92d9-1eb3bb9d4ca9","year":2019},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:419bd8f95d9b2910dd0bc17b9f8bbc00ce1d3bff92191e480a57772c6857c722","observation_id":"f593855e-1962-4d13-8321-18784fb2855d","resolution":{"observed_at":"2026-05-15T16:30:37.066099Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"725e251d-bee9-474d-adb6-064d692b9fea","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:a771ab9324430b79dce7eaf1aa0f6852e16b2be6ac08b647ffcd5e1291b9af23","observation_id":"16e994e3-9300-41f4-9a4a-0c30c7fae6d9","resolution":{"observed_at":"2026-05-15T16:30:37.069826Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"A., Stephenson, C., Tang, H., Kim, Y., and Chung, S","venue":null,"work_id":"fd980f81-a0dd-442c-9831-e851bfecafa4","year":2020},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:c91e5becf4d86a0d0b09b41487237cc7767779c447056b435347896845a382ee","observation_id":"2d5525be-969b-41e8-87d3-f1cb9acd05a8","resolution":{"observed_at":"2026-05-15T16:30:37.073808Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"E., and Biau, G","venue":null,"work_id":"a33e1a53-6df1-4f89-9a76-48f8933cba2c","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:dfc26fcfb88a5fc4fc8d3d03526b83a9e0b5b7f743989a8de6876af8e2b401ec","observation_id":"7e638237-37be-452d-b582-f26034993739","resolution":{"observed_at":"2026-05-15T16:30:37.077829Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Pointer sentinel mixture models","venue":null,"work_id":"d837cf96-4173-45d8-8773-6f02d8e51e74","year":2017},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:3089b7b5e73e854f943070c12471c82ba98d8cd16aec250f43df61407ba61f46","observation_id":"9aeccf3a-b7ab-4c26-8618-feeb9dfa2c8d","resolution":{"observed_at":"2026-05-15T16:30:37.081497Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"MTEB : Massive text embedding benchmark","venue":null,"work_id":"7890d0d6-07cd-4b69-92cd-68dea35bf714","year":2022},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:e2928b6d3d7f442a53e7f14d258f597ab84042ad188969dfe3d6eaa42fc71946","observation_id":"9f89fbd7-2862-4f19-aef7-a7de770fb782","resolution":{"observed_at":"2026-05-15T16:30:37.085002Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"ed4b8016-3742-4211-bf26-a464878cc7a4","year":2018},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:553034c9604224ff2ba867e6235580ca50e13777dbe631e1c1d14abe8f17aff9","observation_id":"70133523-81d8-42ab-a961-10872372a895","resolution":{"observed_at":"2026-05-15T16:30:37.088950Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"DINOv2 : Learning robust visual features without supervision","venue":null,"work_id":"711e915b-66ad-4556-ab1a-85ec6ed298bd","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:a5172470b6d1530038db252113794c64494e2c4b406742be83c5458c2082ef4f","observation_id":"cddf6fa1-a8d0-421b-89fe-3a642202d458","resolution":{"observed_at":"2026-05-15T16:30:37.093264Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"J., Jiang, Y., and Veitch, V","venue":null,"work_id":"604e38b8-af1c-42d7-b4c2-585d5901c2b5","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:1324ff5e4a1923b7a71b7dfa13e0547d2b20f5f6c994595ca7cf8fb0010b843c","observation_id":"a1c4358d-6ce4-4589-9061-ea124867fcb5","resolution":{"observed_at":"2026-05-15T16:30:37.096853Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"J., and Veitch, V","venue":null,"work_id":"b2e89022-ff4c-4462-934e-2b844abac95f","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:6d86383393e34ef9867932525b3033d07c9b3956741aff694601c20041f837e0","observation_id":"948aa19f-04aa-487d-ab32-63f6e98a8fb5","resolution":{"observed_at":"2026-05-15T16:30:37.100961Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"W., Hallacy, C., Ramesh, A., Goh, G., Agarwal, S., Sastry, G., Askell, A., Mishkin, P., Clark, J., et al","venue":null,"work_id":"8af10b68-621f-46c4-924b-2a34bb4a4b8b","year":2021},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:8e9d91f32bfceefbf7b7b7df529f4bf76756cd2cca4162a53106240e3f87cb43","observation_id":"d274b3ba-686b-44d9-a6f4-e8a2aff34b3e","resolution":{"observed_at":"2026-05-15T16:30:37.105526Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"SVCCA : Singular vector canonical correlation analysis for deep learning dynamics and interpretability","venue":null,"work_id":"052ddaa2-38ff-433f-af9d-a6ac36d6891e","year":2017},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:d141becc086141cdfe12b717f070efc618d15c08c97ca2e502f2325e33adabfe","observation_id":"7a07a15d-d55b-4d72-9209-848ff6586e79","resolution":{"observed_at":"2026-05-15T16:30:37.109706Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The shape of learning: Anisotropy and intrinsic dimensions in transformer-based models","venue":null,"work_id":"d637252b-0165-4996-8522-8d44d999be73","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:f36631d83df8216155ff3a6d72b9d1966e40d03364a52178dd08d6e26f7b9155","observation_id":"09a3a45c-3920-474a-9295-ac6fb0456871","resolution":{"observed_at":"2026-05-15T16:30:37.113573Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"On measures of entropy and information","venue":null,"work_id":"8b99d785-3e4f-4392-b015-ebc548bdc3f7","year":1961},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:edf1687ab97f60ee8257d37eef7befc00b3edf491f4b887d1790b3f635bbc280","observation_id":"3f5ebba4-8c44-4db8-bed3-7dc8b457e205","resolution":{"observed_at":"2026-05-15T16:30:37.117793Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"and Vetterli, M","venue":null,"work_id":"f2c3464f-5d5b-41c1-9300-cb263afa6148","year":2007},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:311e32c49ddfac76cf5ee0158e9fc262f95fed6649d651ef8fb784596f8d3376","observation_id":"a4f8ff2b-e03b-4743-ae57-06779f53984f","resolution":{"observed_at":"2026-05-15T16:30:37.121746Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"V., Stadelmann, T., and Grewe, B","venue":null,"work_id":"1eb9a045-0340-4ce8-8859-f85409ff3bf2","year":2025},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:ae6b1092c4043564657cb27f46797b5863db90680e007c7edca9d64784ee2ac1","observation_id":"b7f05245-85a3-407c-96a8-491acaf43cc6","resolution":{"observed_at":"2026-05-15T16:30:37.125541Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"and Smola, A","venue":null,"work_id":"f94d8e02-61d5-4379-a45f-fd3e984bb3d1","year":2018},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:188e3d7fb072a95d7df8eb3023f11b8fcab3b7764293ca22144b0c51b4610360","observation_id":"f323d05b-8b2c-4c87-a286-09f7daced439","resolution":{"observed_at":"2026-05-15T16:30:37.129261Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Information flow in deep neural networks","venue":null,"work_id":"0437405d-3e27-4bbc-a916-40e552cbb82c","year":2022},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:193adece599bb3184af0edde0d6e7a89f6538e24eea8055fe837ebec535436d4","observation_id":"495ff9e2-f293-4ff8-8b38-6af5b0b8a2ef","resolution":{"observed_at":"2026-05-15T16:30:37.133296Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"and Tishby, N","venue":null,"work_id":"4d4fa36b-37c1-4e15-a341-34c138934379","year":2019},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:ee10e4e36f5fbd245aad156986af66d3f358e25006e75bd27e3a01bb6acae854","observation_id":"ab92bbca-ca7c-4244-b509-88149fb6bdc4","resolution":{"observed_at":"2026-05-15T16:30:37.137331Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"G., and LeCun, Y","venue":null,"work_id":"1ef4998c-ce6c-4de3-94fa-99ecf698bd3d","year":2023},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:a61fc07a35e5b4cd600ffa302f63f232a324b7dc136421d113b11200ee2a38e5","observation_id":"81fed7c2-771c-4b72-8649-05def27d8229","resolution":{"observed_at":"2026-05-15T16:30:37.141215Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"77c06990-cf96-40eb-a99f-ed8755e2a2aa","year":2023},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:75508c99bea7521e2a68ddd16ebbdedc4a556f08de0991e28cf17b2dfbc578c3","observation_id":"5eeff69d-7818-4e6a-ad98-0676d60d362f","resolution":{"observed_at":"2026-05-15T16:30:37.144616Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"140c15fe-a771-481b-8bb6-3a3d64b67aff","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:297f611ba6b3c5e4a1d65041a405d58b7594cf22571a7eefa4ea1d01adc024fb","observation_id":"324d99d6-7a06-4edb-ad96-f291bbc045b8","resolution":{"observed_at":"2026-05-15T16:30:37.148126Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Neural representational geometry underlies few-shot concept learning","venue":null,"work_id":"f18c0b8f-a657-44bd-a7ac-215ccbde81e2","year":2022},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:61099a93df97169d0e27ea06a8a7e6b5bc0844e814812109a50add4f4bd5257a","observation_id":"02e443d0-cfe3-4a9d-86d2-eb7d6d157507","resolution":{"observed_at":"2026-05-15T16:30:37.152106Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"BERT rediscovers the classical nlp pipeline","venue":null,"work_id":"4f87c501-3f07-44fd-9c86-b853cec0e5a4","year":2019},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:33850a00cb5a6c31f0e8e2b083dee2e417c1d1102ecbf58e5573c8497ac590d0","observation_id":"0604200b-a240-460b-be68-bd6c07e2415a","resolution":{"observed_at":"2026-05-15T16:30:37.155831Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"M., and Littwin, E","venue":null,"work_id":"0b6b9438-3298-473f-9e04-8fbaec6b8d89","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:69824e1f34bb50cdb21a117bec7a4cd14d8516a82a0aca2f26e2e7d8be06baf8","observation_id":"1a23b411-a2d0-4dfa-a507-a67dc967c7e5","resolution":{"observed_at":"2026-05-15T16:30:37.159470Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Contrastive multiview coding","venue":null,"work_id":"d126844b-2781-4f2e-9412-046ff05e9cae","year":2020},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:b9d684287bba83629b80adb70c9efd7329d5f6f62923ee919991929a2b0ad1e6","observation_id":"19af1b39-f3e9-4972-9750-bc4d126597db","resolution":{"observed_at":"2026-05-15T16:30:37.162962Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Llama 2: Open foundation and fine-tuned chat models","venue":null,"work_id":"a62fb009-108d-4abb-b5a7-216289885a5d","year":2023},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:e3a7a6acec91251a80b075446bdd62d4167144dbf3f6521bf8f10dfdf68af641","observation_id":"f5603c00-bb20-4ccd-9867-0622f9689474","resolution":{"observed_at":"2026-05-15T16:30:37.166787Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The geometry of hidden representations of large transformer models","venue":null,"work_id":"5ffaf269-0d68-498b-bf27-e23e173edb60","year":2023},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:924d0d3792967f817cb286c9494209456099d6faf6d932fa2144bbf9b0011d30","observation_id":"caba7ef3-27fc-48e7-8762-df6e3efa64d9","resolution":{"observed_at":"2026-05-15T16:30:37.170798Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"N., Kaiser, L","venue":null,"work_id":"88d214e1-106b-48b1-bd61-1f79a9c63da0","year":2017},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:4faccdd6440daa484a9ca2ea5086d721ca8f1e3a051cd8c63f55b5ad3dc13bf6","observation_id":"c2f3d49e-2629-4ffb-b13e-8c38a1a4211e","resolution":{"observed_at":"2026-05-15T16:30:37.174738Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The bottom-up evolution of representations in the transformer: A study with machine translation and language modeling objectives","venue":null,"work_id":"eeb007e8-221c-4fad-ac0b-04bfac6b1082","year":2019},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:8ab47ae425f40373d119da549b3c8dcca43133fd018f5d35c7dac2e0fc91675c","observation_id":"87de45f6-2b2f-43db-a873-2648696fdf3b","resolution":{"observed_at":"2026-05-15T16:30:37.178600Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Diff-eRank : A novel rank-based metric for evaluating large language models","venue":null,"work_id":"68f16ab8-0945-4c1a-b1b9-b6c978ff431e","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:0cd88ea814458e0943e58f47915f4b18642dc1c01f8ff34c317f496f11917131","observation_id":"a91308a8-8f5f-49e3-9e98-1a3a2c64cfb5","resolution":{"observed_at":"2026-05-15T16:30:37.182576Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Efficient streaming language models with attention sinks","venue":null,"work_id":"0c782691-1926-4cb0-aad0-9e3fd0554916","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:a514b0303294f9f6e2929376034e1f6f99f43abf99c6ce4187393494e34d6056","observation_id":"34610626-8a4a-4b19-9758-5209053c7929","resolution":{"observed_at":"2026-05-15T16:30:37.186828Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Qwen2.5 technical report","venue":null,"work_id":"757d574e-9a56-4c73-9039-32f83eef20e9","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:fec5c663c9c4508e4b52cb919ee331a7942c0c30cf0c965dd63f93bdcc0a0dde","observation_id":"f6c7bcff-e9c1-4a4c-a3ce-5516f13039b8","resolution":{"observed_at":"2026-05-15T16:30:37.190734Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"9c87efe5-429d-4215-99f4-63aae1658e03","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:91ac033f72582fcf99b87f591a3d1a0b513136e56d19246ea679d0d6f106423b","observation_id":"9bcfc23b-c81a-40a6-a2bb-9cc52007ecdf","resolution":{"observed_at":"2026-05-15T16:30:37.194627Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"and Liu, D","venue":null,"work_id":"06914c96-da07-4753-8f24-6c1c9fc5b192","year":2021},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":71,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:8436bf56962ec2eefa8576e80381035568bbdfb151da0e885a6a72ba014bac85","observation_id":"38bff56f-f42a-4126-9ed1-089f14fdc8db","resolution":{"observed_at":"2026-05-15T16:30:37.199286Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Székely and Maria L","venue":null,"work_id":"0c4ca528-480d-4917-b886-c9e8501d208f","year":2008},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":72,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:648fc6f52e22e670bc2989ebc99e6df7cf78afd492fb6ec3abbaf8aa097205f5","observation_id":"b09dba28-6f2c-4002-8b4a-c3bd12b7acb1","resolution":{"observed_at":"2026-05-15T16:30:37.203513Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"11bf9f8d-e973-44ee-aba2-a70fa8a99fcb","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":73,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:cd492e070d369abfd023c655ec5e25ab7c689e04ee33ce2294df602fd654694c","observation_id":"91cf7f4e-9a92-4ca4-b307-2f194273bff3","resolution":{"observed_at":"2026-05-15T16:30:37.207855Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"7ece975f-16b3-43a0-b9e5-d64223508e1f","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":74,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:7b581676f32d3cdcbd4756c757142c4639fc7f3025ba3fff9628d7620bcadca7","observation_id":"cfecbd5c-d653-4738-a3cc-c0c65a2a3510","resolution":{"observed_at":"2026-05-15T16:30:37.211970Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"DeepSeek-","venue":null,"work_id":"e0af7f09-32ea-4a6a-8fcf-89929e03ff46","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":75,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:f59ca405b1ad056730c8f21803448a1e0503347eb95f8886b7a41b585869bdb0","observation_id":"3dda413d-260f-4e2a-bd56-0e4c86a6afaf","resolution":{"observed_at":"2026-05-15T16:30:37.216283Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T11:06:11.687652Z","title":"2019 , publisher=","venue":null,"work_id":"ee748159-5fa3-499b-830f-b7a66964826a","year":2019},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":76,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:f43372d82b6226ab4145faccb1a0154bc2af81805752fa6a81152f25b295fe49","observation_id":"659acadb-5401-4045-ada1-9e2bb05bb020","resolution":{"observed_at":"2026-05-15T16:30:37.221420Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"e1e85e35-78a4-4dd4-894a-b28ae83bd4f8","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:cf75bc27eb8d856c0de5d3a32ae71bcc9a029682cd479ec641a44a3bc98dbe45","observation_id":"c6f20dd7-9d05-4ee4-a2bd-45c8ee7b3cf3","resolution":{"observed_at":"2026-05-15T16:30:37.225490Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"8b458378-bb5a-4d97-bda0-b603a780bc75","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:395654f18b1bf2a2f077238383dfe241d9baab73a876876783269f59626a3489","observation_id":"ea628efc-4c54-449e-9646-ad097d7aad3d","resolution":{"observed_at":"2026-05-15T16:30:37.229638Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"eb7430c3-6fde-4b2b-934a-67524f85ba60","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":79,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:8e971d48b3f23e1988aa7b6f64be742dcf36c9338fde2361bc974150a3b39dcb","observation_id":"15090b91-5721-4b8d-88f3-6c2b0b8516fd","resolution":{"observed_at":"2026-05-15T16:30:37.233726Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"e04505fd-b2b5-49a3-9a06-71d8cda9270d","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":80,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:f17baa9b472da9ef3bae1932528f571c1332f1cfc8c0a84aa67912e0e2b14a53","observation_id":"b5e38fad-1d6d-4cf8-815f-1cae2bc8cfcd","resolution":{"observed_at":"2026-05-15T16:30:37.237959Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"8de25df2-b679-488c-bb49-410d95b1a8dd","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":81,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:9b8bca82f3fef7fdae8b174f757fd7b45b03b1ce7340cec8e7c936a913181d62","observation_id":"bb47c1c2-2ec7-42bb-9e40-8a8b65b93619","resolution":{"observed_at":"2026-05-15T16:30:37.241988Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"ICML 2024 Workshop on Mechanistic Interpretability , year=","venue":null,"work_id":"62adf8f6-0cbe-4496-b260-bc5d36b300f0","year":2024},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":82,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:60bdc05d40ec3e3127ed6fc6a2500e6837220ad97441270b9e9f2cbe2c84d8fb","observation_id":"b890acff-5080-46cf-bebe-1d1648830410","resolution":{"observed_at":"2026-05-15T16:30:37.246498Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"2a03577b-8efd-45f2-ac76-b42d0202b56b","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":83,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:89a52fb02f4ec8e383214474bd25a3098f497e43a69651f30065ee1cd6cb2f05","observation_id":"b804707e-7d52-4407-90ac-0f0d7a394d5e","resolution":{"observed_at":"2026-05-15T16:30:37.250819Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"d6b97b7a-9b0e-492f-94e9-daf17cc6c2cf","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":84,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:1601d194de1eaa73af30172baddea386d9b710ea568914352ac89e1e8fd503e0","observation_id":"4d293a3b-2f31-4f60-8e88-ad743bcead1a","resolution":{"observed_at":"2026-05-15T16:30:37.255153Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"66a94348-5cfc-4651-a0de-f7771a4c8dae","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":85,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:111dd591c1a932aa28b2840bef58ba771ed95a756152c9bd0e147a324aa43ff6","observation_id":"ad1d07b4-6467-4dd5-8374-fe7a49ad26ab","resolution":{"observed_at":"2026-05-15T16:30:37.259245Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"c28cf776-ca5c-4ee2-8af8-c2326da163fc","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":86,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:cb990b6331ae1486e9e8622024930873e829805d2f85389cf6cdf5d094e6bc0c","observation_id":"83442881-787b-4e21-9b06-4c40854cc957","resolution":{"observed_at":"2026-05-15T16:30:37.263442Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Not all layers of","venue":null,"work_id":"bccd085c-72cd-4c32-a98b-a10581500f5f","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":87,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:98c8ed5cd0512408b8700a0fcf3f7221416e254e30c3265109bf18524935f36a","observation_id":"dc395cf5-5046-4461-9ac0-1228ebb31071","resolution":{"observed_at":"2026-05-15T16:30:37.267429Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"c04fc24d-a3bb-4b5c-a0e6-2a9475ec2e9d","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":88,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:6ea8c44c34c8063ec7011dd229578782c1e8cc6403c54b2b085b00a00f25e5ec","observation_id":"17ddfbdd-0991-4ade-98cd-1a1ea89daf7a","resolution":{"observed_at":"2026-05-15T16:30:37.271601Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"085bd107-8721-45fb-b371-6f16f45906b3","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":89,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:8c3661368321f6c4b16927e7f4516ff780400b03a90e67051aedb06111435d30","observation_id":"a959920d-1283-4eaf-9c9c-5ee82cb29614","resolution":{"observed_at":"2026-05-15T16:30:37.276451Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"f019e9ac-4a76-413e-859b-4db17d45ddbf","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":90,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:f7b859027b9cee78df38a12567c7058914378972cae9cc3b492e3b27bde698c7","observation_id":"5d6545ce-aeb8-49b0-8d20-f066f485b976","resolution":{"observed_at":"2026-05-15T16:30:37.280445Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"d1266e80-2eb0-49fd-b9ed-e14ff74ada9c","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":91,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:72d88c845cb0b47057eb20682338e052080db2181efa17617a796cbac7a8a257","observation_id":"2575f71e-ffe9-44c0-b8de-165bf50ba2c5","resolution":{"observed_at":"2026-05-15T16:30:37.284617Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"7bdd2688-636c-4f9a-b323-4bebc0c07b02","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":92,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:6df25c7d990fb2da577627826b69f73a8936cbf291adbdaa81a57ebea48232b6","observation_id":"7e82daee-83e8-4752-85e1-ecf8f78bb225","resolution":{"observed_at":"2026-05-15T16:30:37.289156Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"0e6bc65d-8908-4002-8c2d-882e183cbf39","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":93,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:713a70efb76f5ee7a5829336cea8ccd761d53ea3b813c891a7944ddfbcde61af","observation_id":"bfa866ba-8056-4e1a-81de-3c5968a5bf7e","resolution":{"observed_at":"2026-05-15T16:30:37.293082Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"acd0aedf-56ab-4d96-a102-f67d58025254","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":94,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:37abb0a48f64a7fd3226eb16b5a4ca89d01c5abe94e3ee61995a0d82f843b720","observation_id":"bba57824-2be1-4067-bb19-75c09bcf3198","resolution":{"observed_at":"2026-05-15T16:30:37.296941Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"1341fd53-1991-4238-8537-55cd3c153a53","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":95,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:2a8166d61b66ec3e7dd925c527a0f76dec41362560252e9093fb83b116f41ec5","observation_id":"b274d130-d757-48a0-9fac-7875adb06368","resolution":{"observed_at":"2026-05-15T16:30:37.301044Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"ba00ad81-69af-4329-b37b-dda8d0c48f6d","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":96,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:f7ead212886d8d4757272dfaadff5592d66d6a7e70f90122f8c549cb32d49057","observation_id":"a4b82bd5-cff8-489f-91f6-8fd29b3ccffe","resolution":{"observed_at":"2026-05-15T16:30:37.304945Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"61333af9-680c-47f5-8883-4d65bc6ca281","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":97,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:7f1751062ad3a4885e5584b06d9964692bec30a6b74125e71b18308bdc2d015b","observation_id":"037eee68-f33c-44f4-9ff8-133beefc2537","resolution":{"observed_at":"2026-05-15T16:30:37.308876Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Entropy , year=","venue":null,"work_id":"3c52ba43-6568-4ee0-b945-f74770e89eb2","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":98,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:39d7e90dbe3338e8c0b1db48db3975565978fa3076fa7ae411d2ecacbb1db720","observation_id":"75c19a4c-aca2-4fe9-807e-a297821c62bc","resolution":{"observed_at":"2026-05-15T16:30:37.312808Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"c2d9d8da-6616-4d11-b7e9-864980de4557","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":99,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:56a955e22ff2560a0e9f59a0361100ba996d7191f79f7f2ff6d541d8622e9862","observation_id":"c2c4fb74-8e8c-42ac-aa2f-2c7921ee2db9","resolution":{"observed_at":"2026-05-15T16:30:37.316345Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"e3d094cc-5028-4e2b-87ae-f397e0d42f91","year":null},"citing_paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models","version":2},"reference_index":100,"source":"arxiv_source","source_observed_at":"2026-05-15T16:30:36.783894Z"},"links":{"citing_paper":"/paper/2502.02013"},"observation_digest":"sha256:1ff271682c434cd26245c21347dcf12cc61af623dc035bf32dc63c0cf31ac2a3","observation_id":"5ca65376-4fe4-418d-a19c-f2a634260e93","resolution":{"observed_at":"2026-05-15T16:30:37.320148Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2502.02013","last_updated":"2025-06-15T18:27:17Z","latest_version":2,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-14T18:24:15.384302Z","submitted_at":"2025-02-04T05:03:42Z","title":"Layer by Layer: Uncovering Hidden Representations in Language Models"},"reference_resolution":{"displayed":100,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":31,"verified_exact":0,"verified_fuzzy":69},"total_outbound_references":173},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"thesis":"As of 20 August 2026, this Paper Citation Record lists 100 of 173 outbound references and 74 inbound Pith citation observations for arXiv:2502.02013."}