{"as_of":"2026-08-08T11:09:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:103f4a1eb5f309b1159b470b2292cfc12f336087caef78f0aa36381aee9d53b2","coverage":[{"denominator":48,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":48,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T13:30:42.197839Z","state":"measured"},{"denominator":48,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":48,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2505.21670/citation-record","integrity":"/paper/2505.21670/integrity","json":"/paper/2505.21670/citation-record.json","paper":"/paper/2505.21670"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-07T13:30:37.386566Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:37.386566Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:bee81cc337b1f3010900f08a81bb9dbdba994b7494a17f44d64287f64deac974","observation_id":"e6ed680b-e67a-4fd8-b4cd-947689887516","resolution":{"observed_at":"2026-08-07T13:30:37.386566Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.15024","last_updated":"2024-02-09T17:59:40Z","snapshot_observed_at":"2026-07-06T17:21:01.918787Z","submitted_at":"2024-01-26T17:35:45Z","title":"SliceGPT: Compress Large Language Models by Deleting Rows and Columns","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.15024","snapshot_observed_at":"2026-08-07T13:30:37.429572Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:37.429572Z"},"links":{"cited_paper":"/paper/2401.15024","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:64d6f90ce9ce346cc65c4312eda74aa9c581ea1616e60ce961fa88e079aee3ec","observation_id":"17b8b949-f33b-444a-a8df-4c307bcd769a","resolution":{"observed_at":"2026-08-07T13:30:37.429572Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.00456","last_updated":"2024-10-29T11:09:12Z","snapshot_observed_at":"2026-08-05T10:45:45.389337Z","submitted_at":"2024-03-30T19:20:06Z","title":"QuaRot: Outlier-Free 4-Bit Inference in Rotated LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.00456","snapshot_observed_at":"2026-08-07T13:30:37.510058Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:37.510058Z"},"links":{"cited_paper":"/paper/2404.00456","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:4b9110764f5a43b3f345d1394ba2988b65115fa5b4581a71256a549428bd367d","observation_id":"85d87d80-afe1-41ab-8a74-c95073baec25","resolution":{"observed_at":"2026-08-07T13:30:37.510058Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.03463","last_updated":"2025-03-07T15:17:02Z","snapshot_observed_at":"2026-08-08T06:22:42.687499Z","submitted_at":"2024-09-05T12:19:07Z","title":"Massive Activations in Graph Neural Networks: Decoding Attention for Domain-Dependent Interpretability","version":3},"cited_work":{"arxiv_id":"2409.03463","doi":null,"metadata_source":"pith","pith_arxiv_id":"2409.03463","snapshot_observed_at":"2026-08-07T13:30:42.655670Z","title":"Massive Activations in Graph Neural Networks: Decoding Attention for Domain-Dependent Interpretability","venue":"cs.LG","work_id":"79650bab-1cc0-4f5d-94e8-f8e9b8a8f95b","year":2024},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:37.660231Z"},"links":{"cited_paper":"/paper/2409.03463","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:a728ae335fc3823b0f34e31bd4f4320f81aea956b8bc10b198b1085cd8ac16b1","observation_id":"0dbb7886-2355-4d97-b4bb-86d2f44b2a1a","resolution":{"observed_at":"2026-08-07T13:30:42.735465Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.12929","last_updated":"2023-11-09T14:05:51Z","snapshot_observed_at":"2026-07-06T15:45:35.947502Z","submitted_at":"2023-06-22T14:39:04Z","title":"Quantizable Transformers: Removing Outliers by Helping Attention Heads Do Nothing","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.12929","snapshot_observed_at":"2026-08-07T13:30:37.736038Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:37.736038Z"},"links":{"cited_paper":"/paper/2306.12929","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:89917b4ab3150b278ea786ca65a89d0bdec7c60699b5651638993ee9df36335e","observation_id":"0081ecd2-ce4c-4844-b8fd-ca56779c8cfb","resolution":{"observed_at":"2026-08-07T13:30:37.736038Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2005.14165","last_updated":"2020-07-22T19:47:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-05-28T17:29:03Z","title":"Language Models are Few-Shot Learners","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2005.14165","snapshot_observed_at":"2026-08-07T13:30:37.833071Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:37.833071Z"},"links":{"cited_paper":"/paper/2005.14165","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:ac08aadd81becf7edc8bb81cbd0c3280160fe91598907efbee68e0fe4199a39a","observation_id":"3fc8a278-1724-4728-a751-5f1e981f4c2a","resolution":{"observed_at":"2026-08-07T13:30:37.833071Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:44.985378Z","title":null,"venue":null,"work_id":"09c1cb88-51af-4916-887a-ae58eb2146cf","year":2020},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:37.957646Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:c0659d50d2958d93e2b306a6597f10538b0a7c2a1e6ecc56455287217a4cbd05","observation_id":"f19d1f75-ea28-4230-a15b-133faf70cd6d","resolution":{"observed_at":"2026-08-07T13:30:45.014264Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05265","last_updated":"2025-01-27T13:39:25Z","snapshot_observed_at":"2026-07-06T19:29:09.714725Z","submitted_at":"2024-10-07T17:59:35Z","title":"PrefixQuant: Eliminating Outliers by Prefixed Tokens for Large Language Models Quantization","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05265","snapshot_observed_at":"2026-08-07T13:30:38.037110Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:38.037110Z"},"links":{"cited_paper":"/paper/2410.05265","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:26c7df73ec368f580be09fce1ef6ec36d72f53ff12e6a31f5b3b2796deaa5a80","observation_id":"7ea76c44-dc7b-453e-9180-b9900e150b74","resolution":{"observed_at":"2026-08-07T13:30:38.037110Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:44.819738Z","title":null,"venue":null,"work_id":"a2c9bae7-6b34-4d21-b366-3ee839455a61","year":2023},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:38.188425Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:df311f7fbaeccfd057ddde5f814a1b718f7322fbeb700f90ee39d2b3cdb1bd9e","observation_id":"be4f1da1-6583-4feb-adfb-1423b22c13c2","resolution":{"observed_at":"2026-08-07T13:30:44.883104Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:38.290882Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:38.290882Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:f9cc778dd144386cf77e27d898d6c496b6f191ab906dde5fc65ca4f3ca88cf24","observation_id":"cb351dd3-a46e-4299-938d-b05eb51aabe5","resolution":{"observed_at":"2026-08-07T13:30:38.290882Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:44.618112Z","title":null,"venue":null,"work_id":"01f9fd57-934c-4410-adb0-6755e79e172f","year":2022},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:38.383161Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:1f1f9d944951ff04b07df6925097e00b704ea96753ddf9f9f692740d3f0ebd6e","observation_id":"4a015767-3b44-42cc-9c5e-96a24ebf9e3f","resolution":{"observed_at":"2026-08-07T13:30:44.722213Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:38.471287Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:38.471287Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:63e0bdc70dab6741af211ba19cc89299ca65a6e4ec00cace9e0597c590b487f3","observation_id":"7f340ca3-cbda-4ae1-bfe1-d3b6714765d7","resolution":{"observed_at":"2026-08-07T13:30:38.471287Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2210.17323","last_updated":"2023-03-22T13:10:47Z","snapshot_observed_at":"2026-08-07T08:38:54.025062Z","submitted_at":"2022-10-31T13:42:40Z","title":"GPTQ: Accurate Post-Training Quantization for Generative Pre-trained Transformers","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.17323","snapshot_observed_at":"2026-08-07T13:30:38.592880Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:38.592880Z"},"links":{"cited_paper":"/paper/2210.17323","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:c58005f3893f4e39c2ac66f22d4ca57dee0e57fe9c28d7e344097f47eb7283d2","observation_id":"733264aa-dfbc-4f30-a0dc-794647ed8c37","resolution":{"observed_at":"2026-08-07T13:30:38.592880Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:44.423037Z","title":null,"venue":null,"work_id":"70514d5b-ac0a-4bb2-93d2-eabbf2b733df","year":2023},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:38.752863Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:cb545766b8a0fb15cc362d0dc0fcc558d75c3662d9d24d47b1dda475c8ff765c","observation_id":"6bb46df9-a8c1-4b7a-a733-5be3f2947c1f","resolution":{"observed_at":"2026-08-07T13:30:44.505758Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1606.08415","last_updated":"2023-06-06T01:53:32Z","snapshot_observed_at":"2026-07-06T05:01:27.910364Z","submitted_at":"2016-06-27T19:20:40Z","title":"Gaussian Error Linear Units (GELUs)","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1606.08415","snapshot_observed_at":"2026-08-07T13:30:38.822949Z","title":null,"venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:38.822949Z"},"links":{"cited_paper":"/paper/1606.08415","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:ac870f47d0e80a9040f68cc5f182e4114b0a983b601130418122ecb4c7bb1bfb","observation_id":"657d263d-4fc6-4823-9f41-a94867da0ec6","resolution":{"observed_at":"2026-08-07T13:30:38.822949Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:44.252513Z","title":null,"venue":null,"work_id":"d3b1b0cb-738e-4683-a146-6b542789fd34","year":2023},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:38.959316Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:40323858a6cd4f5fb34b8e7a477a62abf8887eb64754c7e3d9f657481dcdf5e9","observation_id":"edbd2b05-8620-447c-9fff-5a5c890b8dda","resolution":{"observed_at":"2026-08-07T13:30:44.345786Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:44.081155Z","title":null,"venue":null,"work_id":"b16a72a3-366d-4739-b09f-7ee0d49d1919","year":2022},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:39.073437Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:7d70abfd89cfda8c86af9d1b72db6def2cce3b97956f0b908a2e777f69319325","observation_id":"1bf91cc1-9cee-4d49-aa66-bd4b4ff76f52","resolution":{"observed_at":"2026-08-07T13:30:44.164593Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:43.893008Z","title":null,"venue":null,"work_id":"7aecfc7c-52e3-4114-ae63-7239136d32c3","year":2023},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:39.135052Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:983184f08062e87c05f78cc476ffad7defb91cf47ac246dd413ae37580acff6a","observation_id":"c5ba07d3-8648-4a90-a99e-d3b50b06ae13","resolution":{"observed_at":"2026-08-07T13:30:43.980860Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.18158","last_updated":"2024-06-06T07:30:44Z","snapshot_observed_at":"2026-08-08T01:30:27.397319Z","submitted_at":"2024-02-28T08:43:05Z","title":"Evaluating Quantized Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.18158","snapshot_observed_at":"2026-08-07T13:30:39.231092Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:39.231092Z"},"links":{"cited_paper":"/paper/2402.18158","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:10d9e6f449a860947ec834fe41b13fb786271c5bb5032706d070c6eedbb0d15b","observation_id":"791978d9-f2a4-4084-8f0a-e4bacc79b3eb","resolution":{"observed_at":"2026-08-07T13:30:39.231092Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2102.05426","last_updated":"2021-07-25T09:34:39Z","snapshot_observed_at":"2026-08-04T09:21:27.169501Z","submitted_at":"2021-02-10T13:46:16Z","title":"BRECQ: Pushing the Limit of Post-Training Quantization by Block Reconstruction","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2102.05426","snapshot_observed_at":"2026-08-07T13:30:39.304149Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:39.304149Z"},"links":{"cited_paper":"/paper/2102.05426","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:190fcdf77b791d709a675f776d76f37dd75871972bd6c468be15d007f9f7cd58","observation_id":"721391e0-4c93-4db4-93f6-c8ccbb095ed7","resolution":{"observed_at":"2026-08-07T13:30:39.304149Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:43.649502Z","title":null,"venue":null,"work_id":"5af73e9c-46db-4b17-a663-4273d369041b","year":2024},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:39.420523Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:0586e20c93fdd132213c005dc2e1f6bbf2a74716a2c1857cb29becdddb66ce99","observation_id":"d2354f38-eefd-412e-990b-35f4ef165c99","resolution":{"observed_at":"2026-08-07T13:30:43.765572Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.04532","last_updated":"2025-05-01T02:14:05Z","snapshot_observed_at":"2026-08-08T04:42:16.973896Z","submitted_at":"2024-05-07T17:59:30Z","title":"QServe: W4A8KV4 Quantization and System Co-design for Efficient LLM Serving","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.04532","snapshot_observed_at":"2026-08-07T13:30:39.500025Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:39.500025Z"},"links":{"cited_paper":"/paper/2405.04532","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:bba335cd7877bfd59fe4986933a5f691166f636d3b9612edb3dcb7d8b4e8ee51","observation_id":"5f0c7948-1620-4434-998d-65957b1f6bc8","resolution":{"observed_at":"2026-08-07T13:30:39.500025Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:43.384511Z","title":null,"venue":null,"work_id":"f8070c57-cd48-4962-868d-51f1a93276f7","year":2021},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:39.571929Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:00a1968728ecf99079d7ddde136f0e0e9b1abe343e483d287fd684b4d66b941f","observation_id":"bc00a249-3b5c-4b2a-9e3f-5f4f19a0a6a9","resolution":{"observed_at":"2026-08-07T13:30:43.483521Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.16406","last_updated":"2025-02-20T06:07:00Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-05-26T02:15:49Z","title":"SpinQuant: LLM quantization with learned rotations","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.16406","snapshot_observed_at":"2026-08-07T13:30:39.624756Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:39.624756Z"},"links":{"cited_paper":"/paper/2405.16406","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:e8044f2734054caf7e98ce5a1d285e90058137dc4083c565873c15d29fe2e4cc","observation_id":"c5fad801-c2e1-4448-846f-bbbc31331af3","resolution":{"observed_at":"2026-08-07T13:30:39.624756Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1609.07843","last_updated":"2016-09-26T04:06:13Z","snapshot_observed_at":"2026-07-06T05:12:10.387914Z","submitted_at":"2016-09-26T04:06:13Z","title":"Pointer Sentinel Mixture Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1609.07843","snapshot_observed_at":"2026-08-07T13:30:39.699449Z","title":null,"venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:39.699449Z"},"links":{"cited_paper":"/paper/1609.07843","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:bd82d99a32625cdce9b9c00f6e97349e764c5b7fad290c2ef3f9055436edccf0","observation_id":"d73f50b4-067f-4960-a769-48379c852457","resolution":{"observed_at":"2026-08-07T13:30:39.699449Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:39.776969Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:39.776969Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:34b09bd10aca1cc6b2e4d95056100999dbb3c5f2f2c50615568fa3509c1c52de","observation_id":"30482704-8f73-4c3b-bef2-9bab6166e3fe","resolution":{"observed_at":"2026-08-07T13:30:39.776969Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:43.142940Z","title":null,"venue":null,"work_id":"9fb81acc-2922-48df-b09b-053f1c361378","year":2024},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:39.842607Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:b5048f797978416c6beecd1199474dd5a21979962cbc6a5256e855fabe559e34","observation_id":"4537b1f7-9d32-4cb9-ba76-04f4bfd0dead","resolution":{"observed_at":"2026-08-07T13:30:43.259232Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:39.966935Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:39.966935Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:d17268e4108d10d71d6f7ef5bdf9c066259ceb206b020d5ba7428f04e8ff366e","observation_id":"5ee79d5b-d4ad-4daf-88a5-469b1e4ce3cd","resolution":{"observed_at":"2026-08-07T13:30:39.966935Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:40.033619Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:40.033619Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:d1b4c166e2729bc79c8fc2052994297ba1b74d09c457f1579ea82a115b7db928","observation_id":"753928f7-3051-4627-8dca-627c32237d6e","resolution":{"observed_at":"2026-08-07T13:30:40.033619Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.13137","last_updated":"2024-03-18T05:33:22Z","snapshot_observed_at":"2026-07-06T16:10:15.694898Z","submitted_at":"2023-08-25T02:28:35Z","title":"OmniQuant: Omnidirectionally Calibrated Quantization for Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.13137","snapshot_observed_at":"2026-08-07T13:30:40.146040Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:40.146040Z"},"links":{"cited_paper":"/paper/2308.13137","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:99075b589fb894a2f2ffd9b7f772ea2756df36122b0dffd5ad8e34f5a3cf664c","observation_id":"93816b99-e4bb-4df5-ae25-25ab0721a35c","resolution":{"observed_at":"2026-08-07T13:30:40.146040Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.17762","last_updated":"2024-08-14T16:00:49Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-02-27T18:55:17Z","title":"Massive Activations in Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.17762","snapshot_observed_at":"2026-08-07T13:30:40.241455Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:40.241455Z"},"links":{"cited_paper":"/paper/2402.17762","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:bf8e4f01e47bb08959c1f438a46d6ea3bc9539950c87fc7eb70fe59d0b18404e","observation_id":"51c67f3b-c46e-4aa9-80f0-31229390a4ca","resolution":{"observed_at":"2026-08-07T13:30:40.241455Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:40.350230Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:40.350230Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:90375d6e23a4b4cc4ca68cb20125572ac8174df60563d6a63cf7e9c138a9353e","observation_id":"ab631e27-4636-4296-aa47-f72901efcf97","resolution":{"observed_at":"2026-08-07T13:30:40.350230Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2302.13971","last_updated":"2023-02-27T17:11:15Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-02-27T17:11:15Z","title":"LLaMA: Open and Efficient Foundation Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.13971","snapshot_observed_at":"2026-08-07T13:30:40.441902Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:40.441902Z"},"links":{"cited_paper":"/paper/2302.13971","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:02074e43eeb6bdea742aa21be153e4a4fd3fbe88cde525d92ab0431076acceb1","observation_id":"40f8c7b3-6f28-4961-92f7-cff657164557","resolution":{"observed_at":"2026-08-07T13:30:40.441902Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.09288","last_updated":"2023-07-19T17:08:59Z","snapshot_observed_at":"2026-08-07T12:56:43.323460Z","submitted_at":"2023-07-18T14:31:57Z","title":"Llama 2: Open Foundation and Fine-Tuned Chat Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.09288","snapshot_observed_at":"2026-08-07T13:30:40.588045Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:40.588045Z"},"links":{"cited_paper":"/paper/2307.09288","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:cd5142029d0b12803e26b30f4ef7117ab968a62760b6e190f9b4351a940d81ad","observation_id":"133e7ba1-ecf8-4350-9838-b9ea9e132f30","resolution":{"observed_at":"2026-08-07T13:30:40.588045Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.04396","last_updated":"2024-06-04T04:51:52Z","snapshot_observed_at":"2026-07-06T17:26:31.326736Z","submitted_at":"2024-02-06T20:52:12Z","title":"QuIP#: Even Better LLM Quantization with Hadamard Incoherence and Lattice Codebooks","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.04396","snapshot_observed_at":"2026-08-07T13:30:40.665688Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:40.665688Z"},"links":{"cited_paper":"/paper/2402.04396","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:518b4077a91126e7a6e7bbf4fb75a22ab3a91ce19b1155cb48f12473c42e44d7","observation_id":"8f1435f5-aada-4861-9453-291e952d78cf","resolution":{"observed_at":"2026-08-07T13:30:40.665688Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:42.894699Z","title":null,"venue":null,"work_id":"dc7b2ad0-380b-4c51-9ccb-c2dc81d8fc14","year":2024},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:40.761642Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:0d96ca84debb8db7d913d049f2d255dc88eaa5955c3a521813321d8cd9c85d84","observation_id":"05f19d40-6e52-4015-b2d1-74d64bda934e","resolution":{"observed_at":"2026-08-07T13:30:43.011422Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.09145","last_updated":"2023-10-23T08:48:31Z","snapshot_observed_at":"2026-07-06T15:17:05.623407Z","submitted_at":"2023-04-18T17:34:23Z","title":"Outlier Suppression+: Accurate quantization of large language models by equivalent and optimal shifting and scaling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.09145","snapshot_observed_at":"2026-08-07T13:30:40.860196Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:40.860196Z"},"links":{"cited_paper":"/paper/2304.09145","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:5f491972433ff28ffd0e73909a84359412632ce062e25093975f272df4187e3e","observation_id":"0af6810a-da2d-419c-840e-3bac0083962f","resolution":{"observed_at":"2026-08-07T13:30:40.860196Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:41.004301Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:41.004301Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:bffdedc8349a4ed5bb92ae0be788a69275aabfafb4d53d25b1a52ad7d1bda529","observation_id":"b615385d-56e1-4596-b6f8-4776bad7f24b","resolution":{"observed_at":"2026-08-07T13:30:41.004301Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:41.125255Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:41.125255Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:c9c72afe6922a041ad1f0e4220d1ca3bc0926444566a8ed87ede6c83bc00e622","observation_id":"3b479df9-e616-401c-ae20-6bc47c437e6b","resolution":{"observed_at":"2026-08-07T13:30:41.125255Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08417","last_updated":"2024-06-03T01:28:06Z","snapshot_observed_at":"2026-07-06T17:16:13.055978Z","submitted_at":"2024-01-16T15:04:51Z","title":"Contrastive Preference Optimization: Pushing the Boundaries of LLM Performance in Machine Translation","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.08417","snapshot_observed_at":"2026-08-07T13:30:41.250316Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:41.250316Z"},"links":{"cited_paper":"/paper/2401.08417","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:67db3b190384d4f78ce212c916805e103a459df509d62066860d1ced1461c3aa","observation_id":"6cb40ea7-ff7c-4ca5-aa3a-25c69d8e4f62","resolution":{"observed_at":"2026-08-07T13:30:41.250316Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10671","last_updated":"2024-09-10T13:25:53Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-07-15T12:35:42Z","title":"Qwen2 Technical Report","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.10671","snapshot_observed_at":"2026-08-07T13:30:41.364967Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:41.364967Z"},"links":{"cited_paper":"/paper/2407.10671","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:5a2bad58db67cd94314f6448a36931f98bfcb14698df27d6d547cef999856e77","observation_id":"4c8326f7-a920-418f-9313-60e63c196149","resolution":{"observed_at":"2026-08-07T13:30:41.364967Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:41.501093Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:41.501093Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:ec916030f1ea84c6d9adfa07f8ca7ba8d66bba0c416ea4f8e0381b841c327052","observation_id":"9f9a79d6-b3ef-4194-b8bd-5fe6db3232bd","resolution":{"observed_at":"2026-08-07T13:30:41.501093Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2210.02414","last_updated":"2023-10-25T05:22:43Z","snapshot_observed_at":"2026-08-02T04:26:53.194797Z","submitted_at":"2022-10-05T17:34:44Z","title":"GLM-130B: An Open Bilingual Pre-trained Model","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.02414","snapshot_observed_at":"2026-08-07T13:30:41.624035Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:41.624035Z"},"links":{"cited_paper":"/paper/2210.02414","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:9c3ae3fa3d3e72b6b997c1cac4461082782ae8790bc282145d9d6198472d2dc6","observation_id":"3e155d3f-7401-4f72-a7e2-06c1ef6f43bd","resolution":{"observed_at":"2026-08-07T13:30:41.624035Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.02069","last_updated":"2025-05-15T17:18:12Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-04T07:51:30Z","title":"PyramidKV: Dynamic KV Cache Compression based on Pyramidal Information Funneling","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.02069","snapshot_observed_at":"2026-08-07T13:30:41.724183Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:41.724183Z"},"links":{"cited_paper":"/paper/2406.02069","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:10cb53181ad944029e1e9eb8710dd57f84998475e21cc24af4918357fe4612ad","observation_id":"9c1a0824-00ec-4abe-833e-fb320631cf29","resolution":{"observed_at":"2026-08-07T13:30:41.724183Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.09904","last_updated":"2024-07-31T07:35:42Z","snapshot_observed_at":"2026-08-08T03:10:49.991472Z","submitted_at":"2024-06-14T10:23:45Z","title":"QQQ: Quality Quattuor-Bit Quantization for Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.09904","snapshot_observed_at":"2026-08-07T13:30:41.859346Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:41.859346Z"},"links":{"cited_paper":"/paper/2406.09904","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:e29975d528512c96e0a39b2e1cec7bba9cbe6b73dca66ab1ad45e8440e26a192","observation_id":"290d779b-6d71-411c-bd9d-2c5226e91831","resolution":{"observed_at":"2026-08-07T13:30:41.859346Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.04675","last_updated":"2024-06-14T11:40:52Z","snapshot_observed_at":"2026-07-06T15:14:02.615465Z","submitted_at":"2023-04-10T15:51:30Z","title":"Multilingual Machine Translation with Large Language Models: Empirical Results and Analysis","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.04675","snapshot_observed_at":"2026-08-07T13:30:42.000354Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:42.000354Z"},"links":{"cited_paper":"/paper/2304.04675","citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:0282b5b717d75bca5b0d2c8936d60cc5c775d8c9cf29620adc52b6d200b99db5","observation_id":"52f12420-49ba-487d-8944-b2ccf3f3fdca","resolution":{"observed_at":"2026-08-07T13:30:42.000354Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:42.088259Z","title":"online\" 'onlinestring :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:42.088259Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:766a46a0077d2c65280a28a25d66b8051b1426ae0be1ba68b97c5600a566daee","observation_id":"a1d8adb0-a0a9-40ec-be3d-bedb3b55e3b3","resolution":{"observed_at":"2026-08-07T13:30:42.088259Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T13:30:42.197839Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-07T13:30:42.197839Z"},"links":{"citing_paper":"/paper/2505.21670"},"observation_digest":"sha256:7dbfff5bf99fc116fe2e3eff6fbb57dd84766592e546d1f2ce0835133d3f5c22","observation_id":"fbfe6986-82e6-43af-b607-2e4efc490ef2","resolution":{"observed_at":"2026-08-07T13:30:42.197839Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2505.21670","last_updated":"2025-05-27T18:48:40Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-07T23:42:58.100393Z","submitted_at":"2025-05-27T18:48:40Z","title":"Rethinking the Outlier Distribution in Large Language Models: An In-depth Study"},"reference_resolution":{"displayed":48,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":47,"verified_exact":1,"verified_fuzzy":0},"total_outbound_references":48},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 48 of 48 outbound references and 0 inbound Pith citation observations for arXiv:2505.21670."}