{"as_of":"2026-08-18T09:19:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:f579ef80463b600fafe035b153cbf04ea1de926fd5329eb20f5c347ac396acef","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":42,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":42,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-18T06:34:40.430872+00:00","state":"measured"},{"denominator":42,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":42,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-16T05:09:55.031898Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"pith","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":5,"observed_at":"2026-08-05T02:28:24.338817Z","source":"pith"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-12T19:33:01.033890Z","title":"Powerinfer-2: Fast large lan- guage model inference on a smartphone","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.10640","last_updated":"2024-11-16T00:14:51Z","snapshot_observed_at":"2026-08-17T21:54:44.695295Z","submitted_at":"2024-11-16T00:14:51Z","title":"BlueLM-V-3B: Algorithm and System Co-Design for Multimodal Large Language Models on Mobile Devices","version":1},"reference_index":132,"source":"pdf_text","source_observed_at":"2026-08-12T19:33:01.033890Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2411.10640"},"observation_digest":"sha256:e0b4fb2a88a7bf4e22213d2633d72a1723e059c686a183a092eacd814621ce8e","observation_id":"456d174e-001a-4ade-b2fd-533489973d47","resolution":{"observed_at":"2026-08-12T19:33:01.033890Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-12T04:30:36.119613Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.01380","last_updated":"2025-04-03T13:28:51Z","snapshot_observed_at":"2026-08-15T11:50:56.372845Z","submitted_at":"2024-12-02T11:07:51Z","title":"Efficient LLM Inference using Dynamic Input Pruning and Cache-Aware Masking","version":2},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-12T04:30:36.119613Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2412.01380"},"observation_digest":"sha256:5c7d43c71239b34576cf2481a27930de5c1abd5495a483d47802256a4ca12fc1","observation_id":"8c7d2330-15c8-4845-b7b3-f9906a84d1b9","resolution":{"observed_at":"2026-08-12T04:30:36.119613Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-11T21:40:05.602148Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.04315","last_updated":"2024-12-06T11:39:27Z","snapshot_observed_at":"2026-08-14T20:58:54.206883Z","submitted_at":"2024-12-05T16:31:13Z","title":"Densing Law of LLMs","version":2},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-11T21:40:05.602148Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2412.04315"},"observation_digest":"sha256:c7410b1c941d2ec130bf0b77360c21cbbcb746fbf597c34d2ac1e0968e4ee460","observation_id":"f0aac06a-5260-43f4-bdb5-24caff9e8358","resolution":{"observed_at":"2026-08-11T21:40:05.602148Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-10T16:59:01.428571Z","title":"Powerinfer-2: Fast large language model inference on a smart- phone","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.12689","last_updated":"2025-09-04T06:20:55Z","snapshot_observed_at":"2026-08-16T15:43:30.927019Z","submitted_at":"2025-01-22T07:52:38Z","title":"IC-Cache: Efficient Large Language Model Serving via In-context Caching","version":3},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-10T16:59:01.428571Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2501.12689"},"observation_digest":"sha256:b71bfc0bbc35397f160c06699362af4c9a3cb176d839869298e2436dacab750a","observation_id":"7c569c36-6cf0-4a3b-8347-14970b0e2415","resolution":{"observed_at":"2026-08-10T16:59:01.428571Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-10T21:37:14.534676Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.13111","last_updated":"2025-01-08T14:38:13Z","snapshot_observed_at":"2026-08-14T08:24:24.408165Z","submitted_at":"2025-01-08T14:38:13Z","title":"iServe: An Intent-based Serving System for LLMs","version":1},"reference_index":120,"source":"pdf_text","source_observed_at":"2026-08-10T21:37:14.534676Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2501.13111"},"observation_digest":"sha256:50975a6112b100636cad25a950af42e16557935ee8ed006a71f1be531e2d6b8f","observation_id":"73483b14-f398-4fdc-ae5e-94611d77db4d","resolution":{"observed_at":"2026-08-10T21:37:14.534676Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-16T04:51:31.537714Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.00232","last_updated":"2025-05-01T00:44:13Z","snapshot_observed_at":"2026-08-16T04:45:07.395209Z","submitted_at":"2025-05-01T00:44:13Z","title":"Scaling On-Device GPU Inference for Large Generative Models","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-16T04:51:31.537714Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2505.00232"},"observation_digest":"sha256:95b1cfa2810d11b3307136897722dbc3415b82b5c72d75577826f57fd2b34d3c","observation_id":"339db099-dd19-46ef-be92-964ead0a53a3","resolution":{"observed_at":"2026-08-16T04:51:31.537714Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-16T05:09:55.031898Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.00745","last_updated":"2025-04-30T08:08:15Z","snapshot_observed_at":"2026-08-17T08:28:56.704997Z","submitted_at":"2025-04-30T08:08:15Z","title":"Responsive DNN Adaptation for Video Analytics against Environment Shift via Hierarchical Mobile-Cloud Collaborations","version":1},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-08-16T05:09:55.031898Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2505.00745"},"observation_digest":"sha256:8303e77cb4a33d9eabf8c438fdb1bfa995aa82b33caffa2b08d919fd29563b9e","observation_id":"4b03021e-085b-46b0-b834-2a12b430ecde","resolution":{"observed_at":"2026-08-16T05:09:55.031898Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":"2406.06282","doi":"10.48550/arxiv.2406.06282","metadata_source":"pith","pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":"cs.LG","work_id":"c835ae36-bf61-4365-8156-febeba198d32","year":2024},"citing_paper":{"arxiv_id":"2505.02922","last_updated":"2026-04-27T10:13:35Z","snapshot_observed_at":"2026-08-15T10:34:21.616923Z","submitted_at":"2025-05-05T18:01:17Z","title":"RetroInfer: A Vector Storage Engine for Scalable Long-Context LLM Inference","version":3},"reference_index":111,"source":"pdf_text","source_observed_at":"2026-05-22T15:59:04.724780Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2505.02922"},"observation_digest":"sha256:8b3834adea5df082e57eb108f99e46517330db51201d911877827e6351e93886","observation_id":"46a15fe9-ec8e-41b3-98e1-6d898f5b7c49","resolution":{"observed_at":"2026-05-09T06:35:24.624547Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-15T23:00:18.153988Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.05950","last_updated":"2025-05-12T03:29:12Z","snapshot_observed_at":"2026-08-16T10:42:50.887857Z","submitted_at":"2025-05-09T10:53:47Z","title":"FloE: On-the-Fly MoE Inference on Memory-constrained GPU","version":2},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-08-15T23:00:18.153988Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2505.05950"},"observation_digest":"sha256:dc60e8c4cc41247252bcd718494b87448210ade8ea93743754da6b59392b7b0c","observation_id":"fc9a2071-2c3c-4899-ad54-d6cbf007e740","resolution":{"observed_at":"2026-08-15T23:00:18.153988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-07T14:26:17.250081Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.19235","last_updated":"2025-05-25T17:16:34Z","snapshot_observed_at":"2026-08-15T10:20:44.721421Z","submitted_at":"2025-05-25T17:16:34Z","title":"CoreMatching: A Co-adaptive Sparse Inference Framework with Token and Neuron Pruning for Comprehensive Acceleration of Vision-Language Models","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-07T14:26:17.250081Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2505.19235"},"observation_digest":"sha256:723fba61959a7b7137503ecdcd2b0b8bfecfedd395109724e4db5da7a28b9aa6","observation_id":"0f5c1124-bece-4e9b-a235-fd754a1aaa5e","resolution":{"observed_at":"2026-08-07T14:26:17.250081Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-07T13:08:01.749848Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.22735","last_updated":"2025-05-28T18:00:24Z","snapshot_observed_at":"2026-08-15T23:20:36.799155Z","submitted_at":"2025-05-28T18:00:24Z","title":"TensorShield: Safeguarding On-Device Inference by Shielding Critical DNN Tensors with TEE","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-07T13:08:01.749848Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2505.22735"},"observation_digest":"sha256:28e18c654dbaec54682dee9688d214be229da58860651179e6a85e1e2bfcd112","observation_id":"9f3dcdbe-99c7-4bea-9254-fdd19144a1cb","resolution":{"observed_at":"2026-08-07T13:08:01.749848Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":"2406.06282","doi":"10.48550/arxiv.2406.06282","metadata_source":"pith","pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":"cs.LG","work_id":"c835ae36-bf61-4365-8156-febeba198d32","year":2024},"citing_paper":{"arxiv_id":"2506.02153","last_updated":"2025-09-15T22:15:00Z","snapshot_observed_at":"2026-08-12T13:48:47.535846Z","submitted_at":"2025-06-02T18:35:16Z","title":"Small Language Models are the Future of Agentic AI","version":2},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-05-16T11:55:50.897500Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2506.02153"},"observation_digest":"sha256:af7d4f0edf175479e668985639d818a27c9bb6f83f87fdee6f98d3bad3fd4f04","observation_id":"6491a2a5-891d-4668-ac7c-12812760b38b","resolution":{"observed_at":"2026-05-16T11:55:51.043385Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-07T05:37:52.808116Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07533","last_updated":"2025-06-09T08:16:24Z","snapshot_observed_at":"2026-08-13T05:15:55.937681Z","submitted_at":"2025-06-09T08:16:24Z","title":"MoQAE: Mixed-Precision Quantization for Long-Context LLM Inference via Mixture of Quantization-Aware Experts","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-07T05:37:52.808116Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2506.07533"},"observation_digest":"sha256:940a3c13181f939bf96f8fb1d6a9e774f7bd34362dbf766e659af5b23d0d197c","observation_id":"feb5f522-870d-4ed0-a219-fdf23f4847e2","resolution":{"observed_at":"2026-08-07T05:37:52.808116Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-07T04:59:18.087086Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.09351","last_updated":"2025-06-11T03:05:24Z","snapshot_observed_at":"2026-08-13T18:32:34.334743Z","submitted_at":"2025-06-11T03:05:24Z","title":"DIVE into MoE: Diversity-Enhanced Reconstruction of Large Language Models from Dense into Mixture-of-Experts","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-07T04:59:18.087086Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2506.09351"},"observation_digest":"sha256:d231b71bb3b3e387deaa6851995755cee045d7b7a7f20b0c18434eeaa182f3d1","observation_id":"6af74b8a-00e2-4fe4-9dec-a48a499e844f","resolution":{"observed_at":"2026-08-07T04:59:18.087086Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-07T04:33:41.493324Z","title":"arXiv:2406.06282 [cs.LG] https://arxiv.org/abs/2406.06282","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.10443","last_updated":"2025-06-12T07:45:29Z","snapshot_observed_at":"2026-08-15T11:43:15.940792Z","submitted_at":"2025-06-12T07:45:29Z","title":"MNN-LLM: A Generic Inference Engine for Fast Large Language Model Deployment on Mobile Devices","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-07T04:33:41.493324Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2506.10443"},"observation_digest":"sha256:afab3949b6af706195bdf8d50450c2b464fb5cd674c14315f98a406d5d3d8dad","observation_id":"dfe72331-8322-4198-ad01-cb3db0095e9a","resolution":{"observed_at":"2026-08-07T04:33:41.493324Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-15T19:30:31.172365Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.16500","last_updated":"2025-06-19T17:53:34Z","snapshot_observed_at":"2026-08-15T19:22:40.879212Z","submitted_at":"2025-06-19T17:53:34Z","title":"SparseLoRA: Accelerating LLM Fine-Tuning with Contextual Sparsity","version":1},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-08-15T19:30:31.172365Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2506.16500"},"observation_digest":"sha256:d623a96611232aa2c19a45d50ed98989c12e03e3c4c590f803ac2d8ed4182ca8","observation_id":"04ba6675-87b1-4c23-b053-d03152eb5e06","resolution":{"observed_at":"2026-08-15T19:30:31.172365Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-15T18:40:01.498224Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.19884","last_updated":"2025-06-24T04:50:28Z","snapshot_observed_at":"2026-08-16T17:50:55.734332Z","submitted_at":"2025-06-24T04:50:28Z","title":"MNN-AECS: Energy Optimization for LLM Decoding on Mobile Devices via Adaptive Core Selection","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-15T18:40:01.498224Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2506.19884"},"observation_digest":"sha256:d4be76f1d21431f8191f4b076ad3345af6818165e26894e70e371caf9c99fba6","observation_id":"53e72b54-46d4-42fd-8b28-627f600ef4c0","resolution":{"observed_at":"2026-08-15T18:40:01.498224Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-06T20:44:17.762246Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.02135","last_updated":"2025-07-02T20:47:40Z","snapshot_observed_at":"2026-08-17T18:58:14.704944Z","submitted_at":"2025-07-02T20:47:40Z","title":"Dissecting the Impact of Mobile DVFS Governors on LLM Inference Performance and Energy Efficiency","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T20:44:17.762246Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2507.02135"},"observation_digest":"sha256:0594c916439108611875db5e03a1419cca9feb7e1ba8b4dad4ab4ebef3c8abca","observation_id":"53927732-1194-40e3-ba29-008579fe1cf7","resolution":{"observed_at":"2026-08-06T20:44:17.762246Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-06T18:20:53.546644Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.08505","last_updated":"2025-07-14T08:25:14Z","snapshot_observed_at":"2026-08-18T08:58:00.445501Z","submitted_at":"2025-07-11T11:30:57Z","title":"Efficient Deployment of Vision-Language Models on Mobile Devices: A Case Study on OnePlus 13R","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T18:20:53.546644Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2507.08505"},"observation_digest":"sha256:1f06eba42b061869f0557acb6030af176e11f40c71dbefcfd2ae29f766bcf4f4","observation_id":"0b428182-fc7c-427b-a729-c682db3579a5","resolution":{"observed_at":"2026-08-06T18:20:53.546644Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-06T18:20:00.457313Z","title":"PowerInfer-2 : Fast large language model inference on a smartphone","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.08771","last_updated":"2025-07-30T04:14:15Z","snapshot_observed_at":"2026-08-15T02:00:29.298537Z","submitted_at":"2025-07-11T17:28:56Z","title":"BlockFFN: Towards End-Side Acceleration-Friendly Mixture-of-Experts with Chunk-Level Activation Sparsity","version":2},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-06T18:20:00.457313Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2507.08771"},"observation_digest":"sha256:6e3f84876271151b9fd0d9036aeb546acfdab14bbcd127fab2599bcdae793a23","observation_id":"c3bd469b-6ec3-40fb-a33e-85f97cc595f6","resolution":{"observed_at":"2026-08-06T18:20:00.457313Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-06T05:36:21.155515Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.01540","last_updated":"2025-08-03T01:49:08Z","snapshot_observed_at":"2026-08-15T08:02:02.409123Z","submitted_at":"2025-08-03T01:49:08Z","title":"MagicVL-2B: Empowering Vision-Language Models on Mobile Devices with Lightweight Visual Encoders via Curriculum Learning","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-08-06T05:36:21.155515Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2508.01540"},"observation_digest":"sha256:2d4d3693db5cd89c4b2dbc23aee9c62498a3a0e266d4a6a7fe6511832921b97b","observation_id":"dbb96f30-34c4-438f-addc-dc9902738ff9","resolution":{"observed_at":"2026-08-06T05:36:21.155515Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":"2406.06282","doi":"10.48550/arxiv.2406.06282","metadata_source":"pith","pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":"cs.LG","work_id":"c835ae36-bf61-4365-8156-febeba198d32","year":2024},"citing_paper":{"arxiv_id":"2508.16703","last_updated":"2026-04-08T06:56:34Z","snapshot_observed_at":"2026-08-12T14:34:22.889703Z","submitted_at":"2025-08-22T07:41:35Z","title":"ShadowNPU: System and Algorithm Co-design for NPU-Centric On-Device LLM Inference","version":4},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-05-18T22:03:10.316005Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2508.16703"},"observation_digest":"sha256:d00175d9f7c1ab87b3e5f00cb9449ef629da4a409f29f7bb15faa8d465523371","observation_id":"ab19481e-dffc-40bc-96bd-7084443c15c4","resolution":{"observed_at":"2026-05-18T22:06:52.141041Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":"2406.06282","doi":"10.48550/arxiv.2406.06282","metadata_source":"pith","pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":"cs.LG","work_id":"c835ae36-bf61-4365-8156-febeba198d32","year":2024},"citing_paper":{"arxiv_id":"2601.11652","last_updated":"2026-04-07T01:02:52Z","snapshot_observed_at":"2026-08-11T14:14:01.400167Z","submitted_at":"2026-01-15T16:46:01Z","title":"WISP: Waste- and Interference-Suppressed Distributed Speculative LLM Serving at the Edge via Dynamic Drafting and SLO-Aware Batching","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-16T14:12:06.034679Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2601.11652"},"observation_digest":"sha256:1cf14f730cfb47233aa9e5cc0ba1fbd2604e29218cd014c61088afbb866d34c8","observation_id":"8ad5e15e-bae3-48a4-968e-90f9a3b25f4d","resolution":{"observed_at":"2026-05-16T14:12:58.650668Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":"2406.06282","doi":"10.48550/arxiv.2406.06282","metadata_source":"pith","pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":"cs.LG","work_id":"c835ae36-bf61-4365-8156-febeba198d32","year":2024},"citing_paper":{"arxiv_id":"2604.05571","last_updated":"2026-04-07T08:13:45Z","snapshot_observed_at":"2026-08-17T20:31:40.310305Z","submitted_at":"2026-04-07T08:13:45Z","title":"Understanding User Privacy Perceptions of GenAI Smartphones","version":1},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-05-10T19:18:29.955515Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2604.05571"},"observation_digest":"sha256:8081c6eefa2a6ce7eef3dc510427f32ac9f0b98fe66c7bfee9a8cf32be700021","observation_id":"61e88517-3136-4dab-8730-5c9978488b2c","resolution":{"observed_at":"2026-05-10T23:10:51.598346Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":"2406.06282","doi":"10.48550/arxiv.2406.06282","metadata_source":"pith","pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":"cs.LG","work_id":"c835ae36-bf61-4365-8156-febeba198d32","year":2024},"citing_paper":{"arxiv_id":"2604.21026","last_updated":"2026-04-24T07:54:05Z","snapshot_observed_at":"2026-08-11T00:28:28.451528Z","submitted_at":"2026-04-22T19:18:30Z","title":"MCAP: Deployment-Time Layer Profiling for Memory-Constrained LLM Inference","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-10T00:54:33.897112Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2604.21026"},"observation_digest":"sha256:84cf219619292c99b8899514aedc3aba8702527c78e5a6acb4ac70f38c8d4146","observation_id":"a3e0deb1-2816-4836-8266-7be16572cec3","resolution":{"observed_at":"2026-05-10T00:54:48.401477Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":"2406.06282","doi":"10.48550/arxiv.2406.06282","metadata_source":"pith","pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":"cs.LG","work_id":"c835ae36-bf61-4365-8156-febeba198d32","year":2024},"citing_paper":{"arxiv_id":"2604.27476","last_updated":"2026-06-08T07:15:10Z","snapshot_observed_at":"2026-08-12T20:11:33.502007Z","submitted_at":"2026-04-30T06:18:50Z","title":"EdgeFM: Efficient Edge Inference for Vision-Language Models","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-07T08:30:35.264804Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2604.27476"},"observation_digest":"sha256:530605b013d674fd8159ed1d05b165d4792efaed3d092de9585bd029de1aaeaf","observation_id":"ac7eb187-db2c-43fb-85a5-6fe3d9f45801","resolution":{"observed_at":"2026-05-12T10:01:27.717820Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":"2406.06282","doi":"10.48550/arxiv.2406.06282","metadata_source":"pith","pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":"cs.LG","work_id":"c835ae36-bf61-4365-8156-febeba198d32","year":2024},"citing_paper":{"arxiv_id":"2604.27476","last_updated":"2026-06-08T07:15:10Z","snapshot_observed_at":"2026-08-12T20:11:33.502007Z","submitted_at":"2026-04-30T06:18:50Z","title":"EdgeFM: Efficient Edge Inference for Vision-Language Models","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-07-01T08:54:41.845820Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2604.27476"},"observation_digest":"sha256:6415822b8df329b2e0bd350f81112ff94d9e78837bb0b94f8e63e0e3dd7d0caa","observation_id":"5caa659d-b24b-4210-823e-f444f7ed121d","resolution":{"observed_at":"2026-07-01T08:55:34.742443Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":"2406.06282","doi":"10.48550/arxiv.2406.06282","metadata_source":"pith","pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":"cs.LG","work_id":"c835ae36-bf61-4365-8156-febeba198d32","year":2024},"citing_paper":{"arxiv_id":"2605.16786","last_updated":"2026-05-16T03:43:10Z","snapshot_observed_at":"2026-08-16T01:55:32.674428Z","submitted_at":"2026-05-16T03:43:10Z","title":"Lever: Speculative LLM Inference on Smartphones","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-19T21:44:11.735963Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2605.16786"},"observation_digest":"sha256:3518052b2f23c7ec92066dd3038e93c9464679e4abe509a48d5e1da444b78cd4","observation_id":"1d99e954-1e6b-495a-b3a5-5857d15d1b5d","resolution":{"observed_at":"2026-05-19T21:47:48.483707Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":"2406.06282","doi":"10.48550/arxiv.2406.06282","metadata_source":"pith","pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":"cs.LG","work_id":"c835ae36-bf61-4365-8156-febeba198d32","year":2024},"citing_paper":{"arxiv_id":"2605.27435","last_updated":"2026-05-22T10:39:35Z","snapshot_observed_at":"2026-08-08T01:27:25.329843Z","submitted_at":"2026-05-22T10:39:35Z","title":"When NPUs Are Not Always Faster: A Stage-Level Analysis of Mobile LLM Inference","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-30T15:03:31.289211Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2605.27435"},"observation_digest":"sha256:bbe08803c65a76e08691c223a39731ac42300f12461c88a37f2a5875d9510004","observation_id":"e6e60fec-694c-40db-a75a-e7b687591a5c","resolution":{"observed_at":"2026-06-30T15:04:46.182427Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":"2406.06282","doi":"10.48550/arxiv.2406.06282","metadata_source":"pith","pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":"cs.LG","work_id":"c835ae36-bf61-4365-8156-febeba198d32","year":2024},"citing_paper":{"arxiv_id":"2606.11357","last_updated":"2026-07-01T18:57:44Z","snapshot_observed_at":"2026-08-07T05:07:04.133282Z","submitted_at":"2026-06-09T18:33:14Z","title":"TileFuse: A Fused Mixed-Precision Kernel Library for Efficient Quantized LLM Inference on AMD NPUs","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-06-27T11:26:56.230283Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2606.11357"},"observation_digest":"sha256:5bad027371ca2ff9700b827d6db541e268fff59493589a909f2fe253964c331f","observation_id":"af2b36e5-e4b7-4408-bc44-cfb1e0fbb080","resolution":{"observed_at":"2026-07-03T07:57:44.941264Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":"2406.06282","doi":"10.48550/arxiv.2406.06282","metadata_source":"pith","pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":"cs.LG","work_id":"c835ae36-bf61-4365-8156-febeba198d32","year":2024},"citing_paper":{"arxiv_id":"2606.11357","last_updated":"2026-07-01T18:57:44Z","snapshot_observed_at":"2026-08-07T05:07:04.133282Z","submitted_at":"2026-06-09T18:33:14Z","title":"TileFuse: A Fused Mixed-Precision Kernel Library for Efficient Quantized LLM Inference on AMD NPUs","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-07-03T23:37:25.391853Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2606.11357"},"observation_digest":"sha256:030eb33dc360dfe514008ddd7608bbd62d254ddeaef42804e2533e14a9a38d01","observation_id":"1a444fdc-ed0e-44c7-a1b4-b1cc05198f42","resolution":{"observed_at":"2026-07-03T23:39:02.905070Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":"2406.06282","doi":"10.48550/arxiv.2406.06282","metadata_source":"pith","pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":"cs.LG","work_id":"c835ae36-bf61-4365-8156-febeba198d32","year":2024},"citing_paper":{"arxiv_id":"2606.21428","last_updated":"2026-07-09T17:35:35Z","snapshot_observed_at":"2026-08-14T19:40:18.846187Z","submitted_at":"2026-06-19T13:45:45Z","title":"Does Mixture-of-Experts Actually Help Inference on Consumer and Edge Hardware? An Empirical Study","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-06-26T12:30:55.628115Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2606.21428"},"observation_digest":"sha256:1eaa84d4b13af95dcc658193942b000282f3d806c8c51b58ea747d88f1c0192c","observation_id":"071dec5c-0a4d-49ed-b77c-d6b98206327d","resolution":{"observed_at":"2026-07-04T07:59:39.607441Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-07-12T13:05:17.273287Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.21428","last_updated":"2026-07-09T17:35:35Z","snapshot_observed_at":"2026-08-14T19:40:18.846187Z","submitted_at":"2026-06-19T13:45:45Z","title":"Does Mixture-of-Experts Actually Help Inference on Consumer and Edge Hardware? An Empirical Study","version":3},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-07-12T13:05:17.273287Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2606.21428"},"observation_digest":"sha256:5570af44701727b9db13c0e6ed843a1f52570bc700acd5650394d3732191bce1","observation_id":"38bbb818-13f7-46c1-afa4-c6412830500d","resolution":{"observed_at":"2026-07-12T13:05:17.273287Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":"2406.06282","doi":"10.48550/arxiv.2406.06282","metadata_source":"pith","pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":"cs.LG","work_id":"c835ae36-bf61-4365-8156-febeba198d32","year":2024},"citing_paper":{"arxiv_id":"2606.22496","last_updated":"2026-06-21T13:35:57Z","snapshot_observed_at":"2026-08-06T18:59:30.397360Z","submitted_at":"2026-06-21T13:35:57Z","title":"Enabling Cloud-Level Accuracy in Edge AI through IoT Data Preprocessing","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-06-26T09:34:00.058213Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2606.22496"},"observation_digest":"sha256:e820aa567425666426e380a9fa9033091b1abb791bf536b51670bf4df498fc41","observation_id":"a4708dba-eaae-4573-9c1b-af9c4fd2dc6b","resolution":{"observed_at":"2026-07-04T09:49:44.091015Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":"2406.06282","doi":"10.48550/arxiv.2406.06282","metadata_source":"pith","pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":"cs.LG","work_id":"c835ae36-bf61-4365-8156-febeba198d32","year":2024},"citing_paper":{"arxiv_id":"2606.23001","last_updated":"2026-06-24T14:38:38Z","snapshot_observed_at":"2026-08-16T08:38:42.892021Z","submitted_at":"2026-06-22T08:16:19Z","title":"EnerInfer: Energy-Aware On-Device LLM Inference","version":2},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-06-26T07:58:29.150228Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2606.23001"},"observation_digest":"sha256:bf70451d54f0081d36d00357e0c71eec5aafc04fcd1e09210975ef1586d7d387","observation_id":"02212a14-d014-4c2e-9473-8d7696156391","resolution":{"observed_at":"2026-07-04T11:29:50.714252Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-07-11T12:14:57.342551Z","title":"PowerInfer-2: Fast large language model inference on a smartphone.arXiv preprint arXiv:2406.06282, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.05475","last_updated":"2026-07-06T09:42:34Z","snapshot_observed_at":"2026-08-15T17:08:32.463717Z","submitted_at":"2026-07-06T09:42:34Z","title":"Is Your NPU Ready for LLMs? Dissecting the Hidden Efficiency Bottlenecks in Mobile LLM Inference","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-07-11T12:14:57.342551Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2607.05475"},"observation_digest":"sha256:bb0121d7b2b74939b2ea6581a7d4e4edebb8d9d864eca15dd38cc334899d2ce2","observation_id":"9cd432d8-7f33-45ec-a4fb-8fb12919a0e3","resolution":{"observed_at":"2026-07-11T12:14:57.342551Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":"2406.06282","doi":"10.48550/arxiv.2406.06282","metadata_source":"pith","pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Powerinfer-2: Fast large language model inference on a smartphone","venue":"cs.LG","work_id":"c835ae36-bf61-4365-8156-febeba198d32","year":2024},"citing_paper":{"arxiv_id":"2607.07046","last_updated":"2026-07-08T06:23:04Z","snapshot_observed_at":"2026-08-14T15:57:28.599272Z","submitted_at":"2026-07-08T06:23:04Z","title":"Voltron: Enabling Elastic Multi-Device Execution of LLM Inference for Empowered Edge Intelligence","version":1},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-07-09T21:08:24.293077Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2607.07046"},"observation_digest":"sha256:cbf0d14ce927d4a6e9b82bd18f086fc2113cc84c348aa708f82a6e13c3d6d8ae","observation_id":"868c65d1-9304-493d-af0e-5d65aefd8baf","resolution":{"observed_at":"2026-07-09T21:16:34.319320Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-02T06:22:22.488014Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.12839","last_updated":"2026-07-25T21:10:17Z","snapshot_observed_at":"2026-08-12T05:00:11.844991Z","submitted_at":"2026-07-14T14:56:14Z","title":"HeteroMosaic: Exposing and Exploiting Heterogeneous Execution Opportunities for Energy-Efficient Edge LLM Inference","version":4},"reference_index":90,"source":"pdf_text","source_observed_at":"2026-08-02T06:22:22.488014Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2607.12839"},"observation_digest":"sha256:9c4a4d447182dfd2c9500e8bc473893b49c47ad9b29d1e5d595ccd1b4286acab","observation_id":"6db40602-1fd7-4e89-b1d1-26327047031e","resolution":{"observed_at":"2026-08-02T06:22:22.488014Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-01T18:04:57.453503Z","title":"Powerinfer-2: Fast large language model inference on a smartphone,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.17415","last_updated":"2026-07-19T21:22:22Z","snapshot_observed_at":"2026-08-15T00:58:21.574607Z","submitted_at":"2026-07-19T21:22:22Z","title":"Transition-Aware Backend Dispatch for Edge LLM Inference","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-01T18:04:57.453503Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2607.17415"},"observation_digest":"sha256:f04d16aa3e825c8f4e115c303bc8403c5202a7f18d3179b406289d84c8dbb711","observation_id":"2c38e300-665a-4aef-a38d-3b57963224b7","resolution":{"observed_at":"2026-08-01T18:04:57.453503Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-01T16:14:37.007544Z","title":"Powerinfer-2: Fast large language model inference on a smartphone, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.18081","last_updated":"2026-07-20T15:48:33Z","snapshot_observed_at":"2026-08-14T15:51:31.448973Z","submitted_at":"2026-07-20T15:48:33Z","title":"SelectInfer: Selective Neuron Loading and Computation for On-Device LLMs","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-01T16:14:37.007544Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2607.18081"},"observation_digest":"sha256:0c55731b5454fee63722671531ca4f077e2bdff3fa0eb24298b14987faa7915f","observation_id":"88c1eab1-7bb5-4783-88e3-211bbbeb280d","resolution":{"observed_at":"2026-08-01T16:14:37.007544Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-08T04:21:30.070796Z","title":"PowerInfer-2: Fast large language model inference on a smartphone,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.02995","last_updated":"2026-08-04T01:21:29Z","snapshot_observed_at":"2026-08-16T18:28:53.013440Z","submitted_at":"2026-08-04T01:21:29Z","title":"SparSEEty: Extracting Tokens from Sparsity-Exploiting LLM Serving Systems via Deterministic Side Channels","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-08T04:21:30.070796Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2608.02995"},"observation_digest":"sha256:14a915e43965fa781b830db3dcaac2d84f27d6afe469be91172462e4dde4424b","observation_id":"57cc3a2f-510a-4154-936f-42009b093336","resolution":{"observed_at":"2026-08-08T04:21:30.070796Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06282","snapshot_observed_at":"2026-08-16T00:22:44.065594Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.12114","last_updated":"2026-08-12T14:35:21Z","snapshot_observed_at":"2026-08-17T07:38:53.830631Z","submitted_at":"2026-08-12T14:35:21Z","title":"The Ingestion Tax: Adopting File-Backed Weights in Tensor Frameworks","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-16T00:22:44.065594Z"},"links":{"cited_paper":"/paper/2406.06282","citing_paper":"/paper/2608.12114"},"observation_digest":"sha256:437f950dd230adea50fe3c22ab3ea7e12799e0a1d7433623dd3cda8ca6f8b861","observation_id":"38c75e83-06e6-482d-a605-704dc4492a13","resolution":{"observed_at":"2026-08-16T00:22:44.065594Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2406.06282/citation-record","integrity":"/paper/2406.06282/integrity","json":"/paper/2406.06282/citation-record.json","paper":"/paper/2406.06282"},"outbound":[],"paper":{"arxiv_id":"2406.06282","last_updated":"2024-12-12T12:24:18Z","latest_version":3,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-16T13:44:33.550467Z","submitted_at":"2024-06-10T14:01:21Z","title":"PowerInfer-2: Fast Large Language Model Inference on a Smartphone"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"thesis":"As of 18 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 42 inbound Pith citation observations for arXiv:2406.06282."}