{"as_of":"2026-08-11T15:31:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:3829eff7ed94eab8246accc5e080ed50b2a0f81405f1919f848057f01e606bd5","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":40,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":40,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-11T06:34:44.6726+00:00","state":"measured"},{"denominator":40,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":40,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-11T13:55:45.180184Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-10T11:37:03.185820Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":"2310.08659","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-07-10T11:37:03.185820Z","title":"Loftq: Lora- fine-tuning-aware quantization for large language models","venue":"cs.CL","work_id":"ed0c9106-2097-43fb-854d-d9945b1b9488","year":2023},"citing_paper":{"arxiv_id":"2403.14608","last_updated":"2024-09-16T02:54:50Z","snapshot_observed_at":"2026-08-04T09:07:42.158421Z","submitted_at":"2024-03-21T17:55:50Z","title":"Parameter-Efficient Fine-Tuning for Large Models: A Comprehensive Survey","version":7},"reference_index":125,"source":"pdf_text","source_observed_at":"2026-05-13T11:32:36.738536Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2403.14608"},"observation_digest":"sha256:9324aa07cbca544dc5808e29db0cb47fe103bb22eaca750fa223eb20909b856b","observation_id":"7d305197-e35e-4ecd-bec4-e98de6928c66","resolution":{"observed_at":"2026-05-13T11:32:37.151309Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":"2310.08659","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-07-10T11:37:03.185820Z","title":"Loftq: Lora- fine-tuning-aware quantization for large language models","venue":"cs.CL","work_id":"ed0c9106-2097-43fb-854d-d9945b1b9488","year":2023},"citing_paper":{"arxiv_id":"2404.14294","last_updated":"2024-07-19T04:47:36Z","snapshot_observed_at":"2026-07-06T18:03:47.096406Z","submitted_at":"2024-04-22T15:53:08Z","title":"A Survey on Efficient Inference for Large Language Models","version":3},"reference_index":191,"source":"pdf_text","source_observed_at":"2026-05-15T02:39:33.007894Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2404.14294"},"observation_digest":"sha256:925a1eb516f6940432374b9f30fb826a8fff4a9812d5dbee776eb62718783fd1","observation_id":"495db1fa-9662-48ee-845b-a18999a3ee4c","resolution":{"observed_at":"2026-05-15T02:39:33.195960Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":"2310.08659","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-07-10T11:37:03.185820Z","title":"Loftq: Lora- fine-tuning-aware quantization for large language models","venue":"cs.CL","work_id":"ed0c9106-2097-43fb-854d-d9945b1b9488","year":2023},"citing_paper":{"arxiv_id":"2409.18169","last_updated":"2026-04-23T18:48:49Z","snapshot_observed_at":"2026-08-10T21:15:46.437989Z","submitted_at":"2024-09-26T17:55:22Z","title":"Harmful Fine-tuning Attacks and Defenses for Large Language Models: A Survey","version":6},"reference_index":90,"source":"pdf_text","source_observed_at":"2026-05-23T20:58:16.237327Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2409.18169"},"observation_digest":"sha256:9a90860dc5c2683babc60c3056386fc362fc97b26cc1c53111131b3b64daa1eb","observation_id":"24bcf98d-e982-4731-b82c-a36fa93da434","resolution":{"observed_at":"2026-05-23T20:58:26.059834Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-11T13:55:45.180184Z","title":"Loftq: Lora-fine-tuning-aware quantization for large lan- guage models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.12687","last_updated":"2025-03-18T10:50:58Z","snapshot_observed_at":"2026-08-11T13:48:03.652411Z","submitted_at":"2024-12-17T09:08:18Z","title":"Uncertainty-Aware Hybrid Inference with On-Device Small and Remote Large Language Models","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-11T13:55:45.180184Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2412.12687"},"observation_digest":"sha256:84e7d50b57b45e23807bf7eb6efc5a1aaa94e1e6228493661356ce42962cd3ea","observation_id":"09eb085e-17ca-4abc-ac74-87654ef31a8f","resolution":{"observed_at":"2026-08-11T13:55:45.180184Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-11T13:09:46.615711Z","title":"Loftq: Lora-fine-tuning-aware quantization for large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.13437","last_updated":"2024-12-18T02:15:31Z","snapshot_observed_at":"2026-08-11T13:05:31.067210Z","submitted_at":"2024-12-18T02:15:31Z","title":"Deploying Foundation Model Powered Agent Services: A Survey","version":1},"reference_index":199,"source":"pdf_text","source_observed_at":"2026-08-11T13:09:46.615711Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2412.13437"},"observation_digest":"sha256:c738802b7348e16ae6c1c7526e69bf5f8da621153d8c9610089db5ce96b0dab8","observation_id":"b0ddbe91-aa63-4b18-94e8-81102fc49f8a","resolution":{"observed_at":"2026-08-11T13:09:46.615711Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-11T04:33:58.725286Z","title":"Loftq: Lora-fine-tuning-aware quantization for large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.18729","last_updated":"2024-12-25T01:10:25Z","snapshot_observed_at":"2026-08-11T04:29:47.428422Z","submitted_at":"2024-12-25T01:10:25Z","title":"Optimizing Large Language Models with an Enhanced LoRA Fine-Tuning Algorithm for Efficiency and Robustness in NLP Tasks","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-11T04:33:58.725286Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2412.18729"},"observation_digest":"sha256:cc209493fa35128889e12e6bc355d40826760b93a3376cd588286c3a2bdcd7aa","observation_id":"580bf108-4ee5-468a-ab14-f6d3f7105bab","resolution":{"observed_at":"2026-08-11T04:33:58.725286Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-10T23:20:15.321577Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for large lan- guage models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.20772","last_updated":"2025-03-07T05:24:42Z","snapshot_observed_at":"2026-08-11T06:43:55.686481Z","submitted_at":"2024-12-30T07:47:30Z","title":"Large Language Model Enabled Multi-Task Physical Layer Network","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-10T23:20:15.321577Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2412.20772"},"observation_digest":"sha256:2f9be6c024ef1657337ca38ac4e4fdda31ca2ae03aec3503d22db646d8309d27","observation_id":"741b03ff-afe5-4020-9fd4-7e25ea565df4","resolution":{"observed_at":"2026-08-10T23:20:15.321577Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-10T23:43:48.086587Z","title":"Loftq: Lora- fine-tuning-aware quantization for large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.00054","last_updated":"2024-12-28T04:44:07Z","snapshot_observed_at":"2026-08-10T23:36:32.553565Z","submitted_at":"2024-12-28T04:44:07Z","title":"AdvAnchor: Enhancing Diffusion Model Unlearning with Adversarial Anchors","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-10T23:43:48.086587Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2501.00054"},"observation_digest":"sha256:e3928665e3451b1802bf7c20f9b80940c7163624b0266a488a596e9885eb9d1a","observation_id":"283f9c67-e33e-4ae8-ad25-a3165e22c142","resolution":{"observed_at":"2026-08-10T23:43:48.086587Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-10T20:56:20.726481Z","title":"Loftq: Lora-fine-tuning-aware quantization for large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.07139","last_updated":"2025-01-13T08:58:00Z","snapshot_observed_at":"2026-08-10T23:44:44.763612Z","submitted_at":"2025-01-13T08:58:00Z","title":"FlexQuant: Elastic Quantization Framework for Locally Hosted LLM on Edge Devices","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-10T20:56:20.726481Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2501.07139"},"observation_digest":"sha256:f8544109547f12c17de184a38c826e83abc8ef828ef50a9b0c7280aaa5e2b966","observation_id":"f88a258f-7f1d-4def-a702-682ea89bb388","resolution":{"observed_at":"2026-08-10T20:56:20.726481Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-10T13:37:35.321241Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.16302","last_updated":"2025-01-27T18:42:48Z","snapshot_observed_at":"2026-08-10T13:30:38.410970Z","submitted_at":"2025-01-27T18:42:48Z","title":"Matryoshka Re-Ranker: A Flexible Re-Ranking Architecture With Configurable Depth and Width","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-10T13:37:35.321241Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2501.16302"},"observation_digest":"sha256:cdcfdef09d170a3d4140a7f3aabd62df55094f0e11ab0696fffcd018779271bb","observation_id":"f698981f-810b-4879-a20f-b5988f4a26d3","resolution":{"observed_at":"2026-08-10T13:37:35.321241Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-10T14:44:36.202543Z","title":"Loftq: Lora-fine-tuning-aware quantization for large language models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.16385","last_updated":"2025-05-23T05:33:23Z","snapshot_observed_at":"2026-08-11T05:44:38.575800Z","submitted_at":"2025-01-25T06:04:07Z","title":"FBQuant: FeedBack Quantization for Large Language Models","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-10T14:44:36.202543Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2501.16385"},"observation_digest":"sha256:37174b15ec284cfeed02ddf4ad9ab84fa6a3142cd60006308ed9667dc6033f5a","observation_id":"a9ec578e-86d3-4db1-b9c0-de359d3387ba","resolution":{"observed_at":"2026-08-10T14:44:36.202543Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-09T23:25:54.674492Z","title":"Loftq: Lora-fine-tuning-aware quantization for large language models.arXiv preprint arXiv:2310.08659, 2023a","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.18475","last_updated":"2025-08-14T03:12:42Z","snapshot_observed_at":"2026-08-11T03:00:13.330318Z","submitted_at":"2025-01-30T16:48:15Z","title":"CLoQ: Enhancing Fine-Tuning of Quantized LLMs via Calibrated LoRA Initialization","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-09T23:25:54.674492Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2501.18475"},"observation_digest":"sha256:9f9cb1877a5577a873ef1a35590ba5a3547f7f08e565328df62e4742e357c5f4","observation_id":"3a7dc95f-8886-4098-b396-fb5cc10540bb","resolution":{"observed_at":"2026-08-09T23:25:54.674492Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-09T17:26:45.478477Z","title":"Loftq: Lora-fine-tuning-aware quantization for large language models.arXiv preprint arXiv:2310.08659, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.00899","last_updated":"2025-02-02T20:23:32Z","snapshot_observed_at":"2026-08-10T22:34:31.726902Z","submitted_at":"2025-02-02T20:23:32Z","title":"HASSLE-free: A unified Framework for Sparse plus Low-Rank Matrix Decomposition for LLMs","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-09T17:26:45.478477Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2502.00899"},"observation_digest":"sha256:909c6b9bbf84e2c62290d36720ba896f187e7e273b61a5bed740a298ed787d13","observation_id":"c7257759-b6f9-401d-9f4f-720fa22eab29","resolution":{"observed_at":"2026-08-09T17:26:45.478477Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-08T12:26:16.314808Z","title":"Loftq: Lora-fine-tuning-aware quantization for large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.07547","last_updated":"2025-02-11T13:34:09Z","snapshot_observed_at":"2026-08-11T14:09:20.923196Z","submitted_at":"2025-02-11T13:34:09Z","title":"Instance-dependent Early Stopping","version":1},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-08T12:26:16.314808Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2502.07547"},"observation_digest":"sha256:e6dcba4a5ab231b4530aa989ed518de1c78feb8669fbd119c8f7e99e33d54c0d","observation_id":"a3e99755-66b2-4246-b47f-05bdb8fd7346","resolution":{"observed_at":"2026-08-08T12:26:16.314808Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-08T13:44:00.560727Z","title":"Loftq: Lora-fine-tuning-aware quantization for large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.07832","last_updated":"2025-02-11T00:21:40Z","snapshot_observed_at":"2026-08-10T01:10:52.838340Z","submitted_at":"2025-02-11T00:21:40Z","title":"SHARP: Accelerating Language Model Inference by SHaring Adjacent layers with Recovery Parameters","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-08T13:44:00.560727Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2502.07832"},"observation_digest":"sha256:1ef51281e308d101721c72bbfaf1fe4ab937b6fb62c708ada1bdda4fb9b79402","observation_id":"448f8bc4-5c02-4654-a90c-a552d35cd2f8","resolution":{"observed_at":"2026-08-08T13:44:00.560727Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-08T10:26:45.458342Z","title":"Loftq: Lora-fine-tuning-aware quantization for large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.08141","last_updated":"2025-02-12T05:48:26Z","snapshot_observed_at":"2026-08-10T13:36:20.860187Z","submitted_at":"2025-02-12T05:48:26Z","title":"LowRA: Accurate and Efficient LoRA Fine-Tuning of LLMs under 2 Bits","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-08T10:26:45.458342Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2502.08141"},"observation_digest":"sha256:87ce8fbc8bd117ce653bf32ab6c3f2bd9a83fac3f9a2025229b932392c4dbcb2","observation_id":"618795ae-8671-44ed-974f-c1d124e659d9","resolution":{"observed_at":"2026-08-08T10:26:45.458342Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-07T15:43:40.222936Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.14742","last_updated":"2025-05-29T22:04:36Z","snapshot_observed_at":"2026-08-09T14:33:06.239071Z","submitted_at":"2025-05-20T07:19:36Z","title":"Quaff: Quantized Parameter-Efficient Fine-Tuning under Outlier Spatial Stability Hypothesis","version":2},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-07T15:43:40.222936Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2505.14742"},"observation_digest":"sha256:ab9905f88c306324b232788889337ecbdef9e46cfed78b8b1172a282aef849f3","observation_id":"0258a746-0fdd-4232-b58e-f9a605c254f2","resolution":{"observed_at":"2026-08-07T15:43:40.222936Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-07T15:26:44.070397Z","title":"Loftq: Lora- fine-tuning-aware quantization for large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.15304","last_updated":"2025-05-30T10:57:06Z","snapshot_observed_at":"2026-08-07T23:42:58.683483Z","submitted_at":"2025-05-21T09:35:12Z","title":"Saliency-Aware Quantized Imitation Learning for Efficient Robotic Control","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T15:26:44.070397Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2505.15304"},"observation_digest":"sha256:d8d7d9916d6c6ab35656cf910a17fee1a96cafb8889eac3506ce386d208446a7","observation_id":"8d7d5870-8742-4079-8880-eecdc479c54f","resolution":{"observed_at":"2026-08-07T15:26:44.070397Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-07T14:46:00.357902Z","title":"Loftq: Lora-fine-tuning-aware quantization for large language models.arXiv preprint arXiv:2310.08659, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.17872","last_updated":"2025-05-27T07:23:28Z","snapshot_observed_at":"2026-08-08T06:53:10.432858Z","submitted_at":"2025-05-23T13:24:39Z","title":"Mixture of Low Rank Adaptation with Partial Parameter Sharing for Time Series Forecasting","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T14:46:00.357902Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2505.17872"},"observation_digest":"sha256:fe8ed04a2078c70ea84f1c261686ff8d7890e34e8170b97cbb36c539e547c754","observation_id":"91debe91-8a70-4300-af79-51c24fd523bd","resolution":{"observed_at":"2026-08-07T14:46:00.357902Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-07T13:42:47.300407Z","title":"Loftq: Lora-fine-tuning-aware quantization for large language models.arXiv preprint arXiv:2310.08659, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.21382","last_updated":"2025-05-27T16:10:53Z","snapshot_observed_at":"2026-08-09T22:42:41.800500Z","submitted_at":"2025-05-27T16:10:53Z","title":"DeCAF: Decentralized Consensus-And-Factorization for Low-Rank Adaptation of Foundation Models","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T13:42:47.300407Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2505.21382"},"observation_digest":"sha256:036a9633ad1b418eda383d81741218fce5cc9b701d010a58aef073e5e15fcab6","observation_id":"e0b61838-6496-408e-8588-2e9aa86bb6fc","resolution":{"observed_at":"2026-08-07T13:42:47.300407Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-06T22:43:56.321335Z","title":"Loftq: Lora-fine-tuning-aware quantization for large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.21119","last_updated":"2025-06-26T09:37:15Z","snapshot_observed_at":"2026-08-06T22:30:28.321862Z","submitted_at":"2025-06-26T09:37:15Z","title":"Progtuning: Progressive Fine-tuning Framework for Transformer-based Language Models","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T22:43:56.321335Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2506.21119"},"observation_digest":"sha256:0364962f10b1971a164ea487d4b6749468725dc525e90b9d70696b7dee748daa","observation_id":"a57bced4-5871-47f7-ab5c-06b8c8a85b68","resolution":{"observed_at":"2026-08-06T22:43:56.321335Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":"2310.08659","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-07-10T11:37:03.185820Z","title":"Loftq: Lora- fine-tuning-aware quantization for large language models","venue":"cs.CL","work_id":"ed0c9106-2097-43fb-854d-d9945b1b9488","year":2023},"citing_paper":{"arxiv_id":"2507.00029","last_updated":"2026-05-13T05:28:39Z","snapshot_observed_at":"2026-08-11T06:15:13.326412Z","submitted_at":"2025-06-17T14:58:54Z","title":"LoRA-Mixer: Coordinate Modular LoRA Experts Through Serial Attention Routing","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-05-19T09:05:25.236355Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2507.00029"},"observation_digest":"sha256:0b4b3af6cb382d09b1efacd96635d5cd2157b87cf35d68f12d13eb819de0ce5c","observation_id":"bd81917b-8d0c-4ab5-a18c-4bec614a83b3","resolution":{"observed_at":"2026-05-19T09:07:14.572473Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-06T18:48:52.053760Z","title":"Loftq: Lora- fine-tuning-aware quantization for large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.08044","last_updated":"2025-07-09T23:52:31Z","snapshot_observed_at":"2026-08-08T12:59:27.580998Z","submitted_at":"2025-07-09T23:52:31Z","title":"ConsNoTrainLoRA: Data-driven Weight Initialization of Low-rank Adapters using Constraints","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T18:48:52.053760Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2507.08044"},"observation_digest":"sha256:7f35e1e92e5598fcaf35493829c5a5b457ff95c2c0aaf570f8292353c3e4019b","observation_id":"bbe36e6d-c13a-4e40-bd3e-6a6bac78033f","resolution":{"observed_at":"2026-08-06T18:48:52.053760Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":"2310.08659","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-07-10T11:37:03.185820Z","title":"Loftq: Lora- fine-tuning-aware quantization for large language models","venue":"cs.CL","work_id":"ed0c9106-2097-43fb-854d-d9945b1b9488","year":2023},"citing_paper":{"arxiv_id":"2508.06974","last_updated":"2026-05-18T12:47:40Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-08-09T13:00:16Z","title":"Rethinking 1-bit Optimization Leveraging Pre-trained Large Language Models","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-21T23:44:01.953344Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2508.06974"},"observation_digest":"sha256:166365fead5db024a2f9a212dadfaa64f9a7174f8c2d10b4f7b46ad6776e703b","observation_id":"d6550fcd-d067-451e-a83f-53a038108174","resolution":{"observed_at":"2026-05-21T23:44:26.517127Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-05T15:52:44.972851Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.19432","last_updated":"2025-08-26T21:01:45Z","snapshot_observed_at":"2026-08-09T00:35:46.151432Z","submitted_at":"2025-08-26T21:01:45Z","title":"Quantized but Deceptive? A Multi-Dimensional Truthfulness Evaluation of Quantized LLMs","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-05T15:52:44.972851Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2508.19432"},"observation_digest":"sha256:8dc37cfcc6fdb1a15f09dd777c873dbe59183749566292dd2c600373bf62758f","observation_id":"d63e35b2-fe0b-4e0d-bdcc-ba2b240cf54a","resolution":{"observed_at":"2026-08-05T15:52:44.972851Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-03T18:56:57.627233Z","title":"Loftq: Lora-fine-tuning-aware quantization for large language models.arXiv preprint arXiv:2310.08659, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2512.02924","last_updated":"2026-07-20T02:01:50Z","snapshot_observed_at":"2026-08-06T15:38:52.220151Z","submitted_at":"2025-12-02T16:45:25Z","title":"AutoNeural: Co-Designing Vision-Language Models for NPU Inference","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-03T18:56:57.627233Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2512.02924"},"observation_digest":"sha256:15fd21105bd0cfdf5d465fdd7a48bc238aa453cd538cda46dce281f2d3853448","observation_id":"bd700577-3454-4994-b342-9aed805def73","resolution":{"observed_at":"2026-08-03T18:56:57.627233Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-07-13T23:28:12.790404Z","title":"Loftq: Lora-fine-tuning-aware quantization for large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2603.16867","last_updated":"2026-06-03T09:37:20Z","snapshot_observed_at":"2026-08-06T21:47:33.826393Z","submitted_at":"2026-03-17T17:59:51Z","title":"Efficient Reasoning on the Edge","version":2},"reference_index":113,"source":"pdf_text","source_observed_at":"2026-07-13T23:28:12.790404Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2603.16867"},"observation_digest":"sha256:eb7cf5e8ed32b6c5077ed9ee443194b91dad38910a0604ed242be3caa4eab5e6","observation_id":"445d0408-5d1a-425d-a68f-e98aad22d6b2","resolution":{"observed_at":"2026-07-13T23:28:12.790404Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":"2310.08659","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-07-10T11:37:03.185820Z","title":"Loftq: Lora- fine-tuning-aware quantization for large language models","venue":"cs.CL","work_id":"ed0c9106-2097-43fb-854d-d9945b1b9488","year":2023},"citing_paper":{"arxiv_id":"2604.02501","last_updated":"2026-04-02T20:09:05Z","snapshot_observed_at":"2026-07-30T03:07:00.572975Z","submitted_at":"2026-04-02T20:09:05Z","title":"ECG Foundation Models and Medical LLMs for Agentic Cardiovascular Intelligence at the Edge: A Review and Outlook","version":1},"reference_index":95,"source":"pdf_text","source_observed_at":"2026-05-13T20:23:15.138933Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2604.02501"},"observation_digest":"sha256:9a133351861ebac46c308c2e76ab22c7004937595ba57e57357a4a63928ba1fe","observation_id":"f4bae0c1-6444-4dd6-9af7-125d6d9971cd","resolution":{"observed_at":"2026-05-13T20:28:14.293292Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":"2310.08659","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-07-10T11:37:03.185820Z","title":"Loftq: Lora- fine-tuning-aware quantization for large language models","venue":"cs.CL","work_id":"ed0c9106-2097-43fb-854d-d9945b1b9488","year":2023},"citing_paper":{"arxiv_id":"2604.27763","last_updated":"2026-04-30T11:52:50Z","snapshot_observed_at":"2026-07-06T23:13:11.623453Z","submitted_at":"2026-04-30T11:52:50Z","title":"Intent2Tx: Benchmarking LLMs for Translating Natural Language Intents into Ethereum Transactions","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-07T05:49:57.957009Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2604.27763"},"observation_digest":"sha256:e0d5d8f2ed60dc6147ba25906e769e9369fcb34eb4c0d04a94286879b412bd2b","observation_id":"9e16b279-20c3-4dd1-9023-df2b50276d4a","resolution":{"observed_at":"2026-05-12T10:31:29.363812Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":"2310.08659","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-07-10T11:37:03.185820Z","title":"Loftq: Lora- fine-tuning-aware quantization for large language models","venue":"cs.CL","work_id":"ed0c9106-2097-43fb-854d-d9945b1b9488","year":2023},"citing_paper":{"arxiv_id":"2605.05819","last_updated":"2026-05-07T07:57:23Z","snapshot_observed_at":"2026-08-11T01:37:42.197379Z","submitted_at":"2026-05-07T07:57:23Z","title":"HCInfer: An Efficient Inference System via Error Compensation for Resource-Constrained Devices","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-08T14:44:13.936042Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2605.05819"},"observation_digest":"sha256:650ddfa6143ae3e20473cab069d171ad2ba1fe7258954e68b5f7617f39e52c41","observation_id":"ab54e6a1-af95-408d-a599-ead2e56b9bac","resolution":{"observed_at":"2026-05-11T18:41:10.671686Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":"2310.08659","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-07-10T11:37:03.185820Z","title":"Loftq: Lora- fine-tuning-aware quantization for large language models","venue":"cs.CL","work_id":"ed0c9106-2097-43fb-854d-d9945b1b9488","year":2023},"citing_paper":{"arxiv_id":"2605.14841","last_updated":"2026-05-14T13:46:04Z","snapshot_observed_at":"2026-08-09T07:14:33.916800Z","submitted_at":"2026-05-14T13:46:04Z","title":"GPart: End-to-End Isometric Fine-Tuning via Global Parameter Partitioning","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-06-30T21:31:36.692739Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2605.14841"},"observation_digest":"sha256:a06ac74404c0f0b8cba4107382f5ee8904345c480505e31584698763c93856c8","observation_id":"2bd09531-8c53-43d9-87c3-73b4f09c75be","resolution":{"observed_at":"2026-06-30T21:35:04.493416Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":"2310.08659","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-07-10T11:37:03.185820Z","title":"Loftq: Lora- fine-tuning-aware quantization for large language models","venue":"cs.CL","work_id":"ed0c9106-2097-43fb-854d-d9945b1b9488","year":2023},"citing_paper":{"arxiv_id":"2606.00494","last_updated":"2026-06-02T02:20:42Z","snapshot_observed_at":"2026-08-06T11:18:56.095781Z","submitted_at":"2026-05-30T02:54:40Z","title":"ProjQ: Project-and-Quantize for Adapter-Aware LLM Compression","version":2},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-06-28T19:10:58.845006Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2606.00494"},"observation_digest":"sha256:bba679390e1df8346f07003baf695ec465d98f6047ce3ea27407517e42d8f20a","observation_id":"40fe6ea4-dc0c-4d1a-8ed2-ddc95b7162db","resolution":{"observed_at":"2026-06-28T19:12:34.473054Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":"2310.08659","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-07-10T11:37:03.185820Z","title":"Loftq: Lora- fine-tuning-aware quantization for large language models","venue":"cs.CL","work_id":"ed0c9106-2097-43fb-854d-d9945b1b9488","year":2023},"citing_paper":{"arxiv_id":"2606.01412","last_updated":"2026-05-31T19:17:39Z","snapshot_observed_at":"2026-08-05T23:04:35.966323Z","submitted_at":"2026-05-31T19:17:39Z","title":"GPTQ-intrinsic LoRA: A Near-optimal Algorithm for Low-precision Quantization with Low-rank Adaptation","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-06-28T17:28:14.160341Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2606.01412"},"observation_digest":"sha256:3f1397d7fd1237979ec8dd7e8fbd1c92d793bfe4ca756b681c513aa68ea72744","observation_id":"f866fda1-75d6-460b-a520-3e7312368775","resolution":{"observed_at":"2026-07-01T21:06:14.455209Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":"2310.08659","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-07-10T11:37:03.185820Z","title":"Loftq: Lora- fine-tuning-aware quantization for large language models","venue":"cs.CL","work_id":"ed0c9106-2097-43fb-854d-d9945b1b9488","year":2023},"citing_paper":{"arxiv_id":"2606.11244","last_updated":"2026-06-04T22:38:53Z","snapshot_observed_at":"2026-08-08T01:56:34.676724Z","submitted_at":"2026-06-04T22:38:53Z","title":"SPEAR: A System for Post-Quantization Error-Adaptive Recovery Enabling Efficient Low-Bit LLM Serving","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-06-27T23:03:58.078054Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2606.11244"},"observation_digest":"sha256:35807329a89e579ef8ddcfe38def689aa899cfc244de9fadb1969c89f72c433d","observation_id":"55e26d30-5115-4edf-ac60-b3f3571e4ad9","resolution":{"observed_at":"2026-07-02T16:07:08.479523Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":"2310.08659","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-07-10T11:37:03.185820Z","title":"Loftq: Lora- fine-tuning-aware quantization for large language models","venue":"cs.CL","work_id":"ed0c9106-2097-43fb-854d-d9945b1b9488","year":2023},"citing_paper":{"arxiv_id":"2607.08194","last_updated":"2026-07-09T07:50:49Z","snapshot_observed_at":"2026-08-06T01:16:56.040734Z","submitted_at":"2026-07-09T07:50:49Z","title":"Dive Into the Implicit Biases of Low-rank Vision-language Alignment","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-07-10T11:34:56.843122Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2607.08194"},"observation_digest":"sha256:c12d52d610b8e8162e2a0b69a575734fbe21c33b51c96eb225cb40e78ba90cf7","observation_id":"3bf18679-2eff-49e8-8d10-af39729de88f","resolution":{"observed_at":"2026-07-10T11:37:03.187809Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-02T02:01:53.202217Z","title":"Loftq: Lora-fine-tuning-aware quantization for large language models.arXiv preprint arXiv:2310.08659, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.14506","last_updated":"2026-07-16T02:42:24Z","snapshot_observed_at":"2026-08-10T11:42:23.996756Z","submitted_at":"2026-07-16T02:42:24Z","title":"Non-vacuous Generalization Bounds for Reinforcement Learning with Verifiable Rewards","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-02T02:01:53.202217Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2607.14506"},"observation_digest":"sha256:29192734c21ddee33abe924258fabb7d504561b424d7c2406848143b11078b7c","observation_id":"dbf61f83-0c84-4835-8c80-dc3831fcec6a","resolution":{"observed_at":"2026-08-02T02:01:53.202217Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-01T16:42:29.884291Z","title":"arXiv preprint arXiv:2310.08659 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.26404","last_updated":"2026-07-29T02:35:55Z","snapshot_observed_at":"2026-08-03T19:06:45.585998Z","submitted_at":"2026-07-29T02:35:55Z","title":"Examining the Efficacy of Graph Neural Network Message-Passing in Regression Contexts","version":1},"reference_index":239,"source":"arxiv_source","source_observed_at":"2026-08-01T16:42:29.884291Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2607.26404"},"observation_digest":"sha256:cb9d5ea5020b6e698a0ec0003dd61f6122b6e3a2d03aeb6cc2c5af58b90473d1","observation_id":"63ae96a0-4c52-4f9d-afc7-3476d71482ca","resolution":{"observed_at":"2026-08-01T16:42:29.884291Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-06T00:10:39.915105Z","title":"arXiv preprint arXiv:2310.08659 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.01530","last_updated":"2026-08-02T22:49:28Z","snapshot_observed_at":"2026-08-10T11:56:07.531617Z","submitted_at":"2026-08-02T22:49:28Z","title":"ST-LoRA: Single Trajectory LoRA Ensemble for Uncertainty Aware Agricultural Segmentation","version":1},"reference_index":146,"source":"arxiv_source","source_observed_at":"2026-08-06T00:10:39.915105Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2608.01530"},"observation_digest":"sha256:b9db5ecbf0feb6b240f57944b60f62de5c4350795d762759218b215992607a18","observation_id":"cd470429-be81-4e33-8050-4eeef7875a07","resolution":{"observed_at":"2026-08-06T00:10:39.915105Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-08T17:54:40.235752Z","title":"arXiv preprint arXiv:2310.08659 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.05224","last_updated":"2026-08-09T11:16:05Z","snapshot_observed_at":"2026-08-11T15:18:43.928335Z","submitted_at":"2026-08-05T11:34:20Z","title":"Small Foundation Models of Human Cognition and Behaviour","version":1},"reference_index":109,"source":"arxiv_source","source_observed_at":"2026-08-08T17:54:40.235752Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2608.05224"},"observation_digest":"sha256:19a5410b8b54ec71870567d418af52fbb93b11133f34562b34c648f8570d61ad","observation_id":"44f63c31-715f-44d3-9130-7603eed4dba2","resolution":{"observed_at":"2026-08-08T17:54:40.235752Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.08659","snapshot_observed_at":"2026-08-11T04:20:16.502030Z","title":"arXiv preprint arXiv:2310.08659 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.05224","last_updated":"2026-08-09T11:16:05Z","snapshot_observed_at":"2026-08-11T15:18:43.928335Z","submitted_at":"2026-08-05T11:34:20Z","title":"Small Foundation Models of Human Cognition and Behaviour","version":2},"reference_index":109,"source":"arxiv_source","source_observed_at":"2026-08-11T04:20:16.502030Z"},"links":{"cited_paper":"/paper/2310.08659","citing_paper":"/paper/2608.05224"},"observation_digest":"sha256:6d5eae7f3aaf7d363851bd02f6e2effe6edec4ae5abd81bae42c17657a8a4bc1","observation_id":"48a5f62e-59e2-4bc5-b908-d551709fce22","resolution":{"observed_at":"2026-08-11T04:20:16.502030Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2310.08659/citation-record","integrity":"/paper/2310.08659/integrity","json":"/paper/2310.08659/citation-record.json","paper":"/paper/2310.08659"},"outbound":[],"paper":{"arxiv_id":"2310.08659","last_updated":"2023-11-28T16:06:59Z","latest_version":4,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-10T01:10:20.500227Z","submitted_at":"2023-10-12T18:34:08Z","title":"LoftQ: LoRA-Fine-Tuning-Aware Quantization for Large Language Models"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"thesis":"As of 11 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 40 inbound Pith citation observations for arXiv:2310.08659."}