{"as_of":"2026-08-17T13:23:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:9a6fbbce2ff735f54421650b9f6ce2cd284364b4524c39677d7f9126eb1dd9f1","coverage":[{"denominator":72,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":72,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-08T16:44:00.773941Z","state":"measured"},{"denominator":83,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":83,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-17T06:30:58.91139+00:00","state":"measured"},{"denominator":11,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":11,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T20:51:38.685943Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"cited_work":{"arxiv_id":"2502.06130","doi":"10.48550/arxiv.2502.06130","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.06130","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Self- correcting decoding with generative feedback for mitigat- ing hallucinations in large vision-language models","venue":"ArXiv.org","work_id":"cf4a267b-c280-4501-893f-1266b64f9e67","year":2025},"citing_paper":{"arxiv_id":"2404.18930","last_updated":"2025-04-01T18:36:08Z","snapshot_observed_at":"2026-08-06T19:08:14.800394Z","submitted_at":"2024-04-29T17:59:41Z","title":"Hallucination of Multimodal Large Language Models: A Survey","version":2},"reference_index":206,"source":"pdf_text","source_observed_at":"2026-05-11T12:33:32.631346Z"},"links":{"cited_paper":"/paper/2502.06130","citing_paper":"/paper/2404.18930"},"observation_digest":"sha256:4eca7f5b7ad0a62a0518f964097e6089016c59665d460d81400279718cb20f08","observation_id":"8fbf3f47-e14b-4c3a-9ff6-02e9c2d5dd8e","resolution":{"observed_at":"2026-05-11T12:33:33.355426Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.06130","snapshot_observed_at":"2026-08-15T20:51:38.685943Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.17061","last_updated":"2025-06-10T05:05:02Z","snapshot_observed_at":"2026-08-15T20:42:25.761129Z","submitted_at":"2025-05-17T09:44:18Z","title":"Mixture of Decoding: An Attention-Inspired Adaptive Decoding Strategy to Mitigate Hallucinations in Large Vision-Language Models","version":3},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-15T20:51:38.685943Z"},"links":{"cited_paper":"/paper/2502.06130","citing_paper":"/paper/2505.17061"},"observation_digest":"sha256:d1f93b1869ea62978858bcf3326442ef9ab92091ed050d87ad4a0ca41e93e528","observation_id":"4cab3818-2559-4809-9ed5-14e1fc7b6856","resolution":{"observed_at":"2026-08-15T20:51:38.685943Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.06130","snapshot_observed_at":"2026-08-05T19:52:50.367509Z","title":"Self- correcting decoding with generative feedback for mitigat- ing hallucinations in large vision-language models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.11616","last_updated":"2025-08-15T17:29:06Z","snapshot_observed_at":"2026-08-14T10:03:17.394030Z","submitted_at":"2025-08-15T17:29:06Z","title":"Controlling Multimodal LLMs via Reward-guided Decoding","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-05T19:52:50.367509Z"},"links":{"cited_paper":"/paper/2502.06130","citing_paper":"/paper/2508.11616"},"observation_digest":"sha256:2cee48c76107306ea7fa2771e50b3dd6811cc5c61dcf2fb9cdf993f8bfe8f42a","observation_id":"a10d5500-02b9-4cc8-b43e-a2fbc8b451a7","resolution":{"observed_at":"2026-08-05T19:52:50.367509Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"cited_work":{"arxiv_id":"2502.06130","doi":"10.48550/arxiv.2502.06130","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.06130","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Self- correcting decoding with generative feedback for mitigat- ing hallucinations in large vision-language models","venue":"ArXiv.org","work_id":"cf4a267b-c280-4501-893f-1266b64f9e67","year":2025},"citing_paper":{"arxiv_id":"2604.12424","last_updated":"2026-04-14T08:15:44Z","snapshot_observed_at":"2026-08-15T02:00:35.039791Z","submitted_at":"2026-04-14T08:15:44Z","title":"Decoding by Perturbation: Mitigating MLLM Hallucinations via Dynamic Textual Perturbation","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-05-10T15:17:39.344348Z"},"links":{"cited_paper":"/paper/2502.06130","citing_paper":"/paper/2604.12424"},"observation_digest":"sha256:f37350ea463e573c574f9137027104fc740aec63e0a9868b7a29e4def2533657","observation_id":"159c0420-4dc4-4e6c-bd8d-dbef1ae14588","resolution":{"observed_at":"2026-05-11T10:51:03.628937Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"cited_work":{"arxiv_id":"2502.06130","doi":"10.48550/arxiv.2502.06130","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.06130","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Self- correcting decoding with generative feedback for mitigat- ing hallucinations in large vision-language models","venue":"ArXiv.org","work_id":"cf4a267b-c280-4501-893f-1266b64f9e67","year":2025},"citing_paper":{"arxiv_id":"2604.20366","last_updated":"2026-04-22T09:02:17Z","snapshot_observed_at":"2026-08-16T23:23:05.816374Z","submitted_at":"2026-04-22T09:02:17Z","title":"Mitigating Hallucinations in Large Vision-Language Models without Performance Degradation","version":1},"reference_index":161,"source":"arxiv_source","source_observed_at":"2026-05-10T00:33:39.960170Z"},"links":{"cited_paper":"/paper/2502.06130","citing_paper":"/paper/2604.20366"},"observation_digest":"sha256:2c7b75f7aacc378661aeda6873d19fda1b969efec552ef676001ccc3fa7273b4","observation_id":"f69f7d80-3727-4f49-b20e-2d7fb044997d","resolution":{"observed_at":"2026-05-10T00:34:47.833651Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"cited_work":{"arxiv_id":"2502.06130","doi":"10.48550/arxiv.2502.06130","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.06130","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Self- correcting decoding with generative feedback for mitigat- ing hallucinations in large vision-language models","venue":"ArXiv.org","work_id":"cf4a267b-c280-4501-893f-1266b64f9e67","year":2025},"citing_paper":{"arxiv_id":"2604.21027","last_updated":"2026-08-02T16:15:42Z","snapshot_observed_at":"2026-08-10T23:41:30.183421Z","submitted_at":"2026-04-22T19:18:36Z","title":"HypEHR: Hyperbolic Modeling of Electronic Health Records for Efficient Question Answering","version":1},"reference_index":282,"source":"arxiv_source","source_observed_at":"2026-05-09T23:51:47.724033Z"},"links":{"cited_paper":"/paper/2502.06130","citing_paper":"/paper/2604.21027"},"observation_digest":"sha256:0448b55f7dbd017d5b8ca7ba24bb01c53d662f5e99122c292963af8a37008ce2","observation_id":"a5c88132-7129-46b9-9f38-6310f6739f5d","resolution":{"observed_at":"2026-05-09T23:54:45.324151Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"cited_work":{"arxiv_id":"2502.06130","doi":"10.48550/arxiv.2502.06130","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.06130","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Self- correcting decoding with generative feedback for mitigat- ing hallucinations in large vision-language models","venue":"ArXiv.org","work_id":"cf4a267b-c280-4501-893f-1266b64f9e67","year":2025},"citing_paper":{"arxiv_id":"2605.04874","last_updated":"2026-05-06T13:08:12Z","snapshot_observed_at":"2026-08-11T19:12:26.096877Z","submitted_at":"2026-05-06T13:08:12Z","title":"Uncertainty-Aware Exploratory Direct Preference Optimization for Multimodal Large Language Models","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-05-08T17:00:36.574362Z"},"links":{"cited_paper":"/paper/2502.06130","citing_paper":"/paper/2605.04874"},"observation_digest":"sha256:b30e00b886e48c278ccc88df6c33b9c66aba38c0fff09716b9b065e82b192e94","observation_id":"b852acf4-b426-44d7-abcd-46de56601f94","resolution":{"observed_at":"2026-05-11T17:51:09.439758Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"cited_work":{"arxiv_id":"2502.06130","doi":"10.48550/arxiv.2502.06130","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.06130","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Self- correcting decoding with generative feedback for mitigat- ing hallucinations in large vision-language models","venue":"ArXiv.org","work_id":"cf4a267b-c280-4501-893f-1266b64f9e67","year":2025},"citing_paper":{"arxiv_id":"2606.27596","last_updated":"2026-06-25T22:55:46Z","snapshot_observed_at":"2026-08-12T17:23:17.939471Z","submitted_at":"2026-06-25T22:55:46Z","title":"Dismantling Pathological Shortcuts: A Causal Framework for Faithful LVLM Decoding","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-06-29T01:18:13.657975Z"},"links":{"cited_paper":"/paper/2502.06130","citing_paper":"/paper/2606.27596"},"observation_digest":"sha256:0b95e774c5685c2a414e6c393238c5dc03719a4bb3ef38bcdf73a4b0869562d6","observation_id":"e7205243-3903-4706-96c6-1dfcb0a82edf","resolution":{"observed_at":"2026-07-01T19:06:02.566245Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"cited_work":{"arxiv_id":"2502.06130","doi":"10.48550/arxiv.2502.06130","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.06130","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Self- correcting decoding with generative feedback for mitigat- ing hallucinations in large vision-language models","venue":"ArXiv.org","work_id":"cf4a267b-c280-4501-893f-1266b64f9e67","year":2025},"citing_paper":{"arxiv_id":"2606.31054","last_updated":"2026-06-30T02:46:10Z","snapshot_observed_at":"2026-08-07T17:43:53.022630Z","submitted_at":"2026-06-30T02:46:10Z","title":"ADAPT: Attention Dynamics Alignment with Preference Tuning for Faithful MLLMs","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-07-01T06:51:01.371390Z"},"links":{"cited_paper":"/paper/2502.06130","citing_paper":"/paper/2606.31054"},"observation_digest":"sha256:1f521745ee9aab96ad0c7e4c9fbfef3bb66b16d00921022907b4a43df42c5da7","observation_id":"09666c7e-0248-4980-9cf1-b4f5b48fa0b6","resolution":{"observed_at":"2026-07-01T06:55:29.451277Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.06130","snapshot_observed_at":"2026-07-31T23:32:10.389429Z","title":"arXiv preprint arXiv:2502.06130 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.23944","last_updated":"2026-07-27T02:40:49Z","snapshot_observed_at":"2026-08-13T00:45:18.135377Z","submitted_at":"2026-07-27T02:40:49Z","title":"DICA: Dual-Indicator Guided Contrastive Alignment in Multimodal Large Language Models","version":1},"reference_index":71,"source":"arxiv_source","source_observed_at":"2026-07-31T23:32:10.389429Z"},"links":{"cited_paper":"/paper/2502.06130","citing_paper":"/paper/2607.23944"},"observation_digest":"sha256:ff48dde1458d5afcdb68a7dc06bea7147923449f4a2a4df469370fb323840564","observation_id":"aa062bbb-00ab-4ac8-87ac-6ffb9b307a72","resolution":{"observed_at":"2026-07-31T23:32:10.389429Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.06130","snapshot_observed_at":"2026-08-15T14:19:24.819933Z","title":"arXiv preprint arXiv:2502.06130 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.11474","last_updated":"2026-08-11T22:25:22Z","snapshot_observed_at":"2026-08-17T11:52:16.359015Z","submitted_at":"2026-08-11T22:25:22Z","title":"Test-Time Hallucination Control in Large Vision-Language Models","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-15T14:19:24.819933Z"},"links":{"cited_paper":"/paper/2502.06130","citing_paper":"/paper/2608.11474"},"observation_digest":"sha256:f6cced586a97e00ea0154e06f05c8a78094c1f7a65532a21d5000f756bbdd567","observation_id":"b8c614ac-b8ff-4999-82da-d0d1bcd75149","resolution":{"observed_at":"2026-08-15T14:19:24.819933Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2502.06130/citation-record","integrity":"/paper/2502.06130/integrity","json":"/paper/2502.06130/citation-record.json","paper":"/paper/2502.06130"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:00.552417Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.552417Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:04b3aba003fdf2305445e576bd02bf56cd6309592d5e6a7514bd8297208ebfee","observation_id":"7db877c2-cdc7-4ab4-b7fe-784e658c3b8f","resolution":{"observed_at":"2026-08-08T16:44:00.552417Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2112.00390","last_updated":"2022-09-07T04:14:48Z","snapshot_observed_at":"2026-08-16T17:37:55.747156Z","submitted_at":"2021-12-01T10:17:25Z","title":"SegDiff: Image Segmentation with Diffusion Probabilistic Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2112.00390","snapshot_observed_at":"2026-08-08T16:44:00.556993Z","title":"Segdiff: Image segmentation with diffusion probabilistic models","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.556993Z"},"links":{"cited_paper":"/paper/2112.00390","citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:9c9c36c9406cb9912ee765345935df1491db373f1bb36c9d8ea700488c068976","observation_id":"a0040755-0538-4056-a52f-583597614bf9","resolution":{"observed_at":"2026-08-08T16:44:00.556993Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.12966","last_updated":"2023-10-13T02:41:28Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-08-24T17:59:17Z","title":"Qwen-VL: A Versatile Vision-Language Model for Understanding, Localization, Text Reading, and Beyond","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.12966","snapshot_observed_at":"2026-08-08T16:44:00.561438Z","title":"Qwen-vl: A frontier large vision-language model with versatile abilities","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.561438Z"},"links":{"cited_paper":"/paper/2308.12966","citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:b5e104e863f8c9f94bc38beaff4842f8716334b4fa16db9b0df786327b311b1b","observation_id":"ec206db2-aec9-4037-b256-3688c6d940b5","resolution":{"observed_at":"2026-08-08T16:44:00.561438Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.18930","last_updated":"2025-04-01T18:36:08Z","snapshot_observed_at":"2026-08-06T19:08:14.800394Z","submitted_at":"2024-04-29T17:59:41Z","title":"Hallucination of Multimodal Large Language Models: A Survey","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.18930","snapshot_observed_at":"2026-08-08T16:44:00.565830Z","title":"Hallucination of multimodal large language models: A survey","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.565830Z"},"links":{"cited_paper":"/paper/2404.18930","citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:1e1f78b0ac4d9218f35c3534f99a8e036507d51906ac315a4fbcb07741ed1440","observation_id":"97798ece-d41b-4e44-8b4a-442deab1e496","resolution":{"observed_at":"2026-08-08T16:44:00.565830Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.651261Z","title":"Muse: Text-to-image generation via masked generative transformers","venue":null,"work_id":"65f06fb6-b82c-4d14-9057-ae7ecf933120","year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.568916Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:03e90b4cbe01c6d1498bf57f239f437ebe1a5cbef50086b25a9c164d3d7091f1","observation_id":"1aab79ab-6a0a-48c3-8e11-319de53e660b","resolution":{"observed_at":"2026-08-08T16:44:01.654734Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.642127Z","title":"Alleviating hallucinations in large vision-language models through hallucination-induced optimization","venue":null,"work_id":"ca70d15e-c212-4849-b1ab-305017e7834b","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.571384Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:9696c4ad5be42a3f789fea555aead2e6af816ccb7d6d9449257fda3826150712","observation_id":"a9cd9d8c-a77e-4b6e-9864-b9be6388cdd4","resolution":{"observed_at":"2026-08-08T16:44:01.645391Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10185","last_updated":"2026-04-27T07:59:24Z","snapshot_observed_at":"2026-07-06T18:31:08.584682Z","submitted_at":"2024-06-14T17:14:22Z","title":"Detecting and Evaluating Medical Hallucinations in Large Vision Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10185","snapshot_observed_at":"2026-08-08T16:44:00.574564Z","title":"Detecting and evaluating medical hallucinations in large vision language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.574564Z"},"links":{"cited_paper":"/paper/2406.10185","citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:051cb3444caa0b20979b7365fdfb47459238d5e8b50f9d12991c0817bdc8e9b2","observation_id":"b3026656-c2f4-424e-b1ed-3b5f924ae8de","resolution":{"observed_at":"2026-08-08T16:44:00.574564Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.631454Z","title":"Multi-object hallucination in vision language models","venue":null,"work_id":"aaf67ea6-ac7e-4389-acdf-f249e1b2cee7","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.577964Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:f57981329ca73541754767a4ff5491c604773f794acf7e85f7e18c9e44d167f7","observation_id":"1af4a25c-c05f-42a3-a8e5-5f75ff3fd62f","resolution":{"observed_at":"2026-08-08T16:44:01.635299Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.621222Z","title":"HALC : Object hallucination reduction via adaptive focal-contrast decoding","venue":null,"work_id":"079bbfd6-c006-4092-aca6-6bdfaff02f82","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.580731Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:f26aa3b99b575109c2e0101fa2ac96b54d77fdee6fb69178668fd4f851bb3a1a","observation_id":"ef810f8e-451c-4897-beaf-05168cf996e3","resolution":{"observed_at":"2026-08-08T16:44:01.625027Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.16479","last_updated":"2023-11-27T09:30:02Z","snapshot_observed_at":"2026-08-16T14:40:09.176614Z","submitted_at":"2023-11-27T09:30:02Z","title":"Mitigating Hallucination in Visual Language Models with Visual Supervision","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.16479","snapshot_observed_at":"2026-08-08T16:44:00.583461Z","title":"Mitigating hallucination in visual language models with visual supervision","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.583461Z"},"links":{"cited_paper":"/paper/2311.16479","citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:38396175e62ba31100e385be02515506bb7a448b8ed210a52a1bc864e6c20407","observation_id":"1b6b05fa-1172-4f3b-8e90-2ca6cf147c20","resolution":{"observed_at":"2026-08-08T16:44:00.583461Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.611224Z","title":"Reproducible scaling laws for contrastive language-image learning","venue":null,"work_id":"f49f53c6-340e-405b-92d1-68eebad0b5cb","year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.586744Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:c7a8adb83f7e4a78e7e321af9a458a233b730fc9ac6d12998819012912e80b4c","observation_id":"94182b5b-c08f-41e9-8b0b-43542756d8df","resolution":{"observed_at":"2026-08-08T16:44:01.614431Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:00.589567Z","title":"Gonzalez, Ion Stoica, and Eric P","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.589567Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:faf6d2dca611551fedb752a91cf9faa6dc65f87207eeca97edb26722225b71ac","observation_id":"b2e6c437-172a-4adf-95d8-723018aed653","resolution":{"observed_at":"2026-08-08T16:44:00.589567Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.595391Z","title":"Palm: Scaling language modeling with pathways","venue":null,"work_id":"db29afa3-727b-4b0e-b404-0db8b3d7b41e","year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.592103Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:4580d68063aef5b0952b362a1bec6cb631e36e1dfd92cd0f536c3c1f689529be","observation_id":"03719736-3f43-46c2-9d6e-511e1bcc1f83","resolution":{"observed_at":"2026-08-08T16:44:01.598845Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.586192Z","title":"Glass, and Pengcheng He","venue":null,"work_id":"ec41f920-8e33-4344-a0bd-17c22c4976ab","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.595277Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:eabebd795fb0ba7cda7d9078a593db9853c1052fa2e7a9997ea9042fdf40d754","observation_id":"5499fc33-cba0-4fc7-bce4-22578fc813cb","resolution":{"observed_at":"2026-08-08T16:44:01.589133Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:00.598342Z","title":"Diffusion models in vision: A survey","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.598342Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:7773bf8546591c1bfcafa1f903202efe301f882e961b411f18ddd0317db1ef13","observation_id":"e265d2c3-8ea0-4d93-935d-c514e7cbc907","resolution":{"observed_at":"2026-08-08T16:44:00.598342Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.571987Z","title":"Instructblip: towards general-purpose vision-language models with instruction tuning","venue":null,"work_id":"ea5f2b4a-0cb1-4488-988c-01b3ae6a6ebd","year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.601023Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:b1e03d1ae99bdc2a89c01f77f2beb85dc045af40c555fc0ddd457f9fef58676f","observation_id":"731c4cbb-6938-4120-a50d-42eb869ba011","resolution":{"observed_at":"2026-08-08T16:44:01.575892Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.15300","last_updated":"2024-04-23T09:32:25Z","snapshot_observed_at":"2026-08-16T14:15:54.682908Z","submitted_at":"2024-02-23T12:57:16Z","title":"Seeing is Believing: Mitigating Hallucination in Large Vision-Language Models via CLIP-Guided Decoding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.15300","snapshot_observed_at":"2026-08-08T16:44:00.603815Z","title":"Seeing is believing: Mitigating hallucination in large vision-language models via clip-guided decoding","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.603815Z"},"links":{"cited_paper":"/paper/2402.15300","citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:b44b0ce7ab33e678c63c3aff28e2acc7783c61cb68009e2a7082cd39212672a5","observation_id":"19836e82-2759-4388-8322-f5942fcc9886","resolution":{"observed_at":"2026-08-08T16:44:00.603815Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.564374Z","title":"Multi-modal hallucination control by visual information grounding","venue":null,"work_id":"68ff190c-5626-4767-a16a-32f9c28ef631","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.607286Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:9b65186fbf407d0167d18b08fbe577aaa0e5fdb07f74cabc1a68e5a1c67e8fc2","observation_id":"fef21a7a-9306-4168-9871-f7c6e0c7f7f9","resolution":{"observed_at":"2026-08-08T16:44:01.566960Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.13394","last_updated":"2025-10-24T02:45:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-06-23T09:22:36Z","title":"MME: A Comprehensive Evaluation Benchmark for Multimodal Large Language Models","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.13394","snapshot_observed_at":"2026-08-08T16:44:00.609957Z","title":"Mme: A comprehensive evaluation benchmark for multimodal large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.609957Z"},"links":{"cited_paper":"/paper/2306.13394","citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:574290881fd17be3b83e557bee1d4ee85754b7e440679be0ac8868c34e6927ce","observation_id":"11c6070d-9e78-446e-8b54-f4986b3bcf3e","resolution":{"observed_at":"2026-08-08T16:44:00.609957Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.556947Z","title":"Expressive text-to-image generation with rich text","venue":null,"work_id":"b67347c8-8bb0-4ea2-8b60-3381eb909772","year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.613273Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:c270f06d43281607142c2a316377e4ebee383f666c9adeb8001708618046a372","observation_id":"ed094f39-52c1-4eba-a6ba-4ca8d837a8da","resolution":{"observed_at":"2026-08-08T16:44:01.559666Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.548229Z","title":"Generative adversarial nets","venue":null,"work_id":"fb0d92d7-02d9-4b08-bb8f-6408933dcaae","year":2014},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.616367Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:0fe5ea77e211789a88e38f480aa042dae00d8982de098545f96486ffd7a79b43","observation_id":"80da1add-0d29-4791-b8ff-ea70e9661e4b","resolution":{"observed_at":"2026-08-08T16:44:01.551459Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.539649Z","title":"Detecting and preventing hallucinations in large vision language models","venue":null,"work_id":"d180a625-d6fb-4adb-bd52-d672eb1af1ea","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.619704Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:8dc89bba1329e28fdd98301dd6d1ebea3cfe2758eb3442e0d8f4ef5b5f14f88d","observation_id":"f5b8216e-d83c-4f28-845e-465d9bb00404","resolution":{"observed_at":"2026-08-08T16:44:01.542874Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.529895Z","title":"Clipscore: A reference-free evaluation metric for image captioning","venue":null,"work_id":"b3820c24-2e26-48bb-834e-3304b8b1c496","year":2021},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.622749Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:399c9282897ef3f58854c8af7f1c352ae1150344c845000fe0b00038825b6a8b","observation_id":"c536c0ce-5a29-4c20-a484-6dae53d2a0a8","resolution":{"observed_at":"2026-08-08T16:44:01.533595Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:00.625989Z","title":"Denoising diffusion probabilistic models","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.625989Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:f89b40425b5461e3b8548e16d566f9dcc172852fb82190a796aea5c3b18c80c7","observation_id":"9d4102d0-290c-44b0-a4af-fb0d090a926e","resolution":{"observed_at":"2026-08-08T16:44:00.625989Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.515466Z","title":"Opera: Alleviating hallucination in multi-modal large language models via over-trust penalty and retrospection-allocation","venue":null,"work_id":"e7ca94d7-5e2b-4b98-bd26-873141b58981","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.628965Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:acdd6fd612ff1a9a3a36367f15f01eb028f3ac33c1a8776099959b15d2de899d","observation_id":"c2f71472-851a-48ff-bde3-f0310ff4db3b","resolution":{"observed_at":"2026-08-08T16:44:01.519185Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.504530Z","title":"Gqa: A new dataset for real-world visual reasoning and compositional question answering","venue":null,"work_id":"c476cea0-9d41-4145-9c8a-b02462a62670","year":2019},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.631769Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:4fb1223be103c793d2093d184a30360795fb347da0bda82659f935a0123347a6","observation_id":"7bfb7245-78c4-419a-830a-f6fd35471e82","resolution":{"observed_at":"2026-08-08T16:44:01.508317Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.495119Z","title":"Hallucination augmented contrastive learning for multimodal large language model","venue":null,"work_id":"2f9f8850-cf71-4553-ae6c-691d9bb8167a","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.635162Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:5b1f6455797efd3fe520d7cdd8f9df6d842ea309f1bb0e210cbd7229d152f2e0","observation_id":"91a6550e-513c-4888-a00f-c9acf6625cfd","resolution":{"observed_at":"2026-08-08T16:44:01.498358Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.04594","last_updated":"2024-12-19T11:04:20Z","snapshot_observed_at":"2026-08-16T13:27:37.548821Z","submitted_at":"2024-08-08T17:10:16Z","title":"Img-Diff: Contrastive Data Synthesis for Multimodal Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.04594","snapshot_observed_at":"2026-08-08T16:44:00.638413Z","title":"Img-diff: Contrastive data synthesis for multimodal large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.638413Z"},"links":{"cited_paper":"/paper/2408.04594","citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:4082478e5ea0ee26255666e76547e62f01a55a794cd2441fde1515cf30858e7c","observation_id":"6ecfd414-2ba0-4de2-bf97-3d0dddc260f4","resolution":{"observed_at":"2026-08-08T16:44:00.638413Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.485265Z","title":"Scaling up gans for text-to-image synthesis","venue":null,"work_id":"e942bd3c-abc6-4bdf-8d6c-0a43ba981896","year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.641275Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:9552a3079f0893bc5535ce4ec476c754184b77c447072cfdabb5205b59812fd6","observation_id":"aed093c5-0b3b-40e6-bdfe-868d319afe40","resolution":{"observed_at":"2026-08-08T16:44:01.489242Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:00.644543Z","title":"Elucidating the design space of diffusion-based generative models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.644543Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:41b90ce128c1959fde27a3d48f6751f02489f7ee333c6bdd92ad26628c2ec0e9","observation_id":"9c492cb8-26fb-4499-8b84-e5fc3b593915","resolution":{"observed_at":"2026-08-08T16:44:00.644543Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.470260Z","title":"Code: Contrasting self-generated description to combat hallucination in large multi-modal models","venue":null,"work_id":"7eb6e026-9bdd-4c87-a155-d890f0554692","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.647827Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:aed6b398293ef138fc6ee3f52c94891d5a3cd22dec954a2452af881acbf91cbf","observation_id":"2190c6ba-4203-4e93-a6fc-085ef781f455","resolution":{"observed_at":"2026-08-08T16:44:01.474101Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.460247Z","title":"Mitigating object hallucinations in large vision-language models through visual contrastive decoding","venue":null,"work_id":"83d76981-3484-4342-9418-58ad2963352b","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.650845Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:2928aa211b1d3601c3ee14f8f1b638476e91196a64b375edf094059cc1136046","observation_id":"f67589ac-5cd7-48ed-b292-83f3a54c9b4e","resolution":{"observed_at":"2026-08-08T16:44:01.463948Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.452521Z","title":"Your diffusion model is secretly a zero-shot classifier","venue":null,"work_id":"ef1abc0d-f8f9-4335-a6fb-7d5997933170","year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.654259Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:1d68654cd1c61ab6f68ae035d84878a1f332845c52b3a81c45b84178d0b1f17f","observation_id":"35e16345-1cf2-475b-a8f7-6f86cce0ce29","resolution":{"observed_at":"2026-08-08T16:44:01.455305Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.444583Z","title":"Blip-2: Bootstrapping language-image pre-training with frozen image encoders and large language models","venue":null,"work_id":"712d891e-2207-4392-a044-41cfe2c3bebb","year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.657836Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:d53f9c5847afec31ad9dd480aa7fe3636a5be851d1dde71614fc404c766ddb20","observation_id":"61a1908a-d05a-42fa-8c81-222842a53e02","resolution":{"observed_at":"2026-08-08T16:44:01.447286Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.436671Z","title":"Contrastive decoding: Open-ended text generation as optimization","venue":null,"work_id":"44d31430-4e26-4cf1-b7a9-19dbcda3206a","year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.660989Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:695572382a1f467e9568b38ac8bc7a0948ae50911a79447e55574eea73f3bc52","observation_id":"cf61d187-39b4-4030-b849-c2ee7d4d1510","resolution":{"observed_at":"2026-08-08T16:44:01.439454Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.428147Z","title":"Evaluating object hallucination in large vision-language models","venue":null,"work_id":"551578c6-b9ad-45a8-8370-1d343c8a91f6","year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.663316Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:0d8ce8596a55af7e6626c1984355eaa7fafb8e8c978f4decea83bb8f62a9c7c7","observation_id":"eb582d46-4e34-4b48-a117-99f0261a5ce6","resolution":{"observed_at":"2026-08-08T16:44:01.431640Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.417650Z","title":"Microsoft coco: Common objects in context","venue":null,"work_id":"5a635876-57b4-40a1-91be-9ccc3b7b6fd8","year":2014},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.665447Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:cf293fb8a50701ff4dbe43514c344f7041c59e0618ba77048598c8fd40d5a516","observation_id":"6e7569df-f61d-4b00-bda6-d8e456c67b8f","resolution":{"observed_at":"2026-08-08T16:44:01.421091Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.407676Z","title":"Mitigating hallucination in large multi-modal models via robust instruction tuning","venue":null,"work_id":"cb58e348-d7e2-4fe7-bd5b-1668856ec520","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.667552Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:ed0f26f15b321edc07b9c814226f1ccd023cf2ab2fdd4afff452a31fad2bd358","observation_id":"d177ef7f-d365-4cba-a921-249e64018e33","resolution":{"observed_at":"2026-08-08T16:44:01.410818Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.00253","last_updated":"2024-05-06T01:10:01Z","snapshot_observed_at":"2026-08-17T09:37:14.144521Z","submitted_at":"2024-02-01T00:33:21Z","title":"A Survey on Hallucination in Large Vision-Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.00253","snapshot_observed_at":"2026-08-08T16:44:00.669687Z","title":"A survey on hallucination in large vision-language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.669687Z"},"links":{"cited_paper":"/paper/2402.00253","citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:efc065181e6f4353bdcbe24afbd46214fa1cef4a0a7012c2f5326fbe394ed67a","observation_id":"320f0378-3f43-41b3-b916-19cae1108ee4","resolution":{"observed_at":"2026-08-08T16:44:00.669687Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:00.672630Z","title":"Visual instruction tuning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.672630Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:c05463d24662f0ef6f91794c1bacce43eddb91b6a86123ecf7384ef1ca5d3f2d","observation_id":"7cfaccff-463b-4781-9a53-463b36ada50d","resolution":{"observed_at":"2026-08-08T16:44:00.672630Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.393149Z","title":"Improved baselines with visual instruction tuning","venue":null,"work_id":"1881a302-3040-4ba0-8c7a-a14b3006535a","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.675054Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:e10617eb41f0216ea643946d68db35d4a4fa463c4474eba07a5b2d28b560aa65","observation_id":"c39134cd-4f5d-4604-82b5-7a2458c6203b","resolution":{"observed_at":"2026-08-08T16:44:01.396292Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.384081Z","title":"Mmbench: Is your multi-modal model an all-around player? In European Conference on Computer Vision, pp.\\ 216--233","venue":null,"work_id":"87282f05-f945-447b-bb62-5b38ca08ecfa","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.677175Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:6a9dac61feba3f58f27ea39ed1e2c37057514857a5557e52c9cf5ef6a8733344","observation_id":"31cda4ef-7255-4a03-9cec-135c44d309ef","resolution":{"observed_at":"2026-08-08T16:44:01.387590Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.375618Z","title":"Glide: Towards photorealistic image generation and editing with text-guided diffusion models","venue":null,"work_id":"135664d1-3235-4480-a27f-51fbde69dc7c","year":2022},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.679341Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:b3c01c8111ab9ebee15dab6f4bdff7107411539814f491d8493696417d2683d9","observation_id":"771650e1-2a63-4809-a255-8427601bd865","resolution":{"observed_at":"2026-08-08T16:44:01.378724Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.365757Z","title":"SDXL : Improving latent diffusion models for high-resolution image synthesis","venue":null,"work_id":"8af0a3a0-4c48-46f3-9eef-ad5866c4e6eb","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.681676Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:0f9c6aeb595d56ca07960370aa82e8ae38dd08ec64344f27588e98a0978a6565","observation_id":"cb087075-3029-41a4-b132-915ce9d4639b","resolution":{"observed_at":"2026-08-08T16:44:01.369407Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.356399Z","title":"Object hallucination in image captioning","venue":null,"work_id":"175b8af6-3a6c-4d10-a892-6643c538cb6b","year":2018},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.684835Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:3e1baa593d401384fff0a9b4270e46641f771b0d0c25c884ec2a226db788a6b8","observation_id":"fbaacb01-9de3-4588-879c-0fb810ec4467","resolution":{"observed_at":"2026-08-08T16:44:01.359473Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:00.688585Z","title":"High-resolution image synthesis with latent diffusion models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.688585Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:e2f3ebeff2d0994fa2450221599c689636748833ffcf8549fac48f5ea7301e41","observation_id":"89340618-965f-48b1-bc79-e8f15b5eab85","resolution":{"observed_at":"2026-08-08T16:44:00.688585Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:00.691385Z","title":"Photorealistic text-to-image diffusion models with deep language understanding","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.691385Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:b3dc3f125c62b07ac3344508f84b3a19c6cf7bfd4fa928a799b5f86ea8d94fba","observation_id":"57b0e47d-5d00-4813-a822-451e19eec175","resolution":{"observed_at":"2026-08-08T16:44:00.691385Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.337950Z","title":"Stylegan-t: Unlocking the power of gans for fast large-scale text-to-image synthesis","venue":null,"work_id":"8ab54403-90da-4df4-920a-281a22db0209","year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.694849Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:0425ea5972fc7a47cb85fb21fea238e45bcc8157807b0e028020e4024bae7db4","observation_id":"95b856b4-bbde-48e7-ab6b-f37e1d2e8ead","resolution":{"observed_at":"2026-08-08T16:44:01.341473Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:00.697716Z","title":"Laion-5b: An open large-scale dataset for training next generation image-text models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.697716Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:6484e910e9a0735dbe1664c2c87f87eb440aec9559fe75f8de7f353693487142","observation_id":"0dfd7662-8d3a-45d8-a5a0-7616425fbbd8","resolution":{"observed_at":"2026-08-08T16:44:00.697716Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.322109Z","title":"A-okvqa: A benchmark for visual question answering using world knowledge","venue":null,"work_id":"360280ba-585b-425d-b371-98ffdb5cc618","year":2022},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.700479Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:e0de28a1c12f358fa3726361f1bce80a59dd394a8481508f29e70692386b30eb","observation_id":"ff47aeb5-8b47-497c-9dae-12fea1f31023","resolution":{"observed_at":"2026-08-08T16:44:01.326306Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.14525","last_updated":"2023-09-25T20:59:33Z","snapshot_observed_at":"2026-08-15T23:42:34.277075Z","submitted_at":"2023-09-25T20:59:33Z","title":"Aligning Large Multimodal Models with Factually Augmented RLHF","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.14525","snapshot_observed_at":"2026-08-08T16:44:00.703182Z","title":"Aligning large multimodal models with factually augmented rlhf","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.703182Z"},"links":{"cited_paper":"/paper/2309.14525","citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:b3c3195b828d958030dad51670866d560948194c100b6d0ae8c247dbb13d942d","observation_id":"12689926-7fe4-4249-8f5a-a5115c46c111","resolution":{"observed_at":"2026-08-08T16:44:00.703182Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:00.708416Z","title":"Eyes wide shut? exploring the visual shortcomings of multimodal llms","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.708416Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:034dac837bdda614aa39db45f8cd9b323c31a9338f87f037e7a86cf39db9d9bb","observation_id":"47998742-14f3-44b3-b0a6-9bfcc2131023","resolution":{"observed_at":"2026-08-08T16:44:00.708416Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2302.13971","last_updated":"2023-02-27T17:11:15Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-02-27T17:11:15Z","title":"LLaMA: Open and Efficient Foundation Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.13971","snapshot_observed_at":"2026-08-08T16:44:00.711345Z","title":"Llama: Open and efficient foundation language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.711345Z"},"links":{"cited_paper":"/paper/2302.13971","citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:64a3e27f00372664307397dff2da4d48c7dff476d185cb8bcaa97909abbac3f7","observation_id":"6d8ddb12-5a0b-4921-b8bf-e3663d1eee0f","resolution":{"observed_at":"2026-08-08T16:44:00.711345Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.305898Z","title":"Mitigating hallucinations in large vision-language models with instruction contrastive decoding","venue":null,"work_id":"aa0b8ba0-a240-422e-99e8-f114f197d05e","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.714111Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:de96f3614499e5a5008ea92610aa937b998635024ed57c72285232ec7dbe07c3","observation_id":"708014e2-0402-4223-a07f-bc0f5ffd39fb","resolution":{"observed_at":"2026-08-08T16:44:01.309670Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.296015Z","title":"Diffusion models for implicit image segmentation ensembles","venue":null,"work_id":"1cebc28c-3629-44bc-8526-1db1fb4da77f","year":2022},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.717683Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:6f5e17e1949fffcd0a0a69fd03ddacb76a5f2a1cb36b0813bf1d341c8c57d237","observation_id":"5eea5bfc-421a-4037-80bd-425a7bb145c5","resolution":{"observed_at":"2026-08-08T16:44:01.299229Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.17821","last_updated":"2024-12-16T10:27:35Z","snapshot_observed_at":"2026-08-16T13:48:47.702039Z","submitted_at":"2024-05-28T04:41:02Z","title":"RITUAL: Random Image Transformations as a Universal Anti-hallucination Lever in Large Vision Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.17821","snapshot_observed_at":"2026-08-08T16:44:00.721183Z","title":"Ritual: Random image transformations as a universal anti-hallucination lever in lvlms","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.721183Z"},"links":{"cited_paper":"/paper/2405.17821","citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:46d95097bc5065fb140ea79c0fecd6583eb3779fe7f1e2d742bf808bb25a7026","observation_id":"c388cd21-0a60-467d-8f38-d6c238de53e6","resolution":{"observed_at":"2026-08-08T16:44:00.721183Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.287128Z","title":"Evaluating and analyzing relationship hallucinations in lvlms","venue":null,"work_id":"5439f926-af71-4f81-b665-8818db4288cc","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.724528Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:cb17071e0996fce82a0d362f0ddaf926e223290d51bc2fdabf5a0e1f68c5996d","observation_id":"7b6ada38-ac89-44b1-8cbf-efbc8aa79369","resolution":{"observed_at":"2026-08-08T16:44:01.289929Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:00.727907Z","title":"Diffusion models: A comprehensive survey of methods and applications","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.727907Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:ff495375e824875838cb0475b5f3c6248ee4718ce5c6f9242c66f6313d86c3d4","observation_id":"048481bf-088e-42a9-b747-36116b30c26e","resolution":{"observed_at":"2026-08-08T16:44:00.727907Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:00.731303Z","title":"mplug-owl2: Revolutionizing multi-modal large language model with modality collaboration","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.731303Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:3b83184eca96dd750ca497724bc8246853dcd2bc60e3bbac6ffd1f0d23b910a2","observation_id":"17bde78e-0cc9-4bc2-98c5-dd7dfbd1bea1","resolution":{"observed_at":"2026-08-08T16:44:00.731303Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.16045","last_updated":"2024-12-11T02:46:29Z","snapshot_observed_at":"2026-08-16T14:49:14.664860Z","submitted_at":"2023-10-24T17:58:07Z","title":"Woodpecker: Hallucination Correction for Multimodal Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.16045","snapshot_observed_at":"2026-08-08T16:44:00.734462Z","title":"Woodpecker: Hallucination correction for multimodal large language models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.734462Z"},"links":{"cited_paper":"/paper/2310.16045","citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:2706d284566b623ea2bfd0a584f7d74d986225ceecedaea83f7d016aa41a2425","observation_id":"57c92872-401b-4ee9-8e28-37f8045be344","resolution":{"observed_at":"2026-08-08T16:44:00.734462Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.269159Z","title":"Scaling autoregressive models for content-rich text-to-image generation","venue":null,"work_id":"90534b44-47be-465d-8798-b2a49381281f","year":2022},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.738634Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:1000cde25fc925d99b2a14dd061a2e403ef178df348b950cce2c908c436d7d12","observation_id":"53ea1921-d351-4b43-94b7-39c03b3e9838","resolution":{"observed_at":"2026-08-08T16:44:01.271751Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.261098Z","title":"MM -vet: Evaluating large multimodal models for integrated capabilities","venue":null,"work_id":"cb73b81a-52e8-4b04-8832-544b49f4180d","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.741578Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:f41be23a4714e7a98bf8e10cc02473796eaa195dd77803144fa98f0b4d72e912","observation_id":"021f322e-0a72-4af9-86ca-9f8311d4e6be","resolution":{"observed_at":"2026-08-08T16:44:01.264173Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.252039Z","title":"Less is more: Mitigating multimodal hallucination from an EOS decision perspective","venue":null,"work_id":"62616663-e1ac-4143-9e25-17d00d50c6a7","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.744316Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:930700e2b4c7f89fc0a9c037e5ba50063cf20a7dd5f4bbcf93982e2736831e90","observation_id":"957ec9ed-d7bd-4db8-8e73-4c5863d75694","resolution":{"observed_at":"2026-08-08T16:44:01.255563Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.242609Z","title":"Multimodal image synthesis and editing: A survey and taxonomy","venue":null,"work_id":"06626c34-64ee-4fcd-8fe8-4de9939e846d","year":2023},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.747100Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:8134789812d6bd5b9a259fd54d49efc1d962211face0913b018a29c8100bcf3c","observation_id":"cdf92f3d-7837-4771-a7fc-fb9df41249a3","resolution":{"observed_at":"2026-08-08T16:44:01.246325Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.232952Z","title":"Reflective instruction tuning: Mitigating hallucinations in large vision-language models","venue":null,"work_id":"1b934c95-0028-446d-b7ed-19d399e10421","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.749982Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:cc325430393e787d4106fe109b22081fad7118c06f03ee2f70fd1584dbc3a3f0","observation_id":"e95e7954-0c15-4053-b8e8-4d5d90e7774f","resolution":{"observed_at":"2026-08-08T16:44:01.236346Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.222592Z","title":"Fact-and-reflection ( F a R ) improves confidence calibration of large language models","venue":null,"work_id":"115708ea-7030-423d-9588-5d852924f709","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.752859Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:1b7e0dba1705c65096cefba3f6ba0df99e3ae27bbebd45c231ac9d0f292c64e8","observation_id":"a921a59b-2919-4be5-a373-c8aafe383fde","resolution":{"observed_at":"2026-08-08T16:44:01.225942Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.212040Z","title":"Analyzing and mitigating object hallucination in large vision-language models","venue":null,"work_id":"9fcc0bbf-6949-46eb-91a8-8c6cb97c7e33","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.756000Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:10534cccfe57a6f66807fa8337fbb6a487d79628ed22a37bd02a7b62a4e67248","observation_id":"90c2a16c-4df1-480c-88ba-13e7215990dd","resolution":{"observed_at":"2026-08-08T16:44:01.215815Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.203157Z","title":"Mini GPT -4: Enhancing vision-language understanding with advanced large language models","venue":null,"work_id":"b4ea15ca-cf51-4735-89ca-6a350a07b4aa","year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.760268Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:214e9088b22d58549b948f9268da88c0068e09bdfd69455fb8f102b10d2f2b43","observation_id":"41eb729b-11fd-489c-8a88-c45cf276d2c0","resolution":{"observed_at":"2026-08-08T16:44:01.206653Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:01.191877Z","title":"Dm-gan: Dynamic memory generative adversarial networks for text-to-image synthesis","venue":null,"work_id":"b9d62aba-d854-4975-b8b7-011e54a65e61","year":2019},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.763394Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:b43b7196137b3f79477e29bb02686eeec047d6b49f3f869d894422d73045d90b","observation_id":"3fa799e2-f4df-4241-892d-0eab0f771d06","resolution":{"observed_at":"2026-08-08T16:44:01.196962Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:00.767233Z","title":"@esa (Ref","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.767233Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:e69b22c491b34170a17660ee9defab431292291fb62d3eb768247ac8509f25d1","observation_id":"88ec5a9b-6782-4ae0-aba6-94b8fd825f08","resolution":{"observed_at":"2026-08-08T16:44:00.767233Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:00.771115Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":71,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.771115Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:b8cd4e7b5a9571d385061b77fae21cdc9b4a4d36d98c7ea7ca7b316b7e49adf7","observation_id":"caa1733b-0062-4d2e-acf5-8851adac619d","resolution":{"observed_at":"2026-08-08T16:44:00.771115Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T16:44:00.773941Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models","version":2},"reference_index":72,"source":"arxiv_source","source_observed_at":"2026-08-08T16:44:00.773941Z"},"links":{"citing_paper":"/paper/2502.06130"},"observation_digest":"sha256:0c32d7ff99bc2165b0a4a90fd26632cf6bf2a8d8d8ef2e8cfd52529a6f530809","observation_id":"c3f53df6-c17c-4d83-95c8-c0f4eca9a0bb","resolution":{"observed_at":"2026-08-08T16:44:00.773941Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2502.06130","last_updated":"2025-09-09T18:19:31Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-15T05:30:21.338398Z","submitted_at":"2025-02-10T03:43:55Z","title":"Self-Correcting Decoding with Generative Feedback for Mitigating Hallucinations in Large Vision-Language Models"},"reference_resolution":{"displayed":72,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":27,"verified_exact":0,"verified_fuzzy":44},"total_outbound_references":72},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"thesis":"As of 17 August 2026, this Paper Citation Record lists 72 of 72 outbound references and 11 inbound Pith citation observations for arXiv:2502.06130."}