{"as_of":"2026-08-07T12:14:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:4e61d2c8179b73ff11f8cdaeed69a5cdf8daf63e49fbbc88cc11133591ff93c7","coverage":[{"denominator":45,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":45,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T12:46:19.590250Z","state":"measured"},{"denominator":45,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":45,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2507.21507/citation-record","integrity":"/paper/2507.21507/integrity","json":"/paper/2507.21507/citation-record.json","paper":"/paper/2507.21507"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.503857Z","title":"A revisit of sparse coding based anomaly detection in stacked rnn framework,","venue":null,"work_id":"a9fa63fe-abee-4212-bb5b-66c8644ee89d","year":2017},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.388869Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:b1cb87d6305469852c15ef23447299ccd3264389f6defb7358c57d4ed4eb7100","observation_id":"ac12060c-6718-4695-bedc-93e5dc3fcfc1","resolution":{"observed_at":"2026-08-06T12:46:20.509015Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.487857Z","title":"Few-shot scene-adaptive anomaly detection,","venue":null,"work_id":"97301873-b2e5-4fd7-b4dd-4a10905743c7","year":2020},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.394430Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:e524f07c0cfcb8e7e613abca617bddc72bb77db8f1cfd13071657cffeec001b4","observation_id":"5a8b79ec-0323-412a-b236-0717a3b63f83","resolution":{"observed_at":"2026-08-06T12:46:20.493041Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.470827Z","title":"Fastano: Fast anomaly de- tection via spatio-temporal patch transformation,","venue":null,"work_id":"73d34635-ce28-42a6-b71e-775bab7ae0ae","year":2022},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.399487Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:4de7c404098a4c11522a327c4dc0f89e6d2a156307bf626b8da103c44e4d7f34","observation_id":"18aeca26-8405-48b8-a8be-9e8d7ca20ea2","resolution":{"observed_at":"2026-08-06T12:46:20.476885Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.454415Z","title":"Improved anomaly detection in surveillance videos with multiple probabilistic models inference,","venue":null,"work_id":"de6f04fc-4eae-4618-9a25-426c28ee2723","year":2022},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.404334Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:fdcab733ff68df61dfeab2d9f9369747e5bc45dc4d4035331a06afb6d09b3e0f","observation_id":"04046219-d9d5-48de-b35e-bd51785d3161","resolution":{"observed_at":"2026-08-06T12:46:20.459525Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.437786Z","title":"Novel applications for vae-based anomaly detection systems,","venue":null,"work_id":"8ccea294-b434-4c45-bdba-c37089ae8f19","year":2022},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.408860Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:3b9d5e45d504ccf9ffae06e5a152e9f7c5c5bf047fab5e1505a4b2860902f23e","observation_id":"1b22a331-538c-4092-9ce5-a41989811fba","resolution":{"observed_at":"2026-08-06T12:46:20.443709Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.421239Z","title":"Making reconstruction-based method great again for video anomaly detection,","venue":null,"work_id":"10343709-b5ef-49b5-b48d-1c82baff6dd4","year":2022},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.413518Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:84b45c36e6692eb3894e5e3eac94eba1da715da7b84074d619ec661cead1d98e","observation_id":"82ce177e-ef5a-44af-964f-a258dd10653e","resolution":{"observed_at":"2026-08-06T12:46:20.426911Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.404091Z","title":"A hybrid video anomaly detection framework via memory-augmented flow reconstruction and flow-guided frame prediction,","venue":null,"work_id":"f7164a6e-bc2e-4f71-81fe-4302c3cc5375","year":2021},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.418851Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:65f1296d0ab270aed17a76fc39512c7c9eac823dfb20ae3a414e31f7df79dc14","observation_id":"a975ff9a-9b8d-47de-94a9-e26dae670d17","resolution":{"observed_at":"2026-08-06T12:46:20.409616Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.387269Z","title":"Video anomaly detection by solving decoupled spatio-temporal jigsaw puz- zles,","venue":null,"work_id":"6dcaa0d1-ce3b-4976-87f0-6db5e1c225e1","year":2022},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.423364Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:a7e90a7e68530ca0d7175cfb2be4353a0a19dd7727617d3e1b328bdb9a28adde","observation_id":"39e091db-f781-4076-80d3-64a9030b0320","resolution":{"observed_at":"2026-08-06T12:46:20.393313Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2302.05160","last_updated":"2023-02-10T10:39:40Z","snapshot_observed_at":"2026-07-06T14:50:26.417454Z","submitted_at":"2023-02-10T10:39:40Z","title":"Dual Memory Units with Uncertainty Regulation for Weakly Supervised Video Anomaly Detection","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.05160","snapshot_observed_at":"2026-08-06T12:46:19.427868Z","title":"Dual memory units with uncertainty reg- ulation for weakly supervised video anomaly detection,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.427868Z"},"links":{"cited_paper":"/paper/2302.05160","citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:c2eab50683820bc9c6f4e82f7fd5f678145f5031e7b2b1a41744b26df455aca6","observation_id":"554243e5-94e0-44c2-8cea-db98cf105c51","resolution":{"observed_at":"2026-08-06T12:46:19.427868Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:19.432860Z","title":"Generative cooperative learning for unsupervised video anomaly detection,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.432860Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:75b32323a568087cadee4daab9f909cc28722253eb1bbe3bf253f930b598d30f","observation_id":"1a3b3bff-03f7-45dc-b857-0d2efce11947","resolution":{"observed_at":"2026-08-06T12:46:19.432860Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.353375Z","title":"Clip-tsa: Clip-assisted temporal self-attention for weakly-supervised video anomaly detection,","venue":null,"work_id":"2dd826b8-7c9f-4453-b406-e67888996a2f","year":2023},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.437746Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:330e5e8e4fe37610f04897d0119d1ce5f449c7afcb84fcfbdb51aadac2b9cf10","observation_id":"40f0cde2-439e-4b44-b2fb-57ba8d8be48d","resolution":{"observed_at":"2026-08-06T12:46:20.359908Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11681","last_updated":"2023-12-15T09:42:25Z","snapshot_observed_at":"2026-07-06T16:09:12.888489Z","submitted_at":"2023-08-22T14:58:36Z","title":"VadCLIP: Adapting Vision-Language Models for Weakly Supervised Video Anomaly Detection","version":3},"cited_work":{"arxiv_id":"2308.11681","doi":null,"metadata_source":"pith","pith_arxiv_id":"2308.11681","snapshot_observed_at":"2026-08-06T12:46:19.916306Z","title":"VadCLIP: Adapting Vision-Language Models for Weakly Supervised Video Anomaly Detection","venue":"cs.CV","work_id":"06aaba1a-c0ba-40e4-958a-06e49c7113ba","year":2023},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.442269Z"},"links":{"cited_paper":"/paper/2308.11681","citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:5ba1fdac5bf9da692cfe9b9e1352ca342ff8983b71a9f7256a6401a6b701b584","observation_id":"299a4172-19c4-47f2-aeb5-75bf12ab6a03","resolution":{"observed_at":"2026-08-06T12:46:19.921853Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.337409Z","title":"Catching both gray and black swans: Open-set supervised anomaly detection,","venue":null,"work_id":"904281f0-5675-4584-bee9-9a4e6b192619","year":2022},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.447042Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:d298a9261354364d345b55165e5d0d0dbe9319c597735312be73587cc010ea64","observation_id":"39df8caf-0576-4e03-8206-934cb341be41","resolution":{"observed_at":"2026-08-06T12:46:20.342710Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.319461Z","title":"Towards open set video anomaly detection,","venue":null,"work_id":"29a974e9-cbf4-4d96-80fa-1571a344525b","year":2022},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.451306Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:f6d0ec1e668fbd5c18abeb8bb37e3ab297c404792d9bf3400236d233ae962905","observation_id":"61fa1ca9-b828-42dd-b865-915229987ebb","resolution":{"observed_at":"2026-08-06T12:46:20.326124Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.12790","last_updated":"2024-03-17T06:10:25Z","snapshot_observed_at":"2026-07-06T16:35:40.224690Z","submitted_at":"2023-10-19T14:47:11Z","title":"Anomaly Heterogeneity Learning for Open-set Supervised Anomaly Detection","version":3},"cited_work":{"arxiv_id":"2310.12790","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.12790","snapshot_observed_at":"2026-08-06T12:46:19.892220Z","title":"Anomaly Heterogeneity Learning for Open-set Supervised Anomaly Detection","venue":"cs.CV","work_id":"789a31c4-a178-4471-aa67-fa0cc346cc70","year":2023},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.455616Z"},"links":{"cited_paper":"/paper/2310.12790","citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:3a69bc246340dbd4036bcc9e008a08d5d4a2b84f0c7d28abd3dac186937ff528","observation_id":"4566ac5b-eaa7-482f-966f-a8f25cd6222f","resolution":{"observed_at":"2026-08-06T12:46:19.898284Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.298739Z","title":"Hawk: Learning to understand open-world video anomalies,","venue":null,"work_id":"9092b467-05ba-47aa-b5d1-596339825fa4","year":2024},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.460149Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:13bd4dc60c75cf323466c06cbecab8369ad04807fbb319ee785b3047aa0d9464","observation_id":"58a65e5c-e619-401a-aa0b-fbe55d3b5257","resolution":{"observed_at":"2026-08-06T12:46:20.306475Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.01095","last_updated":"2025-03-31T20:17:27Z","snapshot_observed_at":"2026-07-06T19:59:56.438747Z","submitted_at":"2024-12-02T04:10:14Z","title":"VERA: Explainable Video Anomaly Detection via Verbalized Learning of Vision-Language Models","version":3},"cited_work":{"arxiv_id":"2412.01095","doi":null,"metadata_source":"pith","pith_arxiv_id":"2412.01095","snapshot_observed_at":"2026-08-06T12:46:19.866042Z","title":"VERA: Explainable Video Anomaly Detection via Verbalized Learning of Vision-Language Models","venue":"cs.AI","work_id":"54145c29-eed0-4195-a826-5a03be4d4d00","year":2024},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.464603Z"},"links":{"cited_paper":"/paper/2412.01095","citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:4a4227ebc15a868279d564cde933d5237549d061cf31da3142dbf4e63292c2bb","observation_id":"6aee25f6-b75d-424d-9a2c-0abb9f796861","resolution":{"observed_at":"2026-08-06T12:46:19.873190Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10326","last_updated":"2025-03-24T20:26:56Z","snapshot_observed_at":"2026-07-06T18:31:13.039726Z","submitted_at":"2024-06-14T17:59:01Z","title":"VANE-Bench: Video Anomaly Evaluation Benchmark for Conversational LMMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10326","snapshot_observed_at":"2026-08-06T12:46:19.469435Z","title":"Vane- bench: Video anomaly evaluation benchmark for conversational lmms,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.469435Z"},"links":{"cited_paper":"/paper/2406.10326","citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:690e971aa25ae3dfb667bba82eab82b26ce6e9da9a641c0773cfa97bcffa7176","observation_id":"b5e45a4b-3d81-4071-bb94-13d2861335c3","resolution":{"observed_at":"2026-08-06T12:46:19.469435Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.281265Z","title":"Suvad: Semantic understanding based video anomaly detection using mllm,","venue":null,"work_id":"db7dd2f7-41ff-405c-b8dc-76b4865546b1","year":2025},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.474029Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:0a0898656bdb4e1c8a536b12e38a96a624df00dc1f9debde722ec81dae197e5b","observation_id":"093cad5b-8c9d-4dc3-b713-689261d3ca22","resolution":{"observed_at":"2026-08-06T12:46:20.286687Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.262018Z","title":"Harness- ing large language models for training-free video anomaly detection,","venue":null,"work_id":"44a57b45-473d-41d0-aaa0-983ec15f261b","year":2024},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.478389Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:9174e3ad0f0aa84ac052e0c9ddd469be4718bfd2021201119b8cae9dcdc69626","observation_id":"8209b336-85d1-48b2-86a6-c1f27909158c","resolution":{"observed_at":"2026-08-06T12:46:20.268430Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:19.482625Z","title":"Anyanomaly: Zero-shot customizable video anomaly detection with lvlm,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.482625Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:9b5df8c501144ba548fbedf88d83225ed16ed28e79f8f3c625258428bf4c500e","observation_id":"ba845819-6f83-4568-b80d-cbc6615c2846","resolution":{"observed_at":"2026-08-06T12:46:19.482625Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.05702","last_updated":"2024-01-11T07:09:44Z","snapshot_observed_at":"2026-07-06T17:14:07.183954Z","submitted_at":"2024-01-11T07:09:44Z","title":"Video Anomaly Detection and Explanation via Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.05702","snapshot_observed_at":"2026-08-06T12:46:19.486962Z","title":"Video anomaly detection and explanation via large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.486962Z"},"links":{"cited_paper":"/paper/2401.05702","citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:b88c44f72979a797d37ba98acdcc369529fad8fd8a12607795ee66ad78fcf392","observation_id":"9539f4a0-c7cd-4201-8999-c33910401d0f","resolution":{"observed_at":"2026-08-06T12:46:19.486962Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.244126Z","title":"Scene-adaptive svad based on multi- modal action-based feature extraction,","venue":null,"work_id":"f240a3e6-d61b-41fa-a555-421458c67ad3","year":2024},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.492025Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:469c7c0053a7ad9c2352516d7d8d9e29e8e7b8e67e83ad12060aa4bc92ba4025","observation_id":"aae62810-4675-4b8d-b9c8-e977376942d1","resolution":{"observed_at":"2026-08-06T12:46:20.250213Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.226991Z","title":"Uncovering what why and how: A comprehensive bench- mark for causation understanding of video anomaly,","venue":null,"work_id":"2b26284f-e159-4f17-8ea7-1cac75f62b4e","year":2024},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.496190Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:db89362eb9ff138a1ae0da3f40c41e5fc480a4955f74128cc89478a2b85cec52","observation_id":"5e197aa8-cec0-4542-a957-dd78bfdee563","resolution":{"observed_at":"2026-08-06T12:46:20.232939Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.10122","last_updated":"2024-10-01T12:07:31Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-11-16T10:59:44Z","title":"Video-LLaVA: Learning United Visual Representation by Alignment Before Projection","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.10122","snapshot_observed_at":"2026-08-06T12:46:19.500298Z","title":"Video-llava: Learning united visual representation by alignment before projection,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.500298Z"},"links":{"cited_paper":"/paper/2311.10122","citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:313dc407f7d7a75caf66e66162f7369ea7bf7581b77ea3c116346e7333660c0e","observation_id":"7b7ea1c3-6087-4195-86ac-7339836a56ec","resolution":{"observed_at":"2026-08-06T12:46:19.500298Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.14485","last_updated":"2024-12-10T12:45:31Z","snapshot_observed_at":"2026-07-06T19:19:28.953455Z","submitted_at":"2024-09-22T15:13:31Z","title":"Video-XL: Extra-Long Vision Language Model for Hour-Scale Video Understanding","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.14485","snapshot_observed_at":"2026-08-06T12:46:19.505318Z","title":"Video-xl: Extra-long vision language model for hour-scale video understanding,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.505318Z"},"links":{"cited_paper":"/paper/2409.14485","citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:6cb2f430168fd56353a06cdb0da5006ad6c95bb7e1de42ccd2b510b0aac8fba6","observation_id":"e27b375b-5317-4540-b42d-d6fe98846f67","resolution":{"observed_at":"2026-08-06T12:46:19.505318Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.05424","last_updated":"2024-06-10T01:36:53Z","snapshot_observed_at":"2026-07-06T15:40:24.127663Z","submitted_at":"2023-06-08T17:59:56Z","title":"Video-ChatGPT: Towards Detailed Video Understanding via Large Vision and Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.05424","snapshot_observed_at":"2026-08-06T12:46:19.510049Z","title":"Video-chatgpt: Towards detailed video understanding via large vision and language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.510049Z"},"links":{"cited_paper":"/paper/2306.05424","citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:42b0e77efbeb7aaa7fc65277b330501a9d3008e99c57a75c139a6d8896e807fe","observation_id":"1be65b63-89a6-412e-a69d-a2328243b8cd","resolution":{"observed_at":"2026-08-06T12:46:19.510049Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.208716Z","title":"Vtimellm: Empower llm to grasp video moments,","venue":null,"work_id":"56fe8e4e-fa1a-477e-a3d5-3e0e5489c249","year":2024},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.514501Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:ccbd813906d90c5ad31d15cc6ec072acaae8afce632d403613d4f9e5d609bfdf","observation_id":"8d54d34d-1ca0-4c07-aa58-5a99f6d09b43","resolution":{"observed_at":"2026-08-06T12:46:20.214791Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.14178","last_updated":"2024-03-29T08:13:38Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-04-27T13:27:01Z","title":"mPLUG-Owl: Modularization Empowers Large Language Models with Multimodality","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.14178","snapshot_observed_at":"2026-08-06T12:46:19.518886Z","title":"mplug-owl: Modularization empowers large language models with multimodality,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.518886Z"},"links":{"cited_paper":"/paper/2304.14178","citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:31e8863e7d76efa7a18d3f39c6fb9540263e6c8530ecc45da419618d8e86edd6","observation_id":"fa0a3282-54b3-4d91-bdec-8d477e67ba6a","resolution":{"observed_at":"2026-08-06T12:46:19.518886Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.190859Z","title":"Region-aware pretraining for open- vocabulary object detection with vision transformers,","venue":null,"work_id":"09dd7e5b-884e-449a-a8e3-7981c9e594e8","year":2023},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.523276Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:171600645b93d9115630b1efbb5fa3ef9c6dfee9b7edc5a0f3f5e57ff19cef4a","observation_id":"6f64ac80-4c6c-4f06-9083-e5928a15529d","resolution":{"observed_at":"2026-08-06T12:46:20.196190Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.173482Z","title":"Dual discriminator generative adver- sarial network for video anomaly detection,","venue":null,"work_id":"36252b1e-504e-4c00-8883-012cabb5131d","year":2020},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.527684Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:52dd6df760717154e8a517ff13e521b392613c4769abf97e54a491efb34d2a29","observation_id":"19044621-7c8a-463f-a97c-bfe626fec1cc","resolution":{"observed_at":"2026-08-06T12:46:20.179147Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.148933Z","title":"Spatio- temporal autoencoder for video anomaly detection,","venue":null,"work_id":"0e26fcdc-686e-498f-90b1-7d5d2a572567","year":2017},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.531898Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:53f6677958e6c3c957f691b1480015ec60697fb87c53525d937d185a38c8c390","observation_id":"b9fc6f55-8248-4f78-93d8-20f207322a55","resolution":{"observed_at":"2026-08-06T12:46:20.162025Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.133559Z","title":"Synthetic temporal anomaly guided end-to-end video anomaly detection,","venue":null,"work_id":"8c772bb0-feb7-49ad-b6f1-949d7798fc4c","year":2021},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.536317Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:0de510d4b3e2e922e51634d8e72813538189bbb1f54bc520b8f004f78393a3fb","observation_id":"a7550118-fd99-440e-84f0-e4d2ad878f2e","resolution":{"observed_at":"2026-08-06T12:46:20.138239Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.115423Z","title":"Chaotic invariants of lagrangian particle trajectories for anomaly detection in crowded scenes,","venue":null,"work_id":"8cecf98f-d86c-438b-82b4-20e5dbae605a","year":2010},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.540805Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:3d43be4f2fb86c7be93cd4d27c58cc2e2c829f0bb9a04a0dcb4aecbc986343a7","observation_id":"9f9a2346-78cc-44b4-bfeb-44368b1d9590","resolution":{"observed_at":"2026-08-06T12:46:20.121786Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.099638Z","title":"Gods: Generalized one-class discriminative subspaces for anomaly detection,","venue":null,"work_id":"c6376239-e7ea-4d30-8324-01a7b8dcc245","year":2019},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.545373Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:88bbd585f130633ea2b513dc62f05408dc87eeb4a8cffa846225b336c10739e1","observation_id":"a817c356-0f1b-4827-a21c-d647c39dc5f3","resolution":{"observed_at":"2026-08-06T12:46:20.104214Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.084444Z","title":"A deep one-class neural network for anomalous event detection in complex scenes,","venue":null,"work_id":"1153b363-404a-47ff-a202-b34a57e50cb3","year":2019},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.549681Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:0bc3e93c42d30db29b58e81455e837969af78b1f6f32e09ec0c966af7f79df2e","observation_id":"1823a973-2d84-45b6-ac03-885c3d56a176","resolution":{"observed_at":"2026-08-06T12:46:20.089408Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.066843Z","title":"Ubnormal: New benchmark for supervised open-set video anomaly detection,","venue":null,"work_id":"5091edf1-dfa8-4197-b706-73f384f27f32","year":2022},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.553951Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:75a80ed4585edfc51d4fd755ea0ff2f71b739f2d28c8a980976bc4f2c0c2761d","observation_id":"650b361d-5862-48cf-a780-896c9b9fc585","resolution":{"observed_at":"2026-08-06T12:46:20.072715Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.048994Z","title":"Follow the rules: reasoning for video anomaly detection with large language models,","venue":null,"work_id":"3dbbaee2-d800-4eeb-806e-516399a00ae0","year":2024},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.558247Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:7d949e3288adbfb800da878f50c1f9c12964efb2cef6565c66740899a4bb2d7d","observation_id":"5633eb77-bab4-47a9-93a7-aba0e2a1c60e","resolution":{"observed_at":"2026-08-06T12:46:20.054809Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12235","last_updated":"2024-06-29T08:15:27Z","snapshot_observed_at":"2026-07-06T18:32:40.207270Z","submitted_at":"2024-06-18T03:19:24Z","title":"Holmes-VAD: Towards Unbiased and Explainable Video Anomaly Detection via Multi-modal LLM","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12235","snapshot_observed_at":"2026-08-06T12:46:19.562834Z","title":"Holmes-vad: Towards unbiased and explain- able video anomaly detection via multi-modal llm,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.562834Z"},"links":{"cited_paper":"/paper/2406.12235","citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:7edec46a6293a6dd05a5d6b99c688c1cf617666056d8f279f0896ef09c5e63c4","observation_id":"224307fc-6724-453f-b374-c1d7b610c401","resolution":{"observed_at":"2026-08-06T12:46:19.562834Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.031788Z","title":"Real-world anomaly detection in surveillance videos,","venue":null,"work_id":"b27547df-5f1e-46b1-8756-036d2a75ee45","year":2018},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.567552Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:08015cc0a673517db50154464557559c24518332c73eb8b877f046d4b92eb762","observation_id":"ef02caf3-3a76-4263-a7e3-e1e24f8721cb","resolution":{"observed_at":"2026-08-06T12:46:20.037031Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:20.014750Z","title":"Not only look, but also listen: Learning multimodal violence detection under weak supervision,","venue":null,"work_id":"c7dd4049-c937-4101-b816-b92999d01abe","year":2020},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.572158Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:010512a71e5834c90c4a75c547ce0ff4932a23f7ca0f9bee4532cd306326e8a4","observation_id":"e6bbbf2a-fb6d-4080-9687-20946354e797","resolution":{"observed_at":"2026-08-06T12:46:20.020135Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:19.997786Z","title":null,"venue":null,"work_id":"da0c73ec-b225-4002-9112-4d17e1c32c6d","year":2020},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.576850Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:4566b2ffc96a38501b113253be381cbb1c5b663c04e9c707feab247c168b28ea","observation_id":"2fa4d571-fdfd-4c2a-a79e-4e56d0137b2b","resolution":{"observed_at":"2026-08-06T12:46:20.002913Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:19.980947Z","title":"Rouge: A package for automatic evaluation of sum- maries,","venue":null,"work_id":"687ac5b3-eb24-4b6d-a66d-5bbc163e5690","year":2004},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.581316Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:f1c07c0d571cc4d271ed7a29ce8650640448b5cd7ffceea39548e72d5cc8e1bc","observation_id":"ae7d06d2-0912-4def-9079-052ebd045f2f","resolution":{"observed_at":"2026-08-06T12:46:19.985779Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:19.585712Z","title":"Bleu: a method for automatic evaluation of machine translation,","venue":null,"work_id":null,"year":2002},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.585712Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:29285290b9d8d28016d77d0643c5f98d13ebabfa74e02425550f0392a29cd069","observation_id":"d6209e0d-38b5-40d3-b08f-b3bfa45a5150","resolution":{"observed_at":"2026-08-06T12:46:19.585712Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T12:46:19.953604Z","title":"Meteor: An automatic metric for mt evalua- tion with improved correlation with human judgments,","venue":null,"work_id":"401dbdf4-3237-40d2-ae1d-c12cfcb20fba","year":2005},"citing_paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-06T12:46:19.590250Z"},"links":{"citing_paper":"/paper/2507.21507"},"observation_digest":"sha256:2b4f5533621792728a5808d4ebb921d68d42c7ca2dd278dde10e5f64c4c04953","observation_id":"2e1f3604-6bdc-4e9b-878e-5ab3bc32fa62","resolution":{"observed_at":"2026-08-06T12:46:19.958557Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2507.21507","last_updated":"2025-07-29T05:17:48Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-06T12:46:18.741677Z","submitted_at":"2025-07-29T05:17:48Z","title":"VAGU & GtS: LLM-Based Benchmark and Framework for Joint Video Anomaly Grounding and Understanding"},"reference_resolution":{"displayed":45,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":12,"verified_exact":3,"verified_fuzzy":30},"total_outbound_references":45},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 45 of 45 outbound references and 0 inbound Pith citation observations for arXiv:2507.21507."}