{"as_of":"2026-08-09T12:37:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:62373076efa35540d94f7e8c5aec410f57ff50dacccc86c38f10fa5471836f02","coverage":[{"denominator":48,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":48,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T23:59:12.892263Z","state":"measured"},{"denominator":48,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":48,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2508.04549/citation-record","integrity":"/paper/2508.04549/integrity","json":"/paper/2508.04549/citation-record.json","paper":"/paper/2508.04549"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:21.143172Z","title":"https://gemini.google.com/","venue":null,"work_id":"e447c228-b724-483b-beb4-51e5f3632a5f","year":null},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:08.457823Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:d9d5a217518afaaffc7a27f65acc12ee00ad6d734712f6b4a619ba467804b7e6","observation_id":"3005b4d7-50e3-40e1-a180-effb7e66f0be","resolution":{"observed_at":"2026-08-05T23:59:21.216559Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:21.001361Z","title":"https://openai.com/index/gpt-4-1/","venue":null,"work_id":"af30f7c9-a180-4799-b96e-03c0da9efc0c","year":null},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:08.538538Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:37b7f62840a1581ec43852918b19d70f34aa9932da3d187564ab0b30fe7c9f0d","observation_id":"6d0b0fca-e099-42ae-98c1-710d3be5cefb","resolution":{"observed_at":"2026-08-05T23:59:21.057480Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:20.777288Z","title":"https://hailuoai.video","venue":null,"work_id":"52c61344-9dd0-4481-b422-b5375cc07fb1","year":null},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:08.642103Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:4d0a93660680d83179ca6642777f054e466795f55aa42d36d2cfa29b278f942a","observation_id":"9d487554-0ba4-4ab7-a90e-ee5f34c248b3","resolution":{"observed_at":"2026-08-05T23:59:20.925571Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:20.540130Z","title":"https://klingai.com/","venue":null,"work_id":"4c6ecea0-20c7-4756-bf43-5cfa07c7a7e8","year":null},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:08.761578Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:5ce0c86c7a1ef7de1bb441e51d872ca666a64b2cb7b9080bea669cdc499c673f","observation_id":"148fc3cd-d57e-4443-b597-f64d620329c3","resolution":{"observed_at":"2026-08-05T23:59:20.643523Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:20.287333Z","title":null,"venue":null,"work_id":"d6e9b3a6-2654-44ae-b90b-5974344e56a4","year":2016},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:08.858832Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:58c539f4d576031d98443d13cc1b7413c679f261ce1faa116753847cb264b7f7","observation_id":"58c8c4f6-f6da-4505-85af-52032d506dcf","resolution":{"observed_at":"2026-08-05T23:59:20.403246Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.09754","last_updated":"2025-04-03T14:52:24Z","snapshot_observed_at":"2026-08-08T13:57:39.347745Z","submitted_at":"2024-12-12T23:10:54Z","title":"ViCaS: A Dataset for Combining Holistic and Pixel-level Video Understanding using Captions with Grounded Segmentation","version":3},"cited_work":{"arxiv_id":"2412.09754","doi":null,"metadata_source":"pith","pith_arxiv_id":"2412.09754","snapshot_observed_at":"2026-08-05T23:59:14.564678Z","title":"ViCaS: A Dataset for Combining Holistic and Pixel-level Video Understanding using Captions with Grounded Segmentation","venue":"cs.CV","work_id":"4645774c-a0e4-4943-af29-61f6a7d8dd37","year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:08.911675Z"},"links":{"cited_paper":"/paper/2412.09754","citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:6cd9bddeb96bff4adf4c78626324d8902fa2c120cfc185cff79a291039e60c44","observation_id":"19af39be-5fed-4716-8f46-c09c679da4b2","resolution":{"observed_at":"2026-08-05T23:59:14.681879Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.12966","last_updated":"2023-10-13T02:41:28Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-08-24T17:59:17Z","title":"Qwen-VL: A Versatile Vision-Language Model for Understanding, Localization, Text Reading, and Beyond","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.12966","snapshot_observed_at":"2026-08-05T23:59:08.971756Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:08.971756Z"},"links":{"cited_paper":"/paper/2308.12966","citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:a0f80440f33950c16e7355d06e77aba78b4a031d909de43751ae35c3391ae77b","observation_id":"dbdc221f-ef67-4303-b778-16c5b018f8e1","resolution":{"observed_at":"2026-08-05T23:59:08.971756Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:09.007152Z","title":null,"venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:09.007152Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:586bd8e5825d5755a9458dfdb23b0e30b4b711f00a1ea0f9de4bfa67ea45c973","observation_id":"a7d07017-6cb8-4a5b-81f6-8899cabb14ff","resolution":{"observed_at":"2026-08-05T23:59:09.007152Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.21201","last_updated":"2025-03-19T15:11:51Z","snapshot_observed_at":"2026-08-07T17:39:16.043241Z","submitted_at":"2025-02-28T16:18:57Z","title":"The PanAf-FGBG Dataset: Understanding the Impact of Backgrounds in Wildlife Behaviour Recognition","version":3},"cited_work":{"arxiv_id":"2502.21201","doi":null,"metadata_source":"pith","pith_arxiv_id":"2502.21201","snapshot_observed_at":"2026-08-05T23:59:14.269207Z","title":"The PanAf-FGBG Dataset: Understanding the Impact of Backgrounds in Wildlife Behaviour Recognition","venue":"cs.CV","work_id":"03562822-ee59-4e06-9dfc-7c59d7bcce29","year":2025},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:09.158746Z"},"links":{"cited_paper":"/paper/2502.21201","citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:d437a13ba767408a270be554e30aa1683d8ddc6aed9376159df1e933d11136d0","observation_id":"b6ba1d4c-0e90-4a5d-a6fa-0dfb7c4568b3","resolution":{"observed_at":"2026-08-05T23:59:14.388620Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:20.046044Z","title":null,"venue":null,"work_id":"5ca4ab45-e45b-49be-9b55-a666440a4900","year":null},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:09.257963Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:7574209d2e9a09853ba6211cc1ea6f33e57335795e29bee45adb904cc1e38968","observation_id":"11857b71-9241-47a1-b8af-61426309f7dd","resolution":{"observed_at":"2026-08-05T23:59:20.157060Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:19.628301Z","title":null,"venue":null,"work_id":"2a5247cf-1115-4885-9725-b5607ee9ff6a","year":2018},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:09.469385Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:f010a9407720d364d556e826723e3d2d03d89f5120f1a8838a8dc9df023702a9","observation_id":"761a8f9a-a433-456b-aedf-4e116ec301ea","resolution":{"observed_at":"2026-08-05T23:59:19.764158Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:19.437907Z","title":null,"venue":null,"work_id":"bc7a1585-94d5-4218-af01-1b5661bc3c7e","year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:09.594320Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:8f6f9888e5f032ff644fe591e0900e332e302680134da1e38b66fbcc91a62989","observation_id":"e44f5a71-4440-4009-a47d-3ee7cafa41e6","resolution":{"observed_at":"2026-08-05T23:59:19.544012Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:19.158671Z","title":null,"venue":null,"work_id":"8b011984-fad4-47cf-adeb-4c126f4091d0","year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:09.695931Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:329f3515ddfa6dd5d423e091ecdf4d1b3f86d43d81194149c107ceced0484d87","observation_id":"fd54cf59-9906-4d3b-8012-311a656bff2b","resolution":{"observed_at":"2026-08-05T23:59:19.271724Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:18.936026Z","title":null,"venue":null,"work_id":"5655bd40-fa63-490e-a59a-3dc423671c8d","year":2019},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:09.814937Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:27d2963305e0aa7f251fd92952c4f83f0d37c06ad5d42504518c81e72f79f320","observation_id":"32b2758e-83ca-4fce-8235-c9220a465b56","resolution":{"observed_at":"2026-08-05T23:59:19.034022Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:09.872872Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:09.872872Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:78d0affbe80718e4f1236c794d0ec829c8bacad2e8cb50890d0d0827d9fe94a8","observation_id":"ac698ceb-0e0a-4d13-b4a6-4ea367bb2ee6","resolution":{"observed_at":"2026-08-05T23:59:09.872872Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2106.04632","last_updated":"2021-08-18T21:55:27Z","snapshot_observed_at":"2026-08-01T20:01:49.108442Z","submitted_at":"2021-06-08T18:34:21Z","title":"VALUE: A Multi-Task Benchmark for Video-and-Language Understanding Evaluation","version":2},"cited_work":{"arxiv_id":"2106.04632","doi":null,"metadata_source":"pith","pith_arxiv_id":"2106.04632","snapshot_observed_at":"2026-08-05T23:59:13.903272Z","title":"VALUE: A Multi-Task Benchmark for Video-and-Language Understanding Evaluation","venue":"cs.CV","work_id":"b6ba9b25-4372-40bc-898b-b650c382c92a","year":2021},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:10.065554Z"},"links":{"cited_paper":"/paper/2106.04632","citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:363e930ba1627485f0fb9d74736d37a507ae7a6ad941266deeeed8ce55d4a02c","observation_id":"c3b64f8d-5b7b-4368-a21d-d868e28e907b","resolution":{"observed_at":"2026-08-05T23:59:14.075656Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:18.667374Z","title":"In Proceedings of the IEEE/CVF international conference on computer vision","venue":null,"work_id":"6ddaa554-6c67-4116-a538-0539947fe97f","year":null},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:09.957738Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:653ff00f9cdae534a802ee80dcf62a0b53c4a60a6ae990dabde34e445217bbf1","observation_id":"d68c9b34-aee0-43a5-a44c-62258c8b9da0","resolution":{"observed_at":"2026-08-05T23:59:18.795953Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:18.115380Z","title":null,"venue":null,"work_id":"01537bdc-82f6-4214-b40e-0ad4a7290ec3","year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:10.244625Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:9765e171d077883e19975dfdc91e24e0f2732f596494724eab59e20ac72922e3","observation_id":"adaf9ef5-1eed-4926-a254-4d86464bde92","resolution":{"observed_at":"2026-08-05T23:59:18.227463Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:18.365438Z","title":null,"venue":null,"work_id":"a8404f8a-a96f-46dc-9a87-1bdc517654c9","year":2023},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:10.148977Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:3a2b2a9be866d79e62b28c483112887bbce2cacea048ea2a02015b67e5c5f349","observation_id":"4f271a68-6e17-4961-9268-4f8908c44483","resolution":{"observed_at":"2026-08-05T23:59:18.500783Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:17.788813Z","title":"Belongie, James Hays, Pietro Perona, Deva Ramanan, Piotr Dollár, and C","venue":null,"work_id":"50f2e9b8-ec34-4f10-90de-88bfda1aa672","year":2014},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:10.398441Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:e8d719b691976430a704b55e5c6eb9a1cd34eca09ef1c23f5a97eba92b07b0b7","observation_id":"1599b09d-c4ea-4a03-82d3-3de2892e27c7","resolution":{"observed_at":"2026-08-05T23:59:17.908722Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:10.280127Z","title":null,"venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:10.280127Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:8888a39aea951fabe8ee3ede84b168121568aeab5ef5c07b810c526e061bf4ca","observation_id":"b477c48a-5db7-4e09-b2ef-c4f693d3cab4","resolution":{"observed_at":"2026-08-05T23:59:10.280127Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.23715","last_updated":"2025-03-31T04:30:34Z","snapshot_observed_at":"2026-08-07T16:26:13.956216Z","submitted_at":"2025-03-31T04:30:34Z","title":"HOIGen-1M: A Large-scale Dataset for Human-Object Interaction Video Generation","version":1},"cited_work":{"arxiv_id":"2503.23715","doi":null,"metadata_source":"pith","pith_arxiv_id":"2503.23715","snapshot_observed_at":"2026-08-05T23:59:13.683321Z","title":"HOIGen-1M: A Large-scale Dataset for Human-Object Interaction Video Generation","venue":"cs.CV","work_id":"5489e02c-1085-45d9-be43-784b0864c0b2","year":2025},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:10.568638Z"},"links":{"cited_paper":"/paper/2503.23715","citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:e2c64d9c6a15f3f38dad12ee2491b2fdcde3a0e9fc3fba472dc05a5124c16c16","observation_id":"61233888-4c42-42fe-9565-541b7737d3b8","resolution":{"observed_at":"2026-08-05T23:59:13.765492Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:10.490588Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:10.490588Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:869bb54ea0822f5fb4eb7c188d8422482b0dad5d43c48c1c3e42fc44c97b8ddd","observation_id":"275e17bc-3eb8-4fd6-8801-f78294b9ff4e","resolution":{"observed_at":"2026-08-05T23:59:10.490588Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:17.499519Z","title":null,"venue":null,"work_id":"2ac509de-efed-498e-b798-17c8e96d2033","year":2025},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:10.787238Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:ced23c39f7cf49f3a8660946fe318bd4f7ba08ed292d4c4ff8db48792420fcae","observation_id":"c9efe0ae-4089-4fbd-b94d-180079fb61a8","resolution":{"observed_at":"2026-08-05T23:59:17.643932Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.05499","last_updated":"2024-07-19T06:00:41Z","snapshot_observed_at":"2026-07-06T15:00:58.804337Z","submitted_at":"2023-03-09T18:52:16Z","title":"Grounding DINO: Marrying DINO with Grounded Pre-Training for Open-Set Object Detection","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.05499","snapshot_observed_at":"2026-08-05T23:59:10.688285Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:10.688285Z"},"links":{"cited_paper":"/paper/2303.05499","citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:64665e24c0e42782b7ee4fd9c72785c52b6d38c1f841ea266b0effe0269dc532","observation_id":"9cf0746b-47fa-49c4-aac5-ab2e222850ff","resolution":{"observed_at":"2026-08-05T23:59:10.688285Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.04923","last_updated":"2025-03-25T10:08:13Z","snapshot_observed_at":"2026-08-04T02:03:57.818642Z","submitted_at":"2024-11-07T17:59:27Z","title":"VideoGLaMM: A Large Multimodal Model for Pixel-Level Visual Grounding in Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.04923","snapshot_observed_at":"2026-08-05T23:59:10.935802Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:10.935802Z"},"links":{"cited_paper":"/paper/2411.04923","citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:784dd02cb243841d0914e866af1fd2c23f974afc4f2d221383ab89b6bde07807","observation_id":"388bf383-4714-45a5-bdcd-8e3fbec35f84","resolution":{"observed_at":"2026-08-05T23:59:10.935802Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.08027","last_updated":"2025-02-27T07:06:50Z","snapshot_observed_at":"2026-08-03T17:41:21.699474Z","submitted_at":"2024-07-10T20:10:56Z","title":"Fish-Vista: A Multi-Purpose Dataset for Understanding & Identification of Traits from Images","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.08027","snapshot_observed_at":"2026-08-05T23:59:10.874559Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:10.874559Z"},"links":{"cited_paper":"/paper/2407.08027","citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:d64095310f2ae3d8cc2f1c99c625e4e1962626588fc3f242594a6db80d9a8a0f","observation_id":"b355bb82-c631-4f45-b1a4-aabfaf5eba8d","resolution":{"observed_at":"2026-08-05T23:59:10.874559Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:11.076632Z","title":null,"venue":null,"work_id":null,"year":2002},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:11.076632Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:aba06895c6fd3b2fd66e35c10d0d1fa83cdef30cd366c15ed7272ece9fb57b59","observation_id":"1da09c90-9dab-43cd-a993-ac96f63e6596","resolution":{"observed_at":"2026-08-05T23:59:11.076632Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.20781","last_updated":"2025-03-26T17:59:02Z","snapshot_observed_at":"2026-08-07T16:34:49.318092Z","submitted_at":"2025-03-26T17:59:02Z","title":"BASKET: A Large-Scale Video Dataset for Fine-Grained Skill Estimation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.20781","snapshot_observed_at":"2026-08-05T23:59:10.991805Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:10.991805Z"},"links":{"cited_paper":"/paper/2503.20781","citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:cf31b03b9003b99d55dce5cc4f90b9bbd0e1d4428f47ea278d2c7c768c706fec","observation_id":"7931311f-e614-4f5e-864a-4d447b6f1efa","resolution":{"observed_at":"2026-08-05T23:59:10.991805Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.00714","last_updated":"2024-10-28T16:37:57Z","snapshot_observed_at":"2026-07-06T18:55:41.459417Z","submitted_at":"2024-08-01T17:00:08Z","title":"SAM 2: Segment Anything in Images and Videos","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.00714","snapshot_observed_at":"2026-08-05T23:59:11.284876Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:11.284876Z"},"links":{"cited_paper":"/paper/2408.00714","citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:30b41d2062bff74be099dade2564a3e85e61c328c71e76b234768261ccaf5d68","observation_id":"e41563d6-78b2-4d02-999c-fc4c2602c6c0","resolution":{"observed_at":"2026-08-05T23:59:11.284876Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:17.158424Z","title":null,"venue":null,"work_id":"f7494ef7-b1e5-4424-8550-fa4848044dfd","year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:11.182132Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:e5c363a9a2990fce758517dc1ebb7b1fc0d1e43ffa3fd8c40d44f5defd202970","observation_id":"c0495d66-c417-48a2-86fc-50e0b2cc316f","resolution":{"observed_at":"2026-08-05T23:59:17.320714Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1811.00347","last_updated":"2018-12-07T07:03:52Z","snapshot_observed_at":"2026-08-04T13:33:48.224677Z","submitted_at":"2018-11-01T12:47:11Z","title":"How2: A Large-scale Dataset for Multimodal Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1811.00347","snapshot_observed_at":"2026-08-05T23:59:11.441634Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:11.441634Z"},"links":{"cited_paper":"/paper/1811.00347","citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:6d326f65dc36665d393b262de79ef25843554212a4c004b403ef35c6217fd51a","observation_id":"00cfd453-54a1-493e-b928-1f6f5c951b94","resolution":{"observed_at":"2026-08-05T23:59:11.441634Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.15683","last_updated":"2024-12-13T20:09:27Z","snapshot_observed_at":"2026-08-05T02:18:35.972794Z","submitted_at":"2024-12-13T20:09:27Z","title":"Results of the 2024 Video Browser Showdown","version":1},"cited_work":{"arxiv_id":"2502.15683","doi":null,"metadata_source":"pith","pith_arxiv_id":"2502.15683","snapshot_observed_at":"2026-08-05T23:59:13.445207Z","title":"Results of the 2024 Video Browser Showdown","venue":"cs.MM","work_id":"7b8a6bf9-a69c-48b6-ba14-cc9d5ddf18a7","year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:11.356423Z"},"links":{"cited_paper":"/paper/2502.15683","citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:98170772132edbd925f89a0fbdd0d7ac197b2904fbb9cd406b9206ac4d7c4d89","observation_id":"e58d9132-ae86-4160-8e83-f69274387249","resolution":{"observed_at":"2026-08-05T23:59:13.516714Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:16.611788Z","title":null,"venue":null,"work_id":"b6b27260-f772-476a-ab52-b27ab48d6e19","year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:11.686484Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:13f1509cdd666196be2f3e1a2ab49c138679163d94446633e91c0ed0048a4565","observation_id":"d18d5931-5102-4e53-9e9b-7fc7788ee817","resolution":{"observed_at":"2026-08-05T23:59:16.716622Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:16.868757Z","title":null,"venue":null,"work_id":"00e3d995-449a-441d-a00e-a5f97758bc53","year":2020},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:11.575271Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:055fc40a6950de4716efdd20a19175bd604124aa14b3e4d501341098caf53d06","observation_id":"fc3c9781-a038-4d52-bf2a-0966420f3453","resolution":{"observed_at":"2026-08-05T23:59:16.995598Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:16.068479Z","title":null,"venue":null,"work_id":"09229def-1952-40a1-9bfd-f92712ea18ee","year":2025},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:11.854961Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:a143f4c53f401b6cbf068f4e05ba10f6bca0429cc7a027ccbb0fd4cf798f1ffc","observation_id":"2909d544-ffeb-4fc5-a8f2-126bc4ed81f7","resolution":{"observed_at":"2026-08-05T23:59:16.195689Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:16.317897Z","title":null,"venue":null,"work_id":"81630964-2d2f-49ae-b1c0-8b42a8d474ac","year":2023},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:11.757097Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:5f81dda22ad4186d839b330b0d056808408acfdd0a9c157613752a24e098a52b","observation_id":"e3c955e8-c118-4bae-b66d-e54257576a9c","resolution":{"observed_at":"2026-08-05T23:59:16.460955Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:12.018176Z","title":null,"venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:12.018176Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:8e39dc85710335b1e713d93f6c53623dc89381c9949c7361aecf127d9562ce45","observation_id":"d7e1a6d2-e20b-4eec-9662-c933c8556d37","resolution":{"observed_at":"2026-08-05T23:59:12.018176Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:15.768872Z","title":null,"venue":null,"work_id":"17a7288b-ffde-4d55-b71b-4fae01a7333f","year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:11.920940Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:4472d7e2f539430b2e78c35aa806ee472f5220f2bb3e9a3341b9ecf4c2932523","observation_id":"6ecc33ac-3986-4b70-9279-f34485eea203","resolution":{"observed_at":"2026-08-05T23:59:15.890330Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2410.20436","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:13.147831Z","title":null,"venue":null,"work_id":"15c1419e-4aa7-4798-b8e2-d1df58c62f09","year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:12.227807Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:e412ee189daff334bf4888f25575cc461813637e9517cb9a8f57bfe4f6208215","observation_id":"7aacfaff-b2a6-44dc-afa9-b2d6714fa687","resolution":{"observed_at":"2026-08-05T23:59:13.286627Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.08260","last_updated":"2025-04-26T17:58:24Z","snapshot_observed_at":"2026-08-01T11:16:30.883699Z","submitted_at":"2024-10-10T17:57:49Z","title":"Koala-36M: A Large-scale Video Dataset Improving Consistency between Fine-grained Conditions and Video Content","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.08260","snapshot_observed_at":"2026-08-05T23:59:12.137674Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:12.137674Z"},"links":{"cited_paper":"/paper/2410.08260","citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:130b91c293eae0d781ee66f249015ebf1ef193b90076c25b7d8640c18233e957","observation_id":"42952be8-2088-43cf-b97b-1b13745d190b","resolution":{"observed_at":"2026-08-05T23:59:12.137674Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.15262","last_updated":"2025-03-31T02:52:56Z","snapshot_observed_at":"2026-08-08T15:17:25.217856Z","submitted_at":"2024-11-22T10:25:08Z","title":"MovieBench: A Hierarchical Movie Level Dataset for Long Video Generation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.15262","snapshot_observed_at":"2026-08-05T23:59:12.480209Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:12.480209Z"},"links":{"cited_paper":"/paper/2411.15262","citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:fd9c2dbf76c357fcc0e1bafd7d1e29cc14683b8c751cfc1a5b8fae515859657e","observation_id":"9549863e-4fa7-4e0d-8380-f9d535c78013","resolution":{"observed_at":"2026-08-05T23:59:12.480209Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2112.04888","last_updated":"2021-12-09T13:21:26Z","snapshot_observed_at":"2026-08-03T23:10:26.584438Z","submitted_at":"2021-12-09T13:21:26Z","title":"A Bilingual, OpenWorld Video Text Dataset and End-to-end Video Text Spotter with Transformer","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2112.04888","snapshot_observed_at":"2026-08-05T23:59:12.338103Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:12.338103Z"},"links":{"cited_paper":"/paper/2112.04888","citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:d6ddd7e4e0767c7e94918e8cc249ff4ecb94d375095a944fe0678ab724ef096d","observation_id":"18477974-b0e3-4edd-9838-927a936e7fdc","resolution":{"observed_at":"2026-08-05T23:59:12.338103Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:15.467373Z","title":null,"venue":null,"work_id":"e6478063-74e1-433a-aa15-2910440ab098","year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:12.670130Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:e03a4484208a8f89231c9caea50f55ac7983e1034b4e923ccf464216e181a8ca","observation_id":"16a3186a-60c7-4737-9de0-ec3a14da51e6","resolution":{"observed_at":"2026-08-05T23:59:15.598618Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.16994","last_updated":"2024-04-29T14:52:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-04-25T19:29:55Z","title":"PLLaVA : Parameter-free LLaVA Extension from Images to Videos for Video Dense Captioning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.16994","snapshot_observed_at":"2026-08-05T23:59:12.577012Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:12.577012Z"},"links":{"cited_paper":"/paper/2404.16994","citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:4a5a87f35c78aa53c3306a0fff4acf6f51b935a59a47efbc2dee3bb520de4610","observation_id":"6ed5d231-08e0-4f6d-a6a8-d7ea0e082437","resolution":{"observed_at":"2026-08-05T23:59:12.577012Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:14.882028Z","title":null,"venue":null,"work_id":"27b36fd2-b23c-41d8-9cf1-7d000039b1ff","year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:12.892263Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:df836339dd891d8aae98c210931260ec27eb74b6199f59cbc956ae51dd7c7980","observation_id":"9c5f42d2-2332-4fa9-bb4e-ee5685830a7b","resolution":{"observed_at":"2026-08-05T23:59:15.008249Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:15.178986Z","title":null,"venue":null,"work_id":"7452fe88-d5f6-4ecf-808a-e14098870343","year":2024},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:12.758400Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:fabbf2e16b49f80249f6fb87f2644ac0181687e6c342a7a77394cb11fc20e2a7","observation_id":"e5becb23-88af-4858-ae86-d96093c77273","resolution":{"observed_at":"2026-08-05T23:59:15.343402Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T23:59:19.858738Z","title":"In Proceedings of the IEEE/CVF international conference on computer vision","venue":null,"work_id":"a19c45fe-2e08-461a-9dd9-eb24b299adbb","year":null},"citing_paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning","version":3},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-05T23:59:09.362654Z"},"links":{"citing_paper":"/paper/2508.04549"},"observation_digest":"sha256:734ddd1c86da54525d62f6c9e138c4b912b47ff3b23223300e6e188c384a21e9","observation_id":"7e3702b2-7c6d-49e4-934f-f889567ea9af","resolution":{"observed_at":"2026-08-05T23:59:19.931143Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2508.04549","last_updated":"2025-09-01T08:36:15Z","latest_version":3,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-06T22:01:15.202286Z","submitted_at":"2025-08-06T15:34:24Z","title":"MSC: A Marine Wildlife Video Dataset with Grounded Segmentation and Clip-Level Captioning"},"reference_resolution":{"displayed":48,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":35,"verified_exact":6,"verified_fuzzy":7},"total_outbound_references":48},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 48 of 48 outbound references and 0 inbound Pith citation observations for arXiv:2508.04549."}