{"as_of":"2026-08-10T16:13:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:bf47ca6b23e4d7dae89843abc44952b4ad9d764dbbf32b3fe09b99b43e32e759","coverage":[{"denominator":51,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":51,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-30T10:45:35.568212Z","state":"measured"},{"denominator":51,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":51,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2606.20140/citation-record","integrity":"/paper/2606.20140/integrity","json":"/paper/2606.20140/citation-record.json","paper":"/paper/2606.20140"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T13:06:15.663166Z","title":"In: Computer Vision– ECCV 2020: 16th European Conference, Glasgow, UK, August 23–28, 2020, Proceedings, Part XIV 16","venue":null,"work_id":"7a143c25-dda8-4e2e-a74c-fb2f6abbc481","year":2020},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:6b4ec9e5424d85711c42af52bbb2f51770e26f6584ab5acf6bdd4dcb3ce5340f","observation_id":"d635651b-e7d8-438a-9dc2-40574ee76107","resolution":{"observed_at":"2026-07-09T13:06:15.664457Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T13:06:15.656371Z","title":"In: European conference on computer vision","venue":null,"work_id":"8427217e-cdc8-4ca7-b6b2-28e482c176c1","year":2020},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:5a44bd0b44501bbe0453d91e6e1ec5f04a1bad2c9c6a3dc5ff2a88de906d8a55","observation_id":"24b6c561-6235-47bf-bebb-8cba59ea1ddf","resolution":{"observed_at":"2026-07-09T13:06:15.657642Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2112.10764","last_updated":"2021-12-20T18:59:59Z","snapshot_observed_at":"2026-08-09T21:31:49.496159Z","submitted_at":"2021-12-20T18:59:59Z","title":"Mask2Former for Video Instance Segmentation","version":1},"cited_work":{"arxiv_id":"2112.10764","doi":null,"metadata_source":"pith","pith_arxiv_id":"2112.10764","snapshot_observed_at":"2026-07-10T03:06:43.601051Z","title":"Mask2former for video instance segmentation","venue":"cs.CV","work_id":"b1131c21-30c6-499a-a423-eef12442cb11","year":2021},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"cited_paper":"/paper/2112.10764","citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:a22f6059d05be5dfcdb8d6cdc8e90b3803f0687203394777ac64c03913f6d1ac","observation_id":"46758175-4b06-49e8-a5a7-1e1fe1319b2f","resolution":{"observed_at":"2026-06-30T10:54:37.012591Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T13:06:15.654555Z","title":"In: Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (June 2020)","venue":null,"work_id":"6e700bea-7447-4826-a573-53fae8187bf5","year":2020},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:6e6ab67bdc38655d8d4df33aba6b468ea127e8b402b08283dda557db00d99ed1","observation_id":"6afe94e9-bf93-43f8-abab-c84804589974","resolution":{"observed_at":"2026-07-09T13:06:15.655789Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T13:06:15.647683Z","title":null,"venue":null,"work_id":"77fe0f79-c849-463b-9033-24c4b8744f45","year":2022},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:cbd4fb4a2d51529fc25a258551a0cc2294681e221f03b824d76e59e3f915a9d3","observation_id":"aa478d70-ceb0-432f-97f0-cc53704997f7","resolution":{"observed_at":"2026-07-09T13:06:15.648776Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T13:06:15.667080Z","title":null,"venue":null,"work_id":"53fb8938-8201-4c03-91a1-0e7b4da030cd","year":2021},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:79ffa2c389e5aa3850be301521c9612a68b49bb287528c229ba017831cbda99c","observation_id":"8ccef650-8855-4126-98c1-c585a17e3b60","resolution":{"observed_at":"2026-07-09T13:06:15.668234Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T13:06:15.651499Z","title":"In: Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","venue":null,"work_id":"f6f4d588-13ae-4978-8bb2-4a74d543f892","year":2023},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:3cddce65f9bd33f37bd4ae9749b974ae201ef1c1c9fe9127631b42abb70fbd80","observation_id":"6c4ccc24-7a33-43dd-bc07-2fdfead4b65b","resolution":{"observed_at":"2026-07-09T13:06:15.653465Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T13:06:15.658248Z","title":"In: Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition","venue":null,"work_id":"cbe21aab-ec86-45e4-bb0b-10828aa49670","year":2023},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:8f5ff25c3316aa4057da8643a25fc9bc2137d72c9c05ccbfe2c69afd8c56cfa1","observation_id":"8c4abee5-1de5-4603-b957-acb54612cb64","resolution":{"observed_at":"2026-07-09T13:06:15.659465Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2024.336107","doi":"10.1109/tcsvt.2024.3361076","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"IEEE Transactions on Circuits and Systems for Video Technology pp","venue":"IEEE Transactions on Circuits and Systems for Video Technology","work_id":"9d025732-ee44-45b3-b676-e5c70381fa2a","year":2024},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:8fe409ba44a721339d8e3bc5677c0a1b4217ff1b39d9c9d9a976e634afa42449","observation_id":"510c1fa1-d04c-4b93-9bd3-fca38afddb09","resolution":{"observed_at":"2026-06-30T10:54:35.914385Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T13:06:15.660196Z","title":null,"venue":null,"work_id":"29670082-d4be-49f8-9714-4c7d456fdc4c","year":2023},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:12dd9d94688b083465837dbf836d7d0198b165d5e0400fc753288c59912c7aaf","observation_id":"6157e3a2-43a6-4b85-af3c-f71339a031c6","resolution":{"observed_at":"2026-07-09T13:06:15.661734Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T13:06:15.665122Z","title":"In: Proceedings of the AAAI Conference on Artificial Intelligence","venue":null,"work_id":"a051837d-248a-490c-9680-e1d37ec67774","year":2021},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:e115af80a6d39354aea3f752d571441ebd236270d998f0062e9030b42c8627ec","observation_id":"de4e9e78-4fa4-48ae-a4f1-b21fe7af54a4","resolution":{"observed_at":"2026-07-09T13:06:15.666491Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.870307Z","title":"In: Proceedings of the IEEE International Conference on Computer Vision","venue":null,"work_id":"a17d81b0-4b72-4981-8781-803b33a91359","year":2017},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:c22f577ed795254968eaf72bb117b1a65fb635cc99fa5a8f71438d7b1f4ae83a","observation_id":"b5327cb4-ac73-4129-b33e-1e403e0b889e","resolution":{"observed_at":"2026-07-09T12:56:15.871539Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.882923Z","title":"In: Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition","venue":null,"work_id":"c598e211-17f1-4526-acba-0af7342e9ad8","year":2022},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:d5883c7d1c49f1d76bbf74aecfdfa53495e7b05e90d8aca49ff83126bb3748e1","observation_id":"a206314a-9377-40d6-ab40-6854c0154d7d","resolution":{"observed_at":"2026-07-09T12:56:15.884115Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17096","last_updated":"2023-05-26T17:10:24Z","snapshot_observed_at":"2026-08-01T19:08:48.964131Z","submitted_at":"2023-05-26T17:10:24Z","title":"GRAtt-VIS: Gated Residual Attention for Auto Rectifying Video Instance Segmentation","version":1},"cited_work":{"arxiv_id":"2305.17096","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2305.17096","snapshot_observed_at":"2026-07-04T03:09:30.488075Z","title":"arXiv preprint arXiv:2305.17096 (2023)","venue":null,"work_id":"ff95ac47-ebaf-4194-ab89-84ab26334b34","year":2023},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"cited_paper":"/paper/2305.17096","citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:a125ee93db36ab6c64ac105d4d12cf54b94b379d5635f10cd3147433e7da6b9e","observation_id":"5c75a7ed-b8fb-414a-bc0b-7640f0e530e1","resolution":{"observed_at":"2026-06-30T10:54:37.007153Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T13:06:15.668860Z","title":"Advances in Neural Information Processing Systems35, 19370–19383 (2022)","venue":null,"work_id":"8a5af04a-a49f-46c7-8f75-115ba259c2fd","year":2022},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:3c3c153ab21a0e4da9dc6101f1c3f2209da8c3fa44b2a694b6ecfa156f0cc09f","observation_id":"7a0a48c4-8bba-4c3b-b93d-381d805b6c8a","resolution":{"observed_at":"2026-07-09T13:06:15.670109Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.876624Z","title":"In: Proceedings of the IEEE International Conference on Computer Vision (ICCV) (Oct 2017)","venue":null,"work_id":"2d816aa9-de6c-4b74-a120-d0c17decb4a0","year":2017},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:4abf5cef444161e85f0cd51db625928066b6a6614f9c2286a15c99529e5c876d","observation_id":"f0c658c6-ce50-4114-8ce3-58c6366635e4","resolution":{"observed_at":"2026-07-09T12:56:15.877727Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.878272Z","title":"In: Pro- ceedings of the IEEE Conference on Computer Vision and Pattern Recognition (CVPR) (June 2016)","venue":null,"work_id":"2a170ef7-d14d-4efd-902b-84a3e417f7e3","year":2016},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:39bd1f62230be2d70b81d5df74590e3c40d3d4527e0bc012d10cd7b306442921","observation_id":"67b2f85c-8f9f-42ac-8392-e4ce07f8b9b9","resolution":{"observed_at":"2026-07-09T12:56:15.879408Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.881437Z","title":"In: CVPR (2023)","venue":null,"work_id":"2c810a0e-6754-430c-b9fb-ce3a2d73f404","year":2023},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:bd6155267ef12e13c4dd7c765e079b8892bbc46b40440346cdbcbb01a32cf357","observation_id":"a1abb3f3-9892-4156-8e39-106134deebff","resolution":{"observed_at":"2026-07-09T12:56:15.882411Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.871789Z","title":"In: Advances in Neural Information Processing Systems (2022) 10","venue":null,"work_id":"a1740e35-bf21-495d-8b23-a728faad8d3e","year":2022},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:d6693e4aacf6b6e8a7e900d12b7afb3cc7c656dbb2e8c635dd0dad8e111aa167","observation_id":"987c8db0-9371-48d1-87b3-2c72c28a5c79","resolution":{"observed_at":"2026-07-09T12:56:15.872926Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.868228Z","title":null,"venue":null,"work_id":"a5b46b29-3baf-4a8f-a42e-94170249a524","year":2022},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:ec5942d4e13420648d65888944a37a5ecbea8f07036087a7cbe573a5952ada31","observation_id":"55a089d4-2887-433e-8b3c-84366a182e58","resolution":{"observed_at":"2026-07-09T12:56:15.869266Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.866523Z","title":"Advances in Neural Information Processing Systems34, 13352– 13363 (2021)","venue":null,"work_id":"74a5fc6b-b183-4547-8d80-575eaf710d2d","year":2021},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:c3494dd9734935f5e431fbc76796dbe686c7416caf1e04f8c30d00160a05ec9f","observation_id":"809811f1-0c7e-4096-bbcc-bf9b01063c23","resolution":{"observed_at":"2026-07-09T12:56:15.867675Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.869919Z","title":"In: Computer Vision – ECCV 2022 Workshops: Tel Aviv, Israel, October 23–27, 2022, Proceedings, Part IV (2023)","venue":null,"work_id":"144d02df-b3dc-4dc7-85a3-3bc380faeede","year":2022},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:6519ee2cc197575e24ab02f1a0026b10587210686cd351b9483013330348c61b","observation_id":"00981fe2-07e2-466f-b852-1713430f29b7","resolution":{"observed_at":"2026-07-09T12:56:15.871205Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.873454Z","title":"In: Proceedings of the IEEE International Conference on Computer Vision","venue":null,"work_id":"6ebc0dbd-0c38-4bad-bac8-46d3a8df5b9a","year":2017},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:2e0a7b27ba5650cf814d22c64c62d877ebcd7640e4434c7fd7cae8eddb02389f","observation_id":"24272057-f678-4764-aab9-4fcd589370ab","resolution":{"observed_at":"2026-07-09T12:56:15.874556Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.879929Z","title":"In: CVPR (2023)","venue":null,"work_id":"12c879c1-e7ca-4574-83cd-88e372cf5dbf","year":2023},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:608582b889bc0c23a38561f6543eff765b0b778a100e6de17899d70e7d8d8d88","observation_id":"22bf91ff-7caa-458a-91d4-d877652c848a","resolution":{"observed_at":"2026-07-09T12:56:15.880929Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2208.10547","last_updated":"2022-08-22T18:54:18Z","snapshot_observed_at":"2026-07-06T13:44:22.049225Z","submitted_at":"2022-08-22T18:54:18Z","title":"InstanceFormer: An Online Video Instance Segmentation Framework","version":1},"cited_work":{"arxiv_id":"2208.10547","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2208.10547","snapshot_observed_at":"2026-07-04T03:09:30.491357Z","title":"arXiv preprint arXiv:2208.10547 (2022)","venue":null,"work_id":"159625e9-a0c1-4701-a3e8-bbe22f486003","year":2022},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"cited_paper":"/paper/2208.10547","citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:648a106d234364c2c70aecf6de3c29df9759f2ff34c27177d6065eb0117cf384","observation_id":"ae686b2b-ebfa-4041-9d70-f74cca7f5dd0","resolution":{"observed_at":"2026-06-30T10:54:37.009864Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.875093Z","title":"In: Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition","venue":null,"work_id":"7ec01396-e2b5-444f-af64-acc9b5426d6c","year":2020},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:787edab1f91ccce356a6f0468fc57ccbc2bca18970f90f9f67b0b25449548655","observation_id":"1002d36c-92d8-40f6-9a92-4567cea7289d","resolution":{"observed_at":"2026-07-09T12:56:15.876129Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T13:06:15.670718Z","title":"In: Computer Vision–ECCV 2014: 13th European Conference, Zurich, Switzerland, September 6-12, 2014, Proceedings, Part V 13","venue":null,"work_id":"bff0f42d-ae0e-4670-b447-3282c82e6c94","year":2014},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:3902e5d98eb3f392d20ef048e10724fd2c4c4f9727f22f4dcb2dab8290a76ef9","observation_id":"f551c391-3b5f-4749-9c39-04c00a70db53","resolution":{"observed_at":"2026-07-09T13:06:15.672426Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.856446Z","title":"In: Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition","venue":null,"work_id":"ba82fa61-a8f7-4a92-8072-21042d21d9d5","year":2021},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:2035ce466515ac144b345080d00dc9f9bb707254aeb57368a011192210922692","observation_id":"8b4c5020-46d3-4a67-8e23-6284839bf648","resolution":{"observed_at":"2026-07-09T12:56:15.857764Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.858357Z","title":"In: Proceedings of the IEEE/CVF international conference on computer vision","venue":null,"work_id":"c4a4e078-f728-47d7-96ae-62dd4bda5dca","year":2021},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:9c9ffb670c33a9b2c74e3f646f0518032c7ad164f3de9d1ee4e2f969937a1bbf","observation_id":"5c85c876-f3be-4df3-8688-0a79e55fef9b","resolution":{"observed_at":"2026-07-09T12:56:15.859617Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.860194Z","title":"In: Proceedings of the IEEE conference on computer vision and pattern recognition","venue":null,"work_id":"1043657c-607a-47b7-80a0-5f9f583a8973","year":2018},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:767cf60af7ed6bbe38637756383b6d957ace40d737a973971c3855a1ff8b1cc7","observation_id":"8af24e73-a2cc-46df-baf6-f343f9186ecd","resolution":{"observed_at":"2026-07-09T12:56:15.861496Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.864796Z","title":"In: Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition","venue":null,"work_id":"5091e1cd-51a6-4304-8a30-398eafaf0b9f","year":2020},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:97ed696037086877252f40383e7073f51fecbbf89abb7ba5753a8efd5014c7f3","observation_id":"ff5ff4e8-27a0-4248-bb69-ef23ba645be6","resolution":{"observed_at":"2026-07-09T12:56:15.865939Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.854628Z","title":"In: Thirty-fifth Conference on Neural Information Processing Systems Datasets and Benchmarks Track (Round","venue":null,"work_id":"4a1a1f8e-4eeb-4590-8913-bd77b6b99810","year":2021},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:242b679cc2bdc9dd38b5d2b4b9ffceeec64a92e1e670ae19f45dc225a8f6fcd1","observation_id":"d3c5a83a-0322-4c17-963a-2523008f6b6d","resolution":{"observed_at":"2026-07-09T12:56:15.855795Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.848987Z","title":null,"venue":null,"work_id":"c3a0bf1d-be96-4367-8065-bea1ba4a658a","year":2021},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:8b66f3992ab0f2edccafa3dc8ef5a6ffc68d0c47fd927c6e6318d784b88e19a1","observation_id":"9ba9db2b-4476-43ba-955c-4ba7a40d4a72","resolution":{"observed_at":"2026-07-09T12:56:15.850076Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.850802Z","title":"In: Com- puter Vision–ECCV 2020: 16th European Conference, Glasgow, UK, August 23–28, 2020, Proceedings, Part I 16","venue":null,"work_id":"bf5ed2b2-858a-4603-aa27-b0cebb7f2bf6","year":2020},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:f1ad737a9fd8431685a4a5f15fe831ae970752ebd946dc8669969d65ea9e9129","observation_id":"47c5728b-6825-4bd0-a3f9-90a2271183a5","resolution":{"observed_at":"2026-07-09T12:56:15.852189Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.847037Z","title":"In: Guyon, I., Luxburg, U.V ., Bengio, S., Wallach, H., Fergus, R., Vishwanathan, S., Garnett, R","venue":null,"work_id":"2c54873e-5c48-4350-813c-6376d95b6ded","year":2017},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:a1a522a0fc56b87e1366188f577836e5ae456cbd9903a3afc266c3f1219bb637","observation_id":"8f16f9a9-73cf-4456-b7c7-faa797adbb1d","resolution":{"observed_at":"2026-07-09T12:56:15.848463Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.841548Z","title":"In: Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (June 2019)","venue":null,"work_id":"bd96d975-9a97-4fb3-9bde-e407b903797b","year":2019},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:bf7f23becfe4b9fc5e527be20bddfab76d19f919bc8c40a6aaeb11326f8f9599","observation_id":"4cf21fd7-fee6-4c50-b51e-3de69902dc76","resolution":{"observed_at":"2026-07-09T12:56:15.842743Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.839734Z","title":"In: Proc","venue":null,"work_id":"9acd5f26-b65b-4ab4-b9fa-86f4492beb02","year":2021},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:eefdc99654a1a04a69260547b0e24ea5e5694a4e9e33c8252980a23d3f9ec779","observation_id":"4b9f6f3b-d416-4c3c-b1c4-1272c3794278","resolution":{"observed_at":"2026-07-09T12:56:15.840907Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.843381Z","title":"In: Proceedings of the IEEE/CVF International Conference on Computer Vision","venue":null,"work_id":"f84e1b52-7b8e-4613-a024-7aa60744b4a0","year":2023},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:780f485f1186519bc3a7947564019b045c850bd3a5674a4779d810e706bcc273","observation_id":"0c96daad-b752-4c96-be50-94b5b6ca120a","resolution":{"observed_at":"2026-07-09T12:56:15.844445Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.845192Z","title":"In: Proceedings of the IEEE/CVF conference on computer vision and pattern recognition","venue":null,"work_id":"bbff285f-662e-4a7a-b4e8-16e4d31fe8bc","year":2021},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:67177d2f216fdde267ca7d28caf20cf3c844b04ff33eba77fd58829593017765","observation_id":"81a9070c-3436-4612-9e67-46db7805b261","resolution":{"observed_at":"2026-07-09T12:56:15.846376Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.852802Z","title":"In: Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition","venue":null,"work_id":"dc4e720c-ce4b-47cc-b8a2-87c02c846eba","year":2022},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:6f05958ae9a5b572cec51692d55a94c74d7413dbbf1a812fd340e56fad44c66c","observation_id":"60cd62b4-0ed1-4ee3-809e-7413d44307a2","resolution":{"observed_at":"2026-07-09T12:56:15.854057Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2112.08275","last_updated":"2022-07-21T17:28:33Z","snapshot_observed_at":"2026-07-06T12:19:07.751329Z","submitted_at":"2021-12-15T17:09:18Z","title":"SeqFormer: Sequential Transformer for Video Instance Segmentation","version":2},"cited_work":{"arxiv_id":"2112.08275","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2112.08275","snapshot_observed_at":"2026-07-04T03:09:30.485482Z","title":"arXiv preprint arXiv:2112.08275 (2021)","venue":null,"work_id":"f87da0a9-b886-4b63-9db1-b193b77b363c","year":2021},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"cited_paper":"/paper/2112.08275","citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:0b2dab1dcafb489615ca8c5e118796defdcad3d8fada6a9cd97b60cb2ab4bd04","observation_id":"552b16d1-2162-484e-813e-ce3396911b9f","resolution":{"observed_at":"2026-06-30T10:54:37.004244Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.862269Z","title":"In: ECCV (2022)","venue":null,"work_id":"ee2ea169-e118-4d12-ad3f-ae8b80244b99","year":2022},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:28a3569524729e9040c438e47c74fb4e34c4ec75a1aeef1b66d64e6209ba6a0c","observation_id":"3800e894-8e81-485d-894c-1905a5882cf1","resolution":{"observed_at":"2026-07-09T12:56:15.863586Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.835817Z","title":"In: Proceedings of the IEEE/CVF International Conference on Computer Vision (ICCV) (October 2019)","venue":null,"work_id":"6b7a20a1-71df-4c15-a283-4ff9d9399e39","year":2019},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:fd840495b74a28d6d71f9efe27609c0392f1c7fc431725096e4b9cfc0809fbef","observation_id":"ab5fca71-99d4-4d3d-97b5-47ae6a054238","resolution":{"observed_at":"2026-07-09T12:56:15.837022Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.833951Z","title":"In: Proceedings of the IEEE/CVF International Conference on Computer Vision (ICCV)","venue":null,"work_id":"caef8a0b-b0b0-48c3-99d4-4c9c1edf1ef2","year":2021},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:32d5fa296d15495736eda7018574cd0ec489bd3ebe585a1c91dcca779038f3c2","observation_id":"d23c2078-e101-4515-be35-6b028d258337","resolution":{"observed_at":"2026-07-09T12:56:15.835156Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.830097Z","title":"In: Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition","venue":null,"work_id":"5a32f42f-4853-4316-ad20-00a27a1b8d51","year":2022},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:e41b56426ae4eebb1006b7daa168905f5099e420b783442863398236ce188c51","observation_id":"3036d527-8df4-436b-a3e7-5acc6bf43770","resolution":{"observed_at":"2026-07-09T12:56:15.831251Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.831939Z","title":"Advances in Neural Information Processing Systems35, 36324–36336 (2022)","venue":null,"work_id":"ad41af7e-5dd1-4709-8528-08be7db42f2d","year":2022},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:3ef0af3ae57f25e29352089759402b9245ce2147d3008b79446d76741af5642d","observation_id":"74440ff9-30fc-4d22-9c45-cab9a5ea6a10","resolution":{"observed_at":"2026-07-09T12:56:15.833289Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.837764Z","title":null,"venue":null,"work_id":"6af753ee-2784-404c-9deb-335910a8c429","year":2023},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:88dadf5f6c9aa12ee0d60c2a09851c497fbf5d7e5ca84f4d719b923366b9cd3d","observation_id":"2534968c-522d-4bba-aff7-bdbdb5496b37","resolution":{"observed_at":"2026-07-09T12:56:15.839005Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.09108","last_updated":"2022-11-16T18:50:14Z","snapshot_observed_at":"2026-08-09T11:20:44.839581Z","submitted_at":"2022-11-16T18:50:14Z","title":"Robust Online Video Instance Segmentation with Track Queries","version":1},"cited_work":{"arxiv_id":"2211.09108","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2211.09108","snapshot_observed_at":"2026-07-04T03:09:30.488781Z","title":"arXiv preprint arXiv:2211.09108 (2022)","venue":null,"work_id":"6abcb8c3-124c-443d-b986-db01490f6c11","year":2022},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"cited_paper":"/paper/2211.09108","citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:46a9a39ead3cdfc3e7c854ba81472a0894d70f2b87fbe53094a10d757b366f2c","observation_id":"51f39141-2824-4e44-8589-62be77a0d065","resolution":{"observed_at":"2026-06-30T10:54:37.003234Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.826492Z","title":"In: Proceedings of the IEEE/CVF International Conference on Computer Vision (ICCV)","venue":null,"work_id":"c5b0e9c8-cadc-496c-9393-56124e2307f0","year":2023},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:638e6a761c1ca1cb66b407a4f73c4d8be6aa6d78e64c103baa5ba331c4f91631","observation_id":"3f25098c-fc88-469b-b06c-a5e9b7d5d3f7","resolution":{"observed_at":"2026-07-09T12:56:15.827553Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.828251Z","title":"In: MultiMedia Modeling: 26th International Conference, MMM 2020, Daejeon, South Korea, January 5–8, 2020, Proceedings, Part I 26","venue":null,"work_id":"d08b45ac-b4dc-4f62-b952-80777bc54f2c","year":2020},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:55ce93dcff6c2a5c7db5da2a8de53e0c9d0900a9920c5bea1ec455b59757edb6","observation_id":"c79d1e83-9e48-4807-b819-af9c25b34e22","resolution":{"observed_at":"2026-07-09T12:56:15.829470Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T12:56:15.824752Z","title":"In: International Conference on Learning Representations (2021),https://openreview.net/forum?id=gZ9hCDWe6ke 12","venue":null,"work_id":"8f93e1f3-32aa-467b-b749-9bd9765291f6","year":2021},"citing_paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-06-30T10:45:35.568212Z"},"links":{"citing_paper":"/paper/2606.20140"},"observation_digest":"sha256:c42fc958033903392d1599e4c8ecc44845fbdf6ddc95cf1396d51bd1bc285488","observation_id":"e7221ae8-0177-448a-b783-a6ab00c0ad95","resolution":{"observed_at":"2026-07-09T12:56:15.825994Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2606.20140","last_updated":"2026-06-29T11:14:49Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-09T15:03:59.427019Z","submitted_at":"2026-06-18T12:03:04Z","title":"SA-VIS: Sparse frame Annotations for training Video Instance Segmentation"},"reference_resolution":{"displayed":51,"state_counts":{"malformed_identifier":0,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":6,"verified_exact":5,"verified_fuzzy":39},"total_outbound_references":51},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 51 of 51 outbound references and 0 inbound Pith citation observations for arXiv:2606.20140."}