{"as_of":"2026-08-02T13:29:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:d88d521f0b8f3c5f5517601386bdd16c4c95a9107c6ae5b3b175ad276fa80d4d","coverage":[{"denominator":48,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":48,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-16T07:16:29.588452Z","state":"measured"},{"denominator":49,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":49,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-02T06:30:47.504484+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-01T11:30:57.042487Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2602.05638","snapshot_observed_at":"2026-08-01T11:30:57.042487Z","title":"& Others UniSurg: A Video-Native Foundation Model for Universal Understanding of Surgical Videos.ArXiv Preprint ArXiv:2602.05638","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.19889","last_updated":"2026-07-22T08:20:52Z","snapshot_observed_at":"2026-08-01T11:30:55.143902Z","submitted_at":"2026-07-22T08:20:52Z","title":"LAVIFT: Latent-Action-Guided Vision Fine-Tuning for Surgical Interaction Recognition","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-01T11:30:57.042487Z"},"links":{"cited_paper":"/paper/2602.05638","citing_paper":"/paper/2607.19889"},"observation_digest":"sha256:600f90054ba04b6eebdbe12d851dc407c09fefb23ef6b35e637ac0b9bd7f6aa7","observation_id":"0ba885f5-863f-45d2-9ca0-d3b792d444ab","resolution":{"observed_at":"2026-08-01T11:30:57.042487Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2602.05638/citation-record","integrity":"/paper/2602.05638/integrity","json":"/paper/2602.05638/citation-record.json","paper":"/paper/2602.05638"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2304.07193","last_updated":"2024-02-02T10:24:09Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-04-14T15:12:19Z","title":"DINOv2: Learning Robust Visual Features without Supervision","version":2},"cited_work":{"arxiv_id":"2304.07193","doi":"10.48550/arxiv.2304.07193","metadata_source":"pith","pith_arxiv_id":"2304.07193","snapshot_observed_at":"2026-07-11T00:07:42.299741Z","title":"DINOv2: Learning Robust Visual Features without Supervision","venue":"cs.CV","work_id":"26b304e5-b54a-4f26-be7e-83299eca52e4","year":2023},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"cited_paper":"/paper/2304.07193","citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:8e81386eddf76b1cf7fb181d40cc7da2f598669d221b30f9f033df3df92bdc59","observation_id":"65c8424e-6828-4ea6-8cf6-2bab48d67022","resolution":{"observed_at":"2026-05-16T07:17:30.329076Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-06T18:32:48.038151Z","title":"Masked autoencoders are scalable vision learners","venue":null,"work_id":"23a19f52-833f-4385-a893-cba047433888","year":2022},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:8b1d76ceccf40b9cd08edf215c28d095c0ce927334a5404b4f168d4aa84361ce","observation_id":"c7eaac3f-bcd2-4ab7-9889-f3bd8f9f0e80","resolution":{"observed_at":"2026-05-16T07:17:31.042736Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Masked autoencoders as spatiotemporal learners","venue":null,"work_id":"865e299e-505f-4a01-893f-33cacf47e6e3","year":2022},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:b4e6131d77a32efe2a7ea8a29c6ce1719b975b8470e6b38c0edd534c99b0fa45","observation_id":"5a94e565-6873-4479-8a3f-d1166deb91bb","resolution":{"observed_at":"2026-05-16T07:17:31.049984Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Endovit: pretraining vision transformers on a large collection of endoscopic images","venue":null,"work_id":"4d636846-6483-451c-875f-5c26818de446","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:e1a783bb9815f973c0fedd1b7a376b551628f9f52a659853640f5ed5275cf466","observation_id":"a8c9586c-5198-4534-9ffa-53d9cfbe7d17","resolution":{"observed_at":"2026-05-16T07:17:31.045207Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Foundation model for endoscopy video analysis via large- scale self-supervised pre-train","venue":null,"work_id":"a6f381e7-048a-47af-9a80-d5db99acd397","year":2023},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:9bfb26b1644524ca4118a91de32361fd8333b46c8e46292f2b2c17214bcd0555","observation_id":"f0317426-8f8a-4b0f-8a0b-9561d3a8b46c","resolution":{"observed_at":"2026-05-16T07:17:31.052145Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.05949","last_updated":"2024-04-12T22:30:54Z","snapshot_observed_at":"2026-07-06T17:42:03.971883Z","submitted_at":"2024-03-09T16:02:46Z","title":"General surgery vision transformer: A video pre-trained foundation model for general surgery","version":3},"cited_work":{"arxiv_id":"2403.05949","doi":"10.48550/arxiv.2403.05949","metadata_source":"arxiv_reference","pith_arxiv_id":"2403.05949","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Schmidgall, J","venue":null,"work_id":"e879f7aa-b06c-432c-949c-e12cf1d790ee","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"cited_paper":"/paper/2403.05949","citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:932a1de503cc305f5aec2a6af93e42db527a296f1f086a729c9c5114f09153dd","observation_id":"bf41add1-0ecf-4883-a5fa-aa24103cf21a","resolution":{"observed_at":"2026-05-16T07:17:30.324703Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Videomae: Masked autoencoders are data-efficient learners for self-supervised video pre-training","venue":null,"work_id":"e9210d53-bd15-4edf-8dbe-3aab9d819a81","year":2022},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:7ecc9991ba9643bb91b24c73091f85470fa5c3a9d4c29960b22587a403bfa4cc","observation_id":"87a0759a-26b3-4e95-bce9-2fc11c56f802","resolution":{"observed_at":"2026-05-16T07:17:31.054747Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Videomae v2: Scaling video masked autoencoders with dual masking","venue":null,"work_id":"21d4295f-c66e-4abf-9660-50a659f6e94f","year":2023},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:eb5179e53f8618747857b2079f699e4926f3df9780a016bd0ab4ec9de03cae7d","observation_id":"d34de7fd-0b84-4b29-b33b-c8fcc05c2a96","resolution":{"observed_at":"2026-05-16T07:17:31.056970Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Dissecting self-supervised learning methods for surgical computer vision","venue":null,"work_id":"a124c126-f3b3-416c-b4e6-745c77523762","year":2023},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:130b255060b1fd1dca00492ca1ce79653c7b7014042c4b37fd44fa42f9d0f38e","observation_id":"098ed1b4-6051-48cc-9cfb-d45625b74eef","resolution":{"observed_at":"2026-05-16T07:17:31.047502Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Endonet: a deep architecture for recognition tasks on laparoscopic videos","venue":null,"work_id":"d77e719b-f339-4187-96b3-0407d900dc23","year":2016},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:777e49036c37cc2a0e02d23be9527d8ada12ba575896c7d66c512b5634afdde3","observation_id":"71913d2d-5ac9-4783-a5dc-4b4654ef6414","resolution":{"observed_at":"2026-05-16T07:17:31.059469Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Pitvis-2023 challenge: Workflow recognition in videos of endoscopic pituitary surgery","venue":null,"work_id":"a4eed530-a7a2-48ed-b460-6b07bcf62c05","year":2023},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:b8fa3731a97632b6c36f4ea492e342a7fce13aa744ee8fda06a45434d57f73fc","observation_id":"100195a1-baa4-47bc-a1c2-67a51dad23bf","resolution":{"observed_at":"2026-05-16T07:17:31.016048Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Egosurgery-phase: a dataset of surgical phase recognition from egocentric open surgery videos","venue":null,"work_id":"9c176e39-5383-40ac-aded-11d59f7e4edf","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:8c374730dc8bd41bd1be899aff5f594195c263ff3264c0e88034863b0f14d25d","observation_id":"b2615ac1-326c-4387-9d6a-beaf33b18a48","resolution":{"observed_at":"2026-05-16T07:17:31.026298Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.08471","last_updated":"2024-02-15T18:59:11Z","snapshot_observed_at":"2026-08-02T05:40:31.086014Z","submitted_at":"2024-02-15T18:59:11Z","title":"Revisiting Feature Prediction for Learning Visual Representations from Video","version":1},"cited_work":{"arxiv_id":"2404.08471","doi":"10.1145/3178876.3185996","metadata_source":"pith","pith_arxiv_id":"2404.08471","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Revisiting Feature Prediction for Learning Visual Representations from Video","venue":"cs.CV","work_id":"f7251dcf-5341-4915-bfe7-27812387b61a","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"cited_paper":"/paper/2404.08471","citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:90a30c5079da97d5938b4f14a185eefb3fb30a47c9aef451e75f2f7846872413","observation_id":"b991f193-b445-45b3-8239-99d6d20e37d5","resolution":{"observed_at":"2026-05-16T07:17:30.319773Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-05-23T01:54:34.489718+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T01:54:34.489718+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.09985","last_updated":"2025-06-11T17:57:09Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-06-11T17:57:09Z","title":"V-JEPA 2: Self-Supervised Video Models Enable Understanding, Prediction and Planning","version":1},"cited_work":{"arxiv_id":"2506.09985","doi":"10.48550/arxiv.2506.09985","metadata_source":"pith","pith_arxiv_id":"2506.09985","snapshot_observed_at":"2026-07-11T02:27:49.493432Z","title":"V-JEPA 2: Self-Supervised Video Models Enable Understanding, Prediction and Planning","venue":"cs.AI","work_id":"a9c28401-f16a-4933-89f0-788e2f94e52b","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"cited_paper":"/paper/2506.09985","citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:1d6d661d7b6fb3cc3985e375f419a6468681a0b712f61ab150a3e720398d407c","observation_id":"b5c27669-8051-4f81-9fca-a33b4e8e9d6e","resolution":{"observed_at":"2026-05-16T07:17:30.309716Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Bootstrap your own latent: A new approach to self-supervised learn- ing","venue":null,"work_id":"fa22a7fc-4e5f-4eef-a115-8d3356ff9f10","year":2020},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:bcf96dd918ed46f15abd732ac6ceac4cfdc9725ecf17cfa5a9b955d4d68ec442","observation_id":"1f629bb9-cf2f-445b-a19a-b81cd415a272","resolution":{"observed_at":"2026-05-16T07:17:31.021009Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Internvideo2: Scaling video foundation models for multimodal video understanding","venue":null,"work_id":"5cf2349d-b2fe-4140-a4e1-2a382a56f596","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:7d2b4a640070e7c943dff31259bdd9d7b9cd2d148020b37820d6638c84a44c97","observation_id":"4e21a4e5-bfbb-4062-9c7e-20bf4afa6a3a","resolution":{"observed_at":"2026-05-16T07:17:31.018463Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2512.01342","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-03T14:28:32.049805Z","title":"Internvideo-next: Towards general video foundation models without video-text supervision","venue":null,"work_id":"d65f61ad-e3f7-4e3f-9b08-da7a6e8b0595","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:c976a58ce02dc3e745ffa7ed17cbfcd418d323dbaeb9b8096f424a499423df33","observation_id":"7214fd01-073e-495c-ba87-88203e7ec15b","resolution":{"observed_at":"2026-05-16T07:17:30.296414Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T23:44:22.359690Z","title":"Emerging properties in self-supervised vision transformers","venue":null,"work_id":"c47546d4-4e91-494a-a5e3-a1e01ff6556f","year":2021},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:3a00799eaa239e87eb4a55e48a1f8c732636770fade25257089423465b3e23fa","observation_id":"8d73de36-f6d3-4ab2-bc36-7e6bf6778573","resolution":{"observed_at":"2026-05-16T07:17:31.023621Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.10104","last_updated":"2025-08-13T18:00:55Z","snapshot_observed_at":"2026-07-06T22:12:35.584339Z","submitted_at":"2025-08-13T18:00:55Z","title":"DINOv3","version":1},"cited_work":{"arxiv_id":"2508.10104","doi":"10.1055/a-2487-1252","metadata_source":"pith","pith_arxiv_id":"2508.10104","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"DINOv3","venue":"cs.CV","work_id":"c8b07deb-8fe7-4e18-9620-f3569d3529ce","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"cited_paper":"/paper/2508.10104","citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:10a01d84963ef2f119fde44277fb40985bfe3c4929b5baf2e961583efc924cb8","observation_id":"1021f588-4279-4938-b861-5e2df567ecaf","resolution":{"observed_at":"2026-05-16T07:17:30.314610Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Gastronet-5m: A multicenter dataset for developing foundation models in gastrointestinal endoscopy","venue":null,"work_id":"5b482dd4-c8fb-4307-8e07-c3f44db1548c","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:66b0c3413292af5007e82ab825b4bfefa4c05abac90013f73f1a41692d3e26ac","observation_id":"4e6d5886-4dd6-4273-a38c-99dd91c43f7c","resolution":{"observed_at":"2026-05-16T07:17:31.028649Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Self-supervised learning for endoscopic video analysis","venue":null,"work_id":"70739a11-dcec-434d-b6cc-b8cbefba7c3e","year":2023},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:6fe405af8ea04cc633de0c984d916fbb77a2e108e9a7c700272ca5a562402fd2","observation_id":"1c07827e-2596-4bfa-be7d-42d423e629cc","resolution":{"observed_at":"2026-05-16T07:17:31.011155Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Endomamba: an efficient founda- tion model for endoscopic videos via hierarchical pre-training","venue":null,"work_id":"bfe27351-6c16-4dbd-87f9-0bda669fd6ba","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:18380245a41af7c618c06e7246cc2b67f6206a4e342b8aa0f9a8e7d08a962f45","observation_id":"ebc9fd4d-172e-470d-8857-de06490ac087","resolution":{"observed_at":"2026-05-16T07:17:31.005691Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2501.09436","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Scaling up self-supervised learning for improved surgical foundation models","venue":null,"work_id":"4271cc8e-489b-44fa-8dfd-c0ae7720647e","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:14ce3cc43c7246d403ba5f116b979b1ff4e4940bcb957e358f32c07f012a16d6","observation_id":"6fd0697f-7fca-40d5-a3fa-587ab0e66800","resolution":{"observed_at":"2026-05-16T07:17:30.292026Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Learn- ing multi-modal representations by watching hundreds of surgical video lectures","venue":null,"work_id":"c2417f7b-a452-4aa6-8abb-f9b5fecdee35","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:c7d8d797b78eccd84ff11903fbb97fc18c366371735d5d7d538d51ae972a3b45","observation_id":"a0969ef8-c80e-4c83-8510-7030bec8b518","resolution":{"observed_at":"2026-05-16T07:17:31.002969Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1610.09278","last_updated":"2017-08-31T14:27:37Z","snapshot_observed_at":"2026-07-06T05:16:33.346094Z","submitted_at":"2016-10-28T15:36:58Z","title":"The TUM LapChole dataset for the M2CAI 2016 workflow challenge","version":2},"cited_work":{"arxiv_id":"1610.09278","doi":null,"metadata_source":"pith","pith_arxiv_id":"1610.09278","snapshot_observed_at":"2026-07-04T20:10:07.403035Z","title":"The TUM LapChole dataset for the M2CAI 2016 workflow challenge","venue":"cs.CV","work_id":"020971ab-45bc-489b-9897-f0255dc6d71f","year":2016},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"cited_paper":"/paper/1610.09278","citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:8197f89326da27ba46b1fd59b42809bcd5731c848b324a95f8ea63bb3e008984","observation_id":"ae41b1ce-250c-4a0b-86ad-af55f87cc6f9","resolution":{"observed_at":"2026-05-16T07:17:30.299900Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Rendezvous: Attention mechanisms for the recognition of surgical action triplets in endoscopic videos","venue":null,"work_id":"117d7265-b07e-4307-a8e9-a450e4519ef0","year":2022},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:219fe53bc976bd40c1eec27d1b7906dafc99e44cb219f89d3ca5cabb0d38e175","observation_id":"881fa467-b602-416e-b807-a77cf03bdad8","resolution":{"observed_at":"2026-05-16T07:17:31.000856Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Autolaparo: A new dataset of integrated multi-tasks for image-guided surgical automation in laparoscopic hysterectomy","venue":null,"work_id":"f9426651-8798-47fd-8b49-0e36d6f648dd","year":2022},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:71d897aa62b17a27e3301c1d056325ab4085a49a338553b389b1581fc255796a","observation_id":"7ee17f12-9e50-4697-b874-622379a6e908","resolution":{"observed_at":"2026-05-16T07:17:30.998477Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Surgical workflow recognition and blocking effectiveness detection in laparoscopic liver resection with pringle maneuver","venue":null,"work_id":"0048b71b-f79f-4709-aa9a-0ac2c7ef5012","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:115ebc8412ad0a2cf07453ed2df9ce7bccbda8e5f9f143ef4caf59970b1e5995","observation_id":"91c4640e-79da-496f-a89b-21038eaea380","resolution":{"observed_at":"2026-05-16T07:17:31.031094Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Ophnet: A large-scale video benchmark for ophthalmic surgical workflow understanding","venue":null,"work_id":"a6fb6497-8c30-4650-b48c-5d17e575ae25","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:3f4f7ae090a685d1f2b08b4572e21410b8e2bffe34ee150c0c3493e56f0904f8","observation_id":"bd78fb99-e91a-4db8-b9b3-a2cb15f72130","resolution":{"observed_at":"2026-05-16T07:17:31.013524Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Analyzing surgical technique in diverse open surgical videos with multitask machine learning","venue":null,"work_id":"b5fb16e1-78bd-470f-96d2-9dd6a7e6a033","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:aad158b8f49fa6ee52603127348b8d9b78fb4c3528492cb22e3e5d1204708ee2","observation_id":"f4bda7a8-62b1-4c61-89a2-f4f9cae6c375","resolution":{"observed_at":"2026-05-16T07:17:30.985662Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"A dataset and benchmarks for segmentation and recognition of gestures in robotic surgery","venue":null,"work_id":"5cf765af-4ac2-4125-8fe3-0181e9c9092b","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:ebab22400a87ed0c33c7165f5316b9e06ce41e11e4b46b75aa7839542ee9b4d9","observation_id":"bcaee95c-e5e9-48f3-97a5-39cdcec28c13","resolution":{"observed_at":"2026-05-16T07:17:31.038165Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Aixsuture: vision-based assessment of open suturing skills","venue":null,"work_id":"28a4f930-c4b9-49db-94ce-274bc2709727","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:5b89a967eec57bb70c5f694d2e4428baf14c574d7c5dfd77ce4b3cd263ef3947","observation_id":"416044c4-e24f-4216-8ad4-2198076a5ebd","resolution":{"observed_at":"2026-05-16T07:17:31.033391Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Video retrieval in laparoscopic video recordings with dynamic content descriptors","venue":null,"work_id":"3b024d39-879a-492e-bb91-feb1f3e58b77","year":2018},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:2d73da1f1292fcb6ea1d826054c0ebf3f6bb91b09e8789f4055cc88ada846aa3","observation_id":"7f67398d-310d-4c54-844f-8d8438a43f0a","resolution":{"observed_at":"2026-05-16T07:17:30.983458Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Contrastive transformer- based multiple instance learning for weakly supervised polyp frame detection","venue":null,"work_id":"3720a984-cc0e-4076-98df-dbb137e4825c","year":2022},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:c1270e124fb0d2ccfaafeacc1981b3e907b932c8a012b3832fd367fce00f494b","observation_id":"0e5c71cd-c19b-4398-98b6-8ffae8dba498","resolution":{"observed_at":"2026-05-16T07:17:31.035865Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Wm-dova maps for accurate polyp highlighting in colonoscopy: Validation vs. saliency maps from physicians","venue":null,"work_id":"8432b16d-9f06-464d-829d-921f9ccea4a6","year":2015},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:8ea3561c646fd54cce074902363dbd22694daf11c532b89072bcadc9fc7c9def","observation_id":"f98b61b0-9374-4ff1-a737-679468ca8c54","resolution":{"observed_at":"2026-05-16T07:17:30.989951Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Implicit domain adaptation with conditional generative adversarial networks for depth prediction in endoscopy","venue":null,"work_id":"2b36d126-2df8-4676-8fa8-a55a58f1334c","year":2019},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:0c9376dd5ad8d9e886bf7a9b242393eacaa1bed48041641a90c4d6750e464357","observation_id":"2cc3a194-d676-4fae-88de-d27188eeae67","resolution":{"observed_at":"2026-05-16T07:17:31.040582Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Colonoscopy 3d video dataset with paired depth from 2d-3d registration","venue":null,"work_id":"f5bcba8c-59c3-47ce-ae77-f34ec63c7d76","year":2023},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:ee2496b37e87d61b8dcbfed9e0592f7ac36a0deb56bf3114b645ebcc1d8dfd6e","observation_id":"7df0091d-7d66-432f-a3d1-bdbd0e6f5b4f","resolution":{"observed_at":"2026-05-16T07:17:31.008591Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.21227/ac97-8m18","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Cataracts: Challenge on automatic tool annotation for cataract surgery","venue":null,"work_id":"0f81b61c-8a7b-4a7b-a2f2-9da6d7c7291d","year":2019},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:8bc58b285b9d735813a66e0a5d8af2b08ac7a62222938dc1ec557033b8282c36","observation_id":"7a94bfda-bcd3-46bc-a550-8a6d427f0949","resolution":{"observed_at":"2026-05-16T07:17:30.179924Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Challenges in multi-centric generalization: phase and step recog- nition in roux-en-y gastric bypass surgery","venue":null,"work_id":"5e4ca10e-ec87-4245-aaa0-72e344cb63d1","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:8baf3e4f92692e54942c7dc4c4332a1afe75829020914b6caf9e9dffc9f2cad5","observation_id":"eef2e08a-40b3-4bf1-b948-f20326cd3f64","resolution":{"observed_at":"2026-05-16T07:17:30.976472Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Copesd: A multi- level surgical motion dataset for training large vision-language models to co-pilot endoscopic submucosal dissection","venue":null,"work_id":"8d7ea69c-56c6-4396-8019-6fc40289e1d9","year":2024},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:669e2a03c69af8c9d376f1467185500354a09d7e6f162081067dde2a5b51d0a7","observation_id":"2d9333d2-f622-4027-8bdd-c18428c6e1b9","resolution":{"observed_at":"2026-05-16T07:17:30.978987Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Towards holistic surgical scene understanding","venue":null,"work_id":"76c9e713-3f89-4ecd-ba64-4bcbc18ecbf9","year":2023},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:c5488a93e0e0d869a50a5acf3f98c39c7d991febf5bb232b8c70ab9eedb8cb86","observation_id":"fe24adf8-d2a4-435a-b1e2-1636d1db2711","resolution":{"observed_at":"2026-05-16T07:17:30.974178Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Kvasir-seg: A segmented polyp dataset","venue":null,"work_id":"b48e79c0-6a76-47e0-9a9a-400f1b927c27","year":2020},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:d083f4334125c0248d0f17ea7161283c476a18c77396d83a750d8219720269e5","observation_id":"00a9411f-b337-447d-82ab-0ab0a47cb83b","resolution":{"observed_at":"2026-05-16T07:17:30.971963Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"A benchmark for endoluminal scene segmentation of colonoscopy images","venue":null,"work_id":"b15940e6-b237-4970-ad13-87c03c4f35bf","year":2017},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:77af917ea238cd8dd002b32761a6572408cafb2b160cb0922bae16d7cfc8b12c","observation_id":"10bfce8b-2270-4931-b1ae-c7ab493c2d0e","resolution":{"observed_at":"2026-05-16T07:17:30.981215Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Towards automatic polyp detection with a polyp appearance model","venue":null,"work_id":"2b1717a8-acd8-4c3d-ad9c-e1782d9a8db5","year":2012},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:f288a8e9583c4e200804243b56044975c990c05d35d30be36bb9a63811855f6b","observation_id":"0d5a1ad6-0bd0-4022-9370-4de51db83397","resolution":{"observed_at":"2026-05-16T07:17:30.987823Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Toward embedded detection of polyps in wce images for early diagnosis of colorectal cancer","venue":null,"work_id":"5334d03e-687a-4d7e-9e11-36e593cb3a94","year":2014},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:b7dbae9533212c065297b3f7881ce955232ee2099afb68edb59275f82d98985d","observation_id":"59a4ce96-7c01-40eb-8b5b-5b2c83f43849","resolution":{"observed_at":"2026-05-16T07:17:30.992302Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Pranet: Parallel reverse attention network for polyp segmentation","venue":null,"work_id":"ed1aa94e-7244-4e6f-8ca2-1d780406bdf4","year":2020},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:45ba1c9e2cdd9b1f0ce101969af542af00a3a1a99e92535d206881f6223bbbf0","observation_id":"91b31e61-a7de-4ceb-8c61-a2b21be828ba","resolution":{"observed_at":"2026-05-16T07:17:30.969766Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Uacanet: Uncertainty augmented context attention for polyp segmentation","venue":null,"work_id":"3fbf992d-3eff-4ef6-b6be-7dc37bfb8ddb","year":2021},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:8b87815375f3037ccaa187c5ffd6ac740bf7275006ea62cee14b88996315080b","observation_id":"4d2cb680-6211-4d5a-b3bc-b180c933ea31","resolution":{"observed_at":"2026-05-16T07:17:30.967308Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2504.10986","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Pranet-v2: Dual-supervised reverse attention for medical image segmentation","venue":null,"work_id":"8d8b11c2-95c7-4fb3-a135-4a75a0d83f52","year":2025},"citing_paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos","version":3},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-05-16T07:16:29.588452Z"},"links":{"citing_paper":"/paper/2602.05638"},"observation_digest":"sha256:e1f4a077f6ab1c1bb5532ce159a78ef3d22c1bcdae55606a3e5e0bb03f6d4657","observation_id":"243a59d5-e604-4e35-950e-9cf958af9d13","resolution":{"observed_at":"2026-05-16T07:17:30.304570Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2602.05638","last_updated":"2026-04-17T08:31:14Z","latest_version":3,"primary_category":"cs.CV","snapshot_observed_at":"2026-07-06T22:44:42.380489Z","submitted_at":"2026-02-05T13:18:33Z","title":"SurgMotion: A Video-Native Foundation Model for Universal Understanding of Surgical Videos"},"reference_resolution":{"displayed":48,"state_counts":{"malformed_identifier":0,"metadata_mismatch":2,"parse_uncertain":0,"unresolved":0,"verified_exact":8,"verified_fuzzy":38},"total_outbound_references":48},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"thesis":"As of 2 August 2026, this Paper Citation Record lists 48 of 48 outbound references and 1 inbound Pith citation observation for arXiv:2602.05638."}