{"as_of":"2026-08-07T07:11:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:f2c1aa9afc7dea2dc4f040744b8ba17ebc261fbf4506dbaca93914d3b808679e","coverage":[{"denominator":40,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":40,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-10T05:45:37.447636Z","state":"measured"},{"denominator":40,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":40,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2607.08537/citation-record","integrity":"/paper/2607.08537/integrity","json":"/paper/2607.08537/citation-record.json","paper":"/paper/2607.08537"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.466151Z","title":"In: European Conference on Computer Vision (ECCV) (2020)","venue":null,"work_id":"049d9bc7-19c8-4ae5-9f3e-c78523adf841","year":2020},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:d92d3d1f83d61986943ad419bfb6ac0d726d66e4ccba502b26d65e0c66a31bf4","observation_id":"e792e6b5-a8c9-44a5-9439-2613c107ecc2","resolution":{"observed_at":"2026-07-10T05:46:50.467227Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.467785Z","title":"Journal of Experimental Algorithmics (JEA)17(2012)","venue":null,"work_id":"06df325e-98f5-46b4-b7a4-ec2b502360f0","year":2012},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:d78a6920db84d1e1e99cb998b1dc45f2df1daf47f38124d4dcdbf40f8cb9d42d","observation_id":"5a84b1cc-51c7-4e6a-a375-9bad2f609df1","resolution":{"observed_at":"2026-07-10T05:46:50.468803Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.461313Z","title":"In: International Conference on Image Analysis and Processing (ICIAP) (2015)","venue":null,"work_id":"5e160667-1e4b-4eaf-93df-5ac178e8dfab","year":2015},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:2cc877a708769cb3979b60296535a453e61899514fffcd7144963a69927015ff","observation_id":"170d79de-a8e3-4f18-8026-d6444a2a843d","resolution":{"observed_at":"2026-07-10T05:46:50.462390Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.463010Z","title":null,"venue":null,"work_id":"eda8b2fb-70b2-4b9d-aa16-30ae06083447","year":2018},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:0f70d25cae1c380f43bc9ddc2caa23773ef9238b1ef11afd75111c64cc3a66a0","observation_id":"06da919a-59e8-4c4b-a287-83e2f7495934","resolution":{"observed_at":"2026-07-10T05:46:50.463989Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.456263Z","title":"In: European Conference on Computer Vision (ECCV) (2024)","venue":null,"work_id":"e1d27179-2b58-46e3-8909-c5971ac9eb46","year":2024},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:75733bbc52ad1035944fae3f3bf14751c9d17c8a4e9c6a0228fe930c1f3f0163","observation_id":"53ea99ef-739a-4a98-a7a4-e57b9c9376d9","resolution":{"observed_at":"2026-07-10T05:46:50.457300Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.457943Z","title":"In: Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2025)","venue":null,"work_id":"95ce23b1-a950-4f0f-a0e5-4d06694245c1","year":2025},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:b77805b0ee84b9aa5d7c4ccdab3ccd0ba4f63084cb9710705d0c0b3f92aeff46","observation_id":"c7cf7521-d987-4a73-9afc-9b2d6a71865b","resolution":{"observed_at":"2026-07-10T05:46:50.459019Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.459549Z","title":"In: Asian Conference on Computer Vision (ACCV) (2024)","venue":null,"work_id":"6e6d704a-ceab-4f37-8635-ccd8dcda663c","year":2024},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:73a723749c04bf5cb73a7bd5a162b1619392f18de0a4b63efb24f52bdc7858d7","observation_id":"a85ce3e5-b725-4a7b-8309-f8d1ba2a31df","resolution":{"observed_at":"2026-07-10T05:46:50.460615Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.464523Z","title":null,"venue":null,"work_id":"e7b8e85f-0821-42ba-b991-2c94a98d1e75","year":2006},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:c507a01cf0f3689236b392ebea89caf398f6defa8114bebddc82c6e95342dee7","observation_id":"da933099-d849-4aeb-93e2-e21c79a14889","resolution":{"observed_at":"2026-07-10T05:46:50.465501Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.469300Z","title":"In: European Conference on Computer Vision (ECCV) (2020)","venue":null,"work_id":"f17f2ddd-8e7e-4db2-a133-ec50c9305d89","year":2020},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:5806f627dc6020f57ba0b8f36cb96f35ba8758ae7b41605974e19a298d262020","observation_id":"f064d9ae-1f59-4a64-92c5-c85fc3617859","resolution":{"observed_at":"2026-07-10T05:46:50.470321Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.449267Z","title":"International Journal of Computer Vision (IJCV)130(2022)","venue":null,"work_id":"43f13f77-1a4e-4249-80a0-2093d6905816","year":2022},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:9f2eddd5b68558f66197add9eb319b520296403c5ecfee63f00f57011296d0fa","observation_id":"a4163001-bac4-46c1-a05b-16b87014f163","resolution":{"observed_at":"2026-07-10T05:46:50.450384Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.447677Z","title":"International Journal of Computer Vision (IJCV)131(2023)","venue":null,"work_id":"65af63c1-aed2-47c1-808b-ed2e108fbfbd","year":2023},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:748aa02e8289b6b499d393f252866a3e06d5b5fb1f601a90001e45a56c808362","observation_id":"fee54797-9af9-4aa2-a6f4-42b370150e28","resolution":{"observed_at":"2026-07-10T05:46:50.448776Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.450841Z","title":null,"venue":null,"work_id":"2edcaa4f-2cf3-4865-9eb6-cad46a7d18d3","year":2025},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:27a411b0842520ec7ac5885857b984f1aa87420d35c0b580a9b46238db24e3e6","observation_id":"91061039-7bd0-462c-9ff1-cb55d1afe5d7","resolution":{"observed_at":"2026-07-10T05:46:50.451870Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.13561","last_updated":"2023-10-01T20:16:22Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-08-24T20:42:21Z","title":"Project Aria: A New Tool for Egocentric Multi-Modal AI Research","version":3},"cited_work":{"arxiv_id":"2308.13561","doi":"10.48550/arxiv.2308.13561","metadata_source":"pith","pith_arxiv_id":"2308.13561","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Project Aria: A New Tool for Egocentric Multi-Modal AI Research","venue":"cs.HC","work_id":"31d9f2a7-f783-4294-9a2c-0a033b2b15fa","year":2023},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"cited_paper":"/paper/2308.13561","citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:a941312f6d4a35fa37ddb2a53958c8ba705f60883cbbcb9222d682cd5e589519","observation_id":"02957dc0-9bfa-4040-b1ff-dc5bc90eea6e","resolution":{"observed_at":"2026-07-10T05:46:50.266450Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.444247Z","title":"In: European Conference on Computer Vision (ECCV) (2024)","venue":null,"work_id":"9a22b83c-c965-4ded-9a55-1f2b16d3048c","year":2024},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:953a5eacd8eca3339e64918f1287a8bf086608828834c4687a3d353a6f207a60","observation_id":"9c309601-61cb-4188-9610-c5a8266cd24a","resolution":{"observed_at":"2026-07-10T05:46:50.445437Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.440569Z","title":"IEEE Transactions on Pattern Analysis and Machine Intelligence (TPAMI)44(2021) Whareformer: Learning to Track What is Where in Long Egocentric Videos 17","venue":null,"work_id":"f7b05800-870d-407f-bd41-d88737c15ef9","year":2021},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:b2c0829061b4c8c38fa5e83fde52a6496759f26486d36f982c8d54f1cc1af92e","observation_id":"55ee5181-9bdb-488a-b0c5-c8442ffe227c","resolution":{"observed_at":"2026-07-10T05:46:50.441760Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.442379Z","title":"Pro- ceedings of the 32nd ACM International Conference on Multimedia (ACM) (2024)","venue":null,"work_id":"037d192c-b1c9-4dd1-b5b0-8aca550bb78a","year":2024},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:6860806d81858bdc5d7a0936821c3a17194d4c80b7e72f83b43aff086869f1d8","observation_id":"60b9ef3c-f2a9-4bde-a137-ac8f3c0ca4ff","resolution":{"observed_at":"2026-07-10T05:46:50.443490Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.445961Z","title":"In: Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2023)","venue":null,"work_id":"f99a6307-386a-4bf3-9329-e8fc66fc5ed1","year":2023},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:01d23cd25081131d7853ae9716b0cfda84c99a7fb97d98d42fbcea2c1315b0e5","observation_id":"fac1f908-56e0-4045-880e-1ced28b527a6","resolution":{"observed_at":"2026-07-10T05:46:50.447163Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.452519Z","title":"In: Proceed- ings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2025)","venue":null,"work_id":"8f94dbb5-9add-4494-88bd-e41f90695fa6","year":2025},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:4cc3923d247559b2f01a576dbe77eb0edb2ca4c6603294e21f9b16256f5c8a1c","observation_id":"92f4b75f-bdf7-40dc-921d-59d26243fd26","resolution":{"observed_at":"2026-07-10T05:46:50.453661Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.433599Z","title":"In: International Conference on Learning Representations (ICLR) (2017)","venue":null,"work_id":"7ef0c7a6-8da4-4eb8-af71-69443d819f21","year":2017},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:65c163bd334044b150108bc1c26562fe747bda7a6ff3bb8018078f9d364d64ca","observation_id":"dd0ce5ce-d813-4a9d-9745-1a31cb3df5e3","resolution":{"observed_at":"2026-07-10T05:46:50.434639Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.435380Z","title":"In: Proceedings of the IEEE/CVF International Conference on Computer Vision (ICCV) (2023)","venue":null,"work_id":"85212697-a68d-49d6-970a-aa8aff2ef95b","year":2023},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:62b4b6c68488ecd54f23749be74a53db7a31ff93c6e4ef05597da7e55f5760ab","observation_id":"ad6c4d08-9bff-4c8f-9364-7dc7b073f955","resolution":{"observed_at":"2026-07-10T05:46:50.436569Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.432059Z","title":"2026 IEEE/CVF Winter Conference on Applications of Computer Vision (WACV) (2024)","venue":null,"work_id":"295d5909-5a57-49ab-aa05-48e98bf31358","year":2026},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:a7d62532c87912ab4666e53c51340f9830c4edd118081123b657a705ec2563d6","observation_id":"e103270c-d2f2-4d75-8400-178b26eaf2c3","resolution":{"observed_at":"2026-07-10T05:46:50.433135Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.437249Z","title":"In: Forty-second International Conference on Machine Learning (ICML) (2025)","venue":null,"work_id":"931ff34b-ad4f-4d95-a6d3-e1d1a9c8705d","year":2025},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:bacbb0ba636190ecb25d9e004b5cac33628b6d8ba467c40fd27e70c74f3d8266","observation_id":"44a1ebb2-b7f6-426c-91a5-5d48fc505a67","resolution":{"observed_at":"2026-07-10T05:46:50.438395Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.428926Z","title":"Journal of The Society for Industrial and Applied Mathematics (SIAM)10(1957)","venue":null,"work_id":"2447d830-2593-45b1-a0a5-dbbb802675c6","year":1957},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:089ee0ea77a990b70a0245be139b73fdc6d1c23a7ffa0023ea4a5fb63f4b4650","observation_id":"33d5191a-0878-4453-9f7c-23b9a596a920","resolution":{"observed_at":"2026-07-10T05:46:50.430116Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.425539Z","title":"Transactions on Ma- chine Learning Research (2024)","venue":null,"work_id":"505a4a86-cbca-433e-aa28-f51f62d27aa8","year":2024},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:0516fbaee3b4511cb7de341d0de754adc07b36014be029c9b6808714c02c490e","observation_id":"157b641e-9bf3-4e85-b44f-5cb69975edf7","resolution":{"observed_at":"2026-07-10T05:46:50.426700Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.427198Z","title":"In: Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (June 2025)","venue":null,"work_id":"c8276511-7bdd-45c2-a687-8bf45acee1f0","year":2025},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:a13a46ef7ea7cb05d473a43a7b8f4b421978db6ef9bc04688c4fa390765d3579","observation_id":"992e03fb-81ef-4534-9d37-1564aeb66d7d","resolution":{"observed_at":"2026-07-10T05:46:50.428427Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.430587Z","title":"In: 2025 International Conference on 3D Vision (3DV) (2025)","venue":null,"work_id":"6641265d-2c43-4348-ada8-a00b30b600ff","year":2025},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:ac95f6514057904dfbd1fcb8e721ee359119b8478393ddd8dd917967ddeaa0ee","observation_id":"2d97431c-d509-4c96-915d-0e27593363f0","resolution":{"observed_at":"2026-07-10T05:46:50.431593Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.438885Z","title":"In: ProceedingsoftheIEEE/CVFConferenceonComputerVisionandPatternRecog- nition (CVPR) (2022)","venue":null,"work_id":"0957077c-9cb3-4e48-a888-96d12183502b","year":2022},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:55fcb2f97f3b5ab17799d985c58b163b53bc54ce468d026bfadafd7aa6927e1a","observation_id":"549d7664-c8f9-44cb-8ab1-aa47371a2ee9","resolution":{"observed_at":"2026-07-10T05:46:50.439958Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.454282Z","title":"In: Proceedings of the IEEE/CVF International Conference on Computer Vision (ICCV) (2023)","venue":null,"work_id":"bbdaffe0-b84c-4c38-8300-9d5b0aa4b2ac","year":2023},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:314970e134c71f67c2de56fed61ef3109e5ffead47992b8fffead9aa44340905","observation_id":"54795cd7-4cdf-4684-847e-7b5b58fc31b6","resolution":{"observed_at":"2026-07-10T05:46:50.455759Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.422250Z","title":"In: Proceedings of the IEEE/CVF Con- ference on Computer Vision and Pattern Recognition (CVPR) (2022)","venue":null,"work_id":"77ff697a-e170-4caa-a075-3ba19b8b2041","year":2022},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:a2ce766d73b6c03051545356abe2364fba36912a86357330529312eead7af487","observation_id":"63e04794-ab4a-4360-9744-f52f60697da5","resolution":{"observed_at":"2026-07-10T05:46:50.423416Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.423906Z","title":"Chalk et al","venue":null,"work_id":"1010ec93-2553-434f-ae9e-86ac351cb4b8","year":2025},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:9d8da65750ef6f2b2e9f06ab1e87b1075cde425cbf07c2c3d26d077e20296f94","observation_id":"bd4a874d-966d-4cce-aefc-be021acd2e2e","resolution":{"observed_at":"2026-07-10T05:46:50.424994Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.420555Z","title":"In: European Conference on Computer Vision (ECCV) (2016)","venue":null,"work_id":"8e2ac09d-268f-40cc-9697-73192537919e","year":2016},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:1b118179a660f1638658035ebd6a60e00ee04720ff0c6fc0c57e6a9763bdae4d","observation_id":"f081ad01-6fe7-4b29-9eab-aa2f2d55bc1e","resolution":{"observed_at":"2026-07-10T05:46:50.421748Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.418471Z","title":"In: Proceedings of the Fourteenth Interna- tional Conference on Artificial Intelligence and Statistics (AISTATS)","venue":null,"work_id":"09f95bfa-acee-4d9d-a95f-25498acbd8a0","year":2011},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:8ae42950537679f504577a1ca9dde05c7aab2bec0f0546d9061ae388f7df105f","observation_id":"399e2653-23bb-42d7-a65d-557cb16d20b4","resolution":{"observed_at":"2026-07-10T05:46:50.419806Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.472445Z","title":"In: European Conference on Computer Vision (ECCV) (2022)","venue":null,"work_id":"65c217f3-997c-4c65-9e6f-acab83861fd3","year":2022},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:7d6507a5c5d6936ef0d6110c0f3af0fd5f8c0f603d1f4db34bff6c5f11d5c47d","observation_id":"4f120cae-8fbd-4d73-bdda-901fb53ae3a2","resolution":{"observed_at":"2026-07-10T05:46:50.473461Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.476442Z","title":"Advances in Neural Information Processing Systems (NeurIPS)36(2023)","venue":null,"work_id":"3efc25d9-3e12-4646-be55-fc0541b676a8","year":2023},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:1562aa4842fefa66d27ab20d74e932d9f9d34683d3b58e93e64d873e98e99d21","observation_id":"c5f3f0aa-d702-462f-8b19-e2360f8a095a","resolution":{"observed_at":"2026-07-10T05:46:50.477903Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.474435Z","title":"IEEE Robotics and Automation Letters (RA-L)7(2022)","venue":null,"work_id":"722f2ab7-e6ec-492b-a432-dff01cd32889","year":2022},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:7894ce98d4299841bb93f137ed57f8f7f9eb25e7a06685b592837e77099a7a54","observation_id":"15a43750-4c7d-4c3e-ba9b-22974e7e3643","resolution":{"observed_at":"2026-07-10T05:46:50.475673Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.481753Z","title":"International Conference on Intelligent Robots and Systems (IROS) (2020)","venue":null,"work_id":"a403d01d-9e52-4adb-aed6-df84fc440ff3","year":2020},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:8fb9fdf98d908a6cbb22aad760b66e32706eb21eae0f9a9c2adc2fd79037a71d","observation_id":"32f39a8d-9adb-481e-bbaf-75dada865ec2","resolution":{"observed_at":"2026-07-10T05:46:50.482834Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.483324Z","title":"IEEE Transactions on Pattern Analysis and Machine Intelligence (TPAMI) 35(2012)","venue":null,"work_id":"e5e54525-2c58-44e0-80ef-e432f51a22c7","year":2012},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:3793312802b274bc383febfc167779f08ddaef556d0d2798e2438c82c32c0b77","observation_id":"efd84cee-81e8-463f-88a4-a60de2e58bdd","resolution":{"observed_at":"2026-07-10T05:46:50.484469Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.478464Z","title":"In: Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2021)","venue":null,"work_id":"20405b75-6730-4551-84d5-9f12d0b52155","year":2021},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:dc76772ef83995e646b89dc8af65bf5f5acf83fc42a8aaacab5cccaf81d0950c","observation_id":"46e710c6-edda-4717-9be7-77e6550d8a44","resolution":{"observed_at":"2026-07-10T05:46:50.479534Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.480006Z","title":"In: European Conference on Computer Vision (ECCV) (2022)","venue":null,"work_id":"dbb228d1-9233-4a23-9908-cad5a21cd4c4","year":2022},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:9d3585f04d9727af3c30d821a7c1ade0c8fbc4f87776604e59c463e80523be1e","observation_id":"05931341-1b15-4587-a2c7-1e4a908554ab","resolution":{"observed_at":"2026-07-10T05:46:50.481046Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.470932Z","title":"what” and “where","venue":null,"work_id":"e9bbad02-e06d-43ee-b24f-ad47d69ba32e","year":2024},"citing_paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-07-10T05:45:37.447636Z"},"links":{"citing_paper":"/paper/2607.08537"},"observation_digest":"sha256:8bd3e28e4158977ebbb8da7ed8cef115fbc6817847996b42e7e11612bfffb79e","observation_id":"57517edd-83b8-4092-8e6d-8890180a926c","resolution":{"observed_at":"2026-07-10T05:46:50.471943Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2607.08537","last_updated":"2026-07-09T14:33:24Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-02T01:48:49.204376Z","submitted_at":"2026-07-09T14:33:24Z","title":"Whareformer: Learning to Track What is Where in Long Egocentric Videos"},"reference_resolution":{"displayed":40,"state_counts":{"malformed_identifier":1,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":3,"verified_exact":0,"verified_fuzzy":35},"total_outbound_references":40},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 40 of 40 outbound references and 0 inbound Pith citation observations for arXiv:2607.08537."}