{"as_of":"2026-08-09T07:20:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:1742447017f574ce58d99ffb8f39d0b7117e1f76fb2c42b35d00c553d351f234","coverage":[{"denominator":59,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":59,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T14:02:05.872722Z","state":"measured"},{"denominator":59,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":59,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2508.21770/citation-record","integrity":"/paper/2508.21770/integrity","json":"/paper/2508.21770/citation-record.json","paper":"/paper/2508.21770"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:14.590184Z","title":"Ubnor- mal: New benchmark for supervised open-set video anomaly detection","venue":null,"work_id":"319068f0-ea10-4134-987f-90da64c29a19","year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-05T14:01:59.564950Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:637b150726843e226734f0366f29f8b6a920656af52cd00459bb81a51a3443c5","observation_id":"c494610a-58b0-4d31-a7b6-0c80420f89ef","resolution":{"observed_at":"2026-08-05T14:02:14.607104Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:14.414086Z","title":"Towards open set deep networks","venue":null,"work_id":"0cdba8d6-0a33-4177-aa76-a50b717ed8b4","year":2016},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-05T14:01:59.741279Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:86489b49f878253f60c7cda4e5a4d98c4c1ee0d71ee7de19df012f7e758f1186","observation_id":"4fccf45d-4b79-4a40-95d8-dd36319a87e3","resolution":{"observed_at":"2026-08-05T14:02:14.498084Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:14.289633Z","title":"Is space-time attention all you need for video understanding? InProceedings of the International Conference on Machine Learning (ICML), July 2021","venue":null,"work_id":"e6029011-6371-4cca-962d-72b4327ee56c","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-05T14:01:59.975483Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:e8e1dabc373b9fbf4de42049601cd382118682f45eecdd16975fb009e0a7d604","observation_id":"cc1f2ea9-6283-4560-9b31-dfbea8634f0f","resolution":{"observed_at":"2026-08-05T14:02:14.361736Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:14.139312Z","title":"Enlarging instance-specific and class-specific information for open-set action recognition","venue":null,"work_id":"dd9e15a3-e759-49e7-a5f2-266cb6114254","year":2023},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:00.109712Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:a06c450afadab90cdd16799677a137fa59dbf82d0afefd18dfc1b2ac031822b1","observation_id":"500f5cfb-9b58-42ce-a045-cd8a50f365ec","resolution":{"observed_at":"2026-08-05T14:02:14.188134Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:00.192110Z","title":"Elaborative rehearsal for zero-shot action recognition","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:00.192110Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:d60dd2776f719abc7a3fe9b32c04a8f998428880ef8749618b9200aa84c3a2dd","observation_id":"3a761539-6bbd-443a-baae-8ed237cb4a19","resolution":{"observed_at":"2026-08-05T14:02:00.192110Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:13.993825Z","title":"Wdiscood: Out-of- distribution detection via whitened linear discriminant analysis","venue":null,"work_id":"a57e6282-bf9c-448e-a099-ad91e591b64d","year":2023},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:00.348788Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:cb0a94841844e4a97a3335a65c3d132d4a8e9bde06a3c4fbf1f98d98028fd005","observation_id":"ce7c0c73-129b-4b78-8fb7-fb430e31df85","resolution":{"observed_at":"2026-08-05T14:02:14.024524Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:13.836386Z","title":"Haa500: Human-centric atomic action dataset with curated videos","venue":null,"work_id":"d3e1b2e0-02b3-44f4-88ec-1610cf2ab0d9","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:00.439028Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:09837da50e3be239c758f8aaa50f332df0799890d1c9b073160022a7a9386923","observation_id":"1c77dc5e-7b2b-4290-8ced-d051ba3fbbdc","resolution":{"observed_at":"2026-08-05T14:02:13.940558Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:13.481648Z","title":"Towards unknown-aware learning with virtual outlier synthesis","venue":null,"work_id":"f4cfd443-06fc-496e-940f-599e60c4217d","year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:00.578152Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:c63f317cfea5ceb6911584b3a79536dd06e7772be6679f37322961d1f804547d","observation_id":"5a4adc85-ee08-4796-916e-c98c8e7905db","resolution":{"observed_at":"2026-08-05T14:02:13.644210Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:13.233276Z","title":"Oops! predicting unintentional ac- tion in video","venue":null,"work_id":"db09e466-2130-42e9-b431-943a5af5efa9","year":2020},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:00.696918Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:168089eb9603b826d0f78b72a815c5c3cae86792f158b9560f899bb6cf1a1493","observation_id":"b5dffb49-5c8c-4155-b694-976d266a0012","resolution":{"observed_at":"2026-08-05T14:02:13.343116Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11736","last_updated":"2023-03-28T18:27:24Z","snapshot_observed_at":"2026-08-07T10:53:37.587234Z","submitted_at":"2022-06-23T14:31:33Z","title":"NovelCraft: A Dataset for Novelty Detection and Discovery in Open Worlds","version":3},"cited_work":{"arxiv_id":"2206.11736","doi":null,"metadata_source":"pith","pith_arxiv_id":"2206.11736","snapshot_observed_at":"2026-08-05T14:02:06.394120Z","title":"NovelCraft: A Dataset for Novelty Detection and Discovery in Open Worlds","venue":"cs.CV","work_id":"5fa8a380-af3c-4648-9983-3369a9ee0ce0","year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:00.857266Z"},"links":{"cited_paper":"/paper/2206.11736","citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:f4c860976aaedbd6e4b0cffd93205a8afdcdb40e70eae4d21a3d710e6633cc39","observation_id":"4b29c84f-1298-4166-b031-c841b3872aa5","resolution":{"observed_at":"2026-08-05T14:02:06.442587Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1803.07728","last_updated":"2018-03-21T03:21:14Z","snapshot_observed_at":"2026-08-09T00:07:35.136537Z","submitted_at":"2018-03-21T03:21:14Z","title":"Unsupervised Representation Learning by Predicting Image Rotations","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1803.07728","snapshot_observed_at":"2026-08-05T14:02:00.976755Z","title":"Unsupervised representation learning by predicting image rotations.arXiv preprint arXiv:1803.07728, 2018","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:00.976755Z"},"links":{"cited_paper":"/paper/1803.07728","citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:8b74c351f8185ed5f0a4362f4e266cff820d8152391cf3aa2cfc3751e05530c3","observation_id":"f335c71d-ca77-4eab-8624-febaea0ce827","resolution":{"observed_at":"2026-08-05T14:02:00.976755Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:12.981696Z","title":"Dense open-set recognition with syn- thetic outliers generated by real nvp","venue":null,"work_id":"18b613c2-6d36-411c-af3a-ac5bba61dec9","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:01.111370Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:47cf5963dfae32c296b15228578f9c8bdd9aa40255742da60aba3e4503b10606","observation_id":"a6b816dc-6ac6-4e48-a175-e2b833a6c4dd","resolution":{"observed_at":"2026-08-05T14:02:13.105602Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:12.875659Z","title":"Learning to discover novel visual categories via deep transfer clustering","venue":null,"work_id":"0fdc18ff-85d8-4856-8977-29a8deeca8ab","year":2019},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:01.186296Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:e81f07fd6ff677c12b15cca002cc7e54484b17671d170a0e7344805a91fb7461","observation_id":"3221c6d5-d2aa-47d4-b4f8-ecda163b8460","resolution":{"observed_at":"2026-08-05T14:02:12.930902Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:12.697963Z","title":"Automatically discovering and learning new visual categories with rank- ing statistics","venue":null,"work_id":"a1701605-8b55-4986-b269-c689742baa7c","year":2020},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:01.306892Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:e85758033b8e2226e196661decc65c43cffb9f352f3992a0981490ef81ad35be","observation_id":"1bdba9f6-4930-4281-bcb3-8e5bfe56a1a1","resolution":{"observed_at":"2026-08-05T14:02:12.767073Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:12.544003Z","title":"Autonovel: Automatically discovering and learning novel visual cate- gories.IEEE Transactions on Pattern Analysis and Machine Intelligence, 44(10): 6767–6781, 2021","venue":null,"work_id":"5c00361e-2506-489d-9383-8f34fffc10dd","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:01.420258Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:faffa94448340cd0004994708e7f7f72f39c26523183b8083ca9f92f7c5fd466","observation_id":"76da8525-94fa-4fce-9c5a-82b214854fe7","resolution":{"observed_at":"2026-08-05T14:02:12.611309Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:12.391135Z","title":"Can spatiotemporal 3d cnns retrace the history of 2d cnns and imagenet? InProceedings of the IEEE conference on Computer Vision and Pattern Recognition, pages 6546–6555, 2018","venue":null,"work_id":"e75d3fa2-7e28-412a-8796-90ef8d5a69f4","year":2018},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:01.568984Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:9e5aebeacf2787563cebe16997b0d611f00de34ca9c58cc3c6b629b9b6a06557","observation_id":"f4e04a7b-f4d4-4f3d-a2c1-be977b73d837","resolution":{"observed_at":"2026-08-05T14:02:12.465106Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:12.199479Z","title":"Video owl-vit: Temporally-consistent open-world localization in video","venue":null,"work_id":"01c38d32-8d64-40a6-9c1e-ad90723f7ee7","year":2023},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:01.711904Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:fc2d15403c9e00c97488eef5056553c161c7c264d6e388bf61a3e69e88402533","observation_id":"2fc44e2b-0ff3-4f97-a762-be465276dca6","resolution":{"observed_at":"2026-08-05T14:02:12.268186Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:12.018722Z","title":"A baseline for detecting misclassified and out- of-distribution examples in neural networks.International Conference on Learning Representations, 2017","venue":null,"work_id":"5ae33304-649a-499b-b8cb-36fea005355f","year":2017},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:01.859472Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:1f8131afdc68f0f13345ee521f7c2fe0dbe2d1cf010b116a4623c962f154900a","observation_id":"dbfb2f24-459b-4e98-935d-61a31caf0d61","resolution":{"observed_at":"2026-08-05T14:02:12.100141Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:11.872303Z","title":"Deep anomaly detection with outlier exposure.International Conference on Learning Representations, 2019","venue":null,"work_id":"eb42356c-72b3-4a4c-904d-0d8430cfc144","year":2019},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:02.004624Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:4f980843826577ca5b263a6adbf092b0ea2721c273727e037f9f0f11046a0abf","observation_id":"eefb184b-9bcd-49a5-9baa-0fa8f33eea86","resolution":{"observed_at":"2026-08-05T14:02:11.936421Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:11.717364Z","title":"Generalized odin: De- tecting out-of-distribution image without learning from out-of-distribution data","venue":null,"work_id":"87e79c76-1034-4b9b-8b7a-0b3933782412","year":2020},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:02.125505Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:f0ffa39e037e8c433c9d79279339eaed6926d61773424e15c5c42af402c7361d","observation_id":"86fe09e9-ec1f-4529-a859-ab67b124e03b","resolution":{"observed_at":"2026-08-05T14:02:11.790544Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1705.06950","last_updated":"2017-05-19T12:07:01Z","snapshot_observed_at":"2026-08-08T17:46:50.107463Z","submitted_at":"2017-05-19T12:07:01Z","title":"The Kinetics Human Action Video Dataset","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1705.06950","snapshot_observed_at":"2026-08-05T14:02:02.276463Z","title":"The ki- netics human action video dataset.arXiv preprint arXiv:1705.06950, 2017","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:02.276463Z"},"links":{"cited_paper":"/paper/1705.06950","citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:8bcdb9cc1ec8340eb624877673961b2599227f74c5f40898b1966aedc2c7aed2","observation_id":"cad9770e-b561-4603-8198-34cf301899d0","resolution":{"observed_at":"2026-08-05T14:02:02.276463Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:11.597749Z","title":"Chal- lenges, evaluation and opportunities for open-world learning.Nature Machine Intelli- gence, 6(6):580–588, 2024","venue":null,"work_id":"42552d8e-bca1-48b7-8bb2-b9ab4ff0d1b2","year":2024},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:02.440137Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:6272dccef4633959da9a128c383386c3f3026dd3f9a27d605b5884f3aac869fb","observation_id":"c5621063-d162-46f4-ba36-664fea4ad350","resolution":{"observed_at":"2026-08-05T14:02:11.655563Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:11.460188Z","title":"Opengan: Open-set recognition via open data genera- tion","venue":null,"work_id":"5cd0c672-6604-4ebd-8026-c2faa37b918f","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:02.623178Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:cc3140dc08bc244da2bcae04539e1af880c695a6d25ef5222236196e61c87111","observation_id":"8e4b0570-c80e-4a39-a6c1-a4438d3c6fd9","resolution":{"observed_at":"2026-08-05T14:02:11.515219Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:11.303215Z","title":"Human action recognition and prediction: A survey.Interna- tional Journal of Computer Vision, 130(5):1366–1401, 2022","venue":null,"work_id":"dc49a8ec-b2f3-4748-8353-219c630048fa","year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:02.758084Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:7be6e3405b035d3188b2595f79c73e1f41bf06150446b1a01e3bb79ed344e81e","observation_id":"22802874-4449-4951-8f6e-2790280e88bd","resolution":{"observed_at":"2026-08-05T14:02:11.384104Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:11.148662Z","title":"Hmdb: a large video database for human motion recognition","venue":null,"work_id":"9886281c-b650-41aa-93cf-27ce30ccddd6","year":2011},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:02.899424Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:762b9ea28ad3c4336c7da035ca8aa88f37102291027b0c9df0ab8126bd5bb6bc","observation_id":"fbd90b24-03dd-4766-b5d8-e0a45039c2de","resolution":{"observed_at":"2026-08-05T14:02:11.220752Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:10.990380Z","title":"Resound: Towards action recognition without representation bias","venue":null,"work_id":"f17260ac-8794-467c-ad61-0acd9df1a000","year":2018},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:02.986988Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:d1c063567c9f41c014d696e01b2d93fa5376cc9788fcd1dfec928e5a1fb8c529","observation_id":"ad014836-6e04-471f-862e-4581f5b01974","resolution":{"observed_at":"2026-08-05T14:02:11.081851Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:10.845329Z","title":"Resource-rational analysis: Understanding hu- man cognition as the optimal use of limited computational resources.Behavioral and brain sciences, 43:e1, 2020","venue":null,"work_id":"f85a3a04-0511-4262-bce2-5eb60975bfee","year":2020},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.039359Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:44eac586fa5098cdb1e65ad12b4941e9f32e4534b7b4fb7dde46704856e6f13f","observation_id":"a0c76e42-ebf0-4991-a94e-cad8cda903de","resolution":{"observed_at":"2026-08-05T14:02:10.916377Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:10.726175Z","title":"Abnormal event detection at 150 fps in matlab","venue":null,"work_id":"17efd3d0-bfb0-4fe0-896a-5db663fc8bf2","year":2013},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.094171Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:edfca163a857b99c3e3e65a68200b5d8f4536e4d6439ec9f05391b7a50ff4769","observation_id":"613db08a-00f3-4c9e-b2cd-feed15629500","resolution":{"observed_at":"2026-08-05T14:02:10.763141Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2010.55398","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:06.175144Z","title":"Anomaly de- tection in crowded scenes","venue":null,"work_id":"b9eaac4a-ee89-451b-bad4-df27b15d1bdb","year":1975},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.136950Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:d0eddfb4a9c9f2d14d9e2e65f438733c66f45697f31f146dda23b5653e0053f7","observation_id":"0d46ae1c-2184-4541-86e7-a7acb2973982","resolution":{"observed_at":"2026-08-05T14:02:06.229836Z","resolver_source":"raw_fallback","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:10.637310Z","title":"HowTo100M: Learning a Text-Video Embedding by Watch- ing Hundred Million Narrated Video Clips","venue":null,"work_id":"b5b24b5a-3f04-46f9-a01a-f26ca5354fe2","year":2019},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.316713Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:cfc94da45be0890113e2ddbd9acb1e695d15332015660ad4dbf8d418f6c523c9","observation_id":"1618ee6f-091a-4d8e-bfd4-0d1ecedda25f","resolution":{"observed_at":"2026-08-05T14:02:10.680199Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:10.478005Z","title":"Poem: Out-of-distribution detection with pos- terior sampling","venue":null,"work_id":"198c0c11-ca33-4004-8061-0074a1187438","year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.507533Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:e2f45ac79c7b739a2865e82ff1624fc23fe884fabebecb1f1684b2c8af1e6a3e","observation_id":"8af1eb24-6367-4916-be5b-33cfbbe09b3c","resolution":{"observed_at":"2026-08-05T14:02:10.551577Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:10.332285Z","title":"Mo- ments in time dataset: one million videos for event understanding.IEEE transactions on pattern analysis and machine intelligence, 42(2):502–508, 2019","venue":null,"work_id":"f971dadd-9329-4c35-a541-5f0c4c10f57d","year":2019},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.553269Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:435699f9aac3a756a62f0e5ab7dc0452937ca53a8a64d25b12f5da3f29995f97","observation_id":"4ab5466a-f3e6-4ce8-9f66-58b1a945f861","resolution":{"observed_at":"2026-08-05T14:02:10.381973Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:03.594472Z","title":"Expanding language-image pretrained models for general video recognition","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.594472Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:0735cb033559953664e4aee94e4b954376ff53c479f52a55346d680bdab37f6a","observation_id":"8673dd7b-28cf-4cb7-a980-3e24a9fa7645","resolution":{"observed_at":"2026-08-05T14:02:03.594472Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:10.190137Z","title":"Outlier exposure with confidence control for out-of-distribution detec- tion.Neurocomputing, 441:138–150, 2021","venue":null,"work_id":"5d87baa2-1cb6-4fb7-b18d-10bae442bb9c","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.663952Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:e7535722a523f73ed331551f43cff122b5fbe3830164838dcb92ec78e74fcddb","observation_id":"5965cb6c-c003-492c-b3eb-b552385ce969","resolution":{"observed_at":"2026-08-05T14:02:10.250264Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:10.043526Z","title":"A survey on vision-based human action recognition.Image and vision computing, 28(6):976–990, 2010","venue":null,"work_id":"cb8d93bd-b46d-4c61-87f5-3c7d048a654a","year":2010},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.762859Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:3f08d79217ba89aaae6174d94a877fdfed54190680b4178c7d1b87e71f7908e9","observation_id":"43b0e379-fa27-46aa-946b-b97cb0ae7e7f","resolution":{"observed_at":"2026-08-05T14:02:10.120128Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:09.964165Z","title":"Learning transferable visual models from natural language supervision","venue":null,"work_id":"a2f95382-ad3f-4605-979a-833380704c30","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.856227Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:7102c37099e9dea27e0ea86519e157ac9539300ed687bb5450e2eeda0c11f798","observation_id":"0c88274e-2914-4fa7-b1ab-7574ef637269","resolution":{"observed_at":"2026-08-05T14:02:09.995951Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:09.860676Z","title":"Fishr: Invariant gradient variances for out-of-distribution generalization","venue":null,"work_id":"5433444b-08c3-4482-90ba-1374eb51d845","year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.957892Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:fd2da81f354781a457a1a02ae3c52166a819aa9c586ed199b65c36ce1a52f34e","observation_id":"15d0e835-873b-42e7-88cc-876999ec6855","resolution":{"observed_at":"2026-08-05T14:02:09.893998Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:04.067249Z","title":"If deep learning is the answer, what is the question?Nature Reviews Neuroscience, 22(1):55–67, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.067249Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:d87075b7652f2006597a44165ebdc17bc69ce6738c9cb47fbbff40bdfa88d0b0","observation_id":"2a510c09-25be-41ba-ad56-6e5e07bf51cb","resolution":{"observed_at":"2026-08-05T14:02:04.067249Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:09.670821Z","title":"Meta- recognition: The theory and practice of recognition score analysis.IEEE transactions on pattern analysis and machine intelligence, 33(8):1689–1695, 2011","venue":null,"work_id":"f174890c-d49b-40b2-9298-3b47f0e82c50","year":2011},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.147325Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:872dd17cfa43e8db0902f412f2e8c131a514395c86209250c7515cd58cd89e34","observation_id":"7f0b9564-2abd-4f87-ab31-6325e71330ce","resolution":{"observed_at":"2026-08-05T14:02:09.751133Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1212.0402","last_updated":"2012-12-03T14:45:31Z","snapshot_observed_at":"2026-07-06T03:01:10.229407Z","submitted_at":"2012-12-03T14:45:31Z","title":"UCF101: A Dataset of 101 Human Actions Classes From Videos in The Wild","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1212.0402","snapshot_observed_at":"2026-08-05T14:02:04.226399Z","title":"Ucf101: A dataset of 101 human actions classes from videos in the wild","venue":null,"work_id":null,"year":2012},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.226399Z"},"links":{"cited_paper":"/paper/1212.0402","citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:6b4997d5818208f8a4fdebc230a7d8a0927b801eae7db34849c8cb9cf483a464","observation_id":"4d9c788d-165c-433b-af06-00aba146c9e3","resolution":{"observed_at":"2026-08-05T14:02:04.226399Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:09.481172Z","title":"Real-world anomaly detection in surveillance videos","venue":null,"work_id":"7c88b126-054f-4e36-826a-83c876b693da","year":2018},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.324052Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:bd9cc44cf5e5feaaa015139146e6f93c131798eeef4a69a01607ab9765fcf40d","observation_id":"c700c04d-6bc3-40c1-b525-fbf7175aa152","resolution":{"observed_at":"2026-08-05T14:02:09.577098Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:09.286183Z","title":"Human action recognition from various data modalities: A review.IEEE transactions on pattern analysis and machine intelligence, 45(3):3200–3225, 2022","venue":null,"work_id":"691a7c3a-e2cf-44a2-b502-5b7d402748e6","year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.402105Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:2ba8a7021010551d24ace1df6a47fe370e2dd9cd5d63d504d81d11d9c65c3249","observation_id":"e856d248-1424-47e0-b9f9-f84507e96a64","resolution":{"observed_at":"2026-08-05T14:02:09.378542Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:09.092110Z","title":"Black, Ivan Laptev, and Cordelia Schmid","venue":null,"work_id":"24e8014e-5193-45e9-952e-a031469880b3","year":2017},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.464878Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:770e8c31a3af4ef0a4fa77c0be5e7344ce2da30fb4a614fd4ec5574ff2cfe3ba","observation_id":"e71b4814-8007-4ffb-a313-a972afcad7fe","resolution":{"observed_at":"2026-08-05T14:02:09.193093Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:08.956029Z","title":"Open-set recognition: A good closed-set classifier is all you need? 2021","venue":null,"work_id":"5cbce019-aca9-4d72-bf2b-5cf73b59288d","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.557425Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:bad972e2289f3a7c2ce2314565d63e47bf1c528e8b9c6d1308058874387a60a3","observation_id":"4afc6e3f-d632-420d-babe-95261d84a10e","resolution":{"observed_at":"2026-08-05T14:02:08.991793Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:08.745086Z","title":"Generalized category discovery","venue":null,"work_id":"0104727d-1eff-4fdc-991b-d1be0fb00a7f","year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.650860Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:497dd3e5a0275d260d0232216e45da2ce2e87708f29bf4712dafb92f79393faf","observation_id":"bedd9d56-d3cb-436c-90a5-96bb781ca5fc","resolution":{"observed_at":"2026-08-05T14:02:08.881327Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:08.563969Z","title":"No representation rules them all in category discovery","venue":null,"work_id":"c3e57b95-7278-4246-a357-315e327786d1","year":2023},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.718218Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:839e17d7c7ad6362a920494933f0c3ff6fcb9e26e334046150773a410ba80f22","observation_id":"fec39636-34ab-453b-9c37-6666dc260ca5","resolution":{"observed_at":"2026-08-05T14:02:08.640276Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:08.404555Z","title":"Self-supervised video representation learning by pace prediction","venue":null,"work_id":"d0271324-488e-4a40-a498-744efa9f2adb","year":2020},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.787888Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:541df95a881ff8ed84237f03a3db26a4ea37eaed997fb9081963d5bdaabecb37","observation_id":"5bac9760-7640-4d5d-8880-f111f096f965","resolution":{"observed_at":"2026-08-05T14:02:08.480574Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:08.190374Z","title":"Action recognition and detection by com- bining motion and appearance features.THUMOS14 Action Recognition Challenge, 1 (2):2, 2014","venue":null,"work_id":"f52e60c6-f241-480a-995b-bf044c073590","year":2014},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.881590Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:010b9b5fe3ed85783d35cd5be4f5b4d18f52936935f31660f344a7e9467b10f1","observation_id":"ac66d5d3-b773-4a9e-b045-3b934958f55b","resolution":{"observed_at":"2026-08-05T14:02:08.289468Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2109.08472","last_updated":"2021-09-17T11:21:34Z","snapshot_observed_at":"2026-07-06T11:48:42.000083Z","submitted_at":"2021-09-17T11:21:34Z","title":"ActionCLIP: A New Paradigm for Video Action Recognition","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2109.08472","snapshot_observed_at":"2026-08-05T14:02:04.935174Z","title":"Actionclip: A new paradigm for video action recognition.arXiv preprint arXiv:2109.08472, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.935174Z"},"links":{"cited_paper":"/paper/2109.08472","citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:4a83920acf3e8f27f38bd618554947bb3252feb5dc6cc2005b8132c06e2776f3","observation_id":"3d76a287-af6c-469c-bde0-b99ccc3cfd3f","resolution":{"observed_at":"2026-08-05T14:02:04.935174Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:07.971277Z","title":"Openood: Benchmark- ing generalized out-of-distribution detection.Advances in Neural Information Pro- cessing Systems, 35:32598–32611, 2022","venue":null,"work_id":"0d082054-aabf-4d75-a230-0f56f2024098","year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.014884Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:4408b8b3e42af24aa90b368f3e07e1f54c5bad41ed7dd93b9f0e8b893bddca88","observation_id":"3a16aaa4-b86f-4306-9abd-019f955ab3f0","resolution":{"observed_at":"2026-08-05T14:02:08.072192Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:07.744214Z","title":"Generalized out-of- distribution detection: A survey.International Journal of Computer Vision, pages 1–28, 2024","venue":null,"work_id":"9a9d406b-a09b-4868-9801-8b7605939932","year":2024},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.103596Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:f4c067d0993db2830b0e170f4376ded3365c556bafc1a9fa5f809da3d42cb4b2","observation_id":"5e5073cc-84dd-4569-b06c-64f73cb7855a","resolution":{"observed_at":"2026-08-05T14:02:07.865968Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:07.548914Z","title":"Un- derstanding deep learning (still) requires rethinking generalization.Communications of the ACM, 64(3):107–115, 2021","venue":null,"work_id":"3d9ca34d-57b4-4088-9722-26ef7488d09c","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.216117Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:08aa3012c39077d7ea6df556007f89669a73e0167dd980a3e342d8b2f4c061ec","observation_id":"9b7c7aa8-c1e7-479a-a0ba-930dae3c597b","resolution":{"observed_at":"2026-08-05T14:02:07.645931Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:07.291430Z","title":"Mixture outlier exposure: Towards out-of-distribution detection in fine-grained en- vironments","venue":null,"work_id":"b291d56d-a4d3-48ad-8e50-e7132cef9f8a","year":2023},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.274180Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:48002b7e71695e1c01dd9ba57ef5d3116dadae5de3e5a4a7ed332e403e7ac4c6","observation_id":"c095afd6-0bd8-45f0-b4fa-af9e02b90dfe","resolution":{"observed_at":"2026-08-05T14:02:07.407597Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:07.138236Z","title":"Open- mix: Reviving known knowledge for discovering novel visual categories in an open world","venue":null,"work_id":"c85f5d75-8862-4ca4-b5db-5638ecaf0de9","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.381623Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:ffa011e48dd9ef0c0d7e60ce0e6b0214fa9c4c80706d39c0fa580f77408cf607","observation_id":"47a6f305-ac19-41bb-9aaf-b037d21b52b0","resolution":{"observed_at":"2026-08-05T14:02:07.199110Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:07.033945Z","title":"Learning placeholders for open-set recognition","venue":null,"work_id":"72e6f60e-48f2-4303-8512-10c0999117e1","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.466295Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:5445b041e6e58ed99961d484dea5788e408e389fc46b9edd2d57daecb9716183","observation_id":"c24e56f8-18e8-42a7-9e9e-934521c96303","resolution":{"observed_at":"2026-08-05T14:02:07.062700Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:06.875909Z","title":"Diversified outlier exposure for out-of-distribution detection via in- formative extrapolation.Advances in Neural Information Processing Systems, 36: 22702–22734, 2023","venue":null,"work_id":"abf00742-d2ab-437a-bde6-7da9b9990d50","year":2023},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.542651Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:2198396bf03e0257a842f1b04dcf0e465d337d6f80cc75d8fd2ac4766fa21c42","observation_id":"aac02e99-d669-4272-bca6-ee4f99b20c7c","resolution":{"observed_at":"2026-08-05T14:02:06.914683Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:06.747330Z","title":"Unleashing mask: Explore the intrinsic out-of-distribution detection capa- bility","venue":null,"work_id":"1cffd789-b024-4ff2-a32d-11168d88da73","year":2023},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.634486Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:299297a7fa946c74ae75f89565ab1dafff33bde6f29f266aa9b6ab0fa31ec068","observation_id":"3c238e71-90f7-44f6-bf6e-4e3031c6bbd7","resolution":{"observed_at":"2026-08-05T14:02:06.803921Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:05.763997Z","title":"Towards universal representation for unseen action recognition","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.763997Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:250a3b010956e380116126be59f5ae9c1f8f53a47e61a3a45d64d65437a992f3","observation_id":"33fcc031-32cb-4f0a-afcd-e5f7ede42976","resolution":{"observed_at":"2026-08-05T14:02:05.763997Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:06.548016Z","title":"Towards open set video anomaly detec- tion","venue":null,"work_id":"1b9515f6-0cb3-4902-9e91-21332b97c59c","year":null},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.872722Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:0d860a8bf30a4e2b5a6acea6bd3c05a9fbb1573a72e7af4a435f1688b601caa1","observation_id":"2722e854-314d-452b-90d6-c81958eb97ce","resolution":{"observed_at":"2026-08-05T14:02:06.642411Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-08T11:57:14.434723Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos"},"reference_resolution":{"displayed":59,"state_counts":{"malformed_identifier":0,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":8,"verified_exact":1,"verified_fuzzy":49},"total_outbound_references":59},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 59 of 59 outbound references and 0 inbound Pith citation observations for arXiv:2508.21770."}