{"as_of":"2026-08-10T12:09:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:2c2c646144a955befcd6d04fc70e1fcbb59d9dfeb4fdaa996d88528fb4de4ae0","coverage":[{"denominator":59,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":59,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T14:02:05.872722Z","state":"measured"},{"denominator":59,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":59,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2508.21770/citation-record","integrity":"/paper/2508.21770/integrity","json":"/paper/2508.21770/citation-record.json","paper":"/paper/2508.21770"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:14.590184Z","title":"Ubnor- mal: New benchmark for supervised open-set video anomaly detection","venue":null,"work_id":"319068f0-ea10-4134-987f-90da64c29a19","year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-05T14:01:59.564950Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:3eba2e42fe706acfc03798b338bdb367bcda7453c27e863b7e2e4a1f45f37a49","observation_id":"c494610a-58b0-4d31-a7b6-0c80420f89ef","resolution":{"observed_at":"2026-08-05T14:02:14.607104Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:14.414086Z","title":"Towards open set deep networks","venue":null,"work_id":"0cdba8d6-0a33-4177-aa76-a50b717ed8b4","year":2016},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-05T14:01:59.741279Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:be67f49c2b30d7f64e730bd7d74ff999d5e7fe9dca195e769022663b598dace0","observation_id":"4fccf45d-4b79-4a40-95d8-dd36319a87e3","resolution":{"observed_at":"2026-08-05T14:02:14.498084Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:14.289633Z","title":"Is space-time attention all you need for video understanding? InProceedings of the International Conference on Machine Learning (ICML), July 2021","venue":null,"work_id":"e6029011-6371-4cca-962d-72b4327ee56c","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-05T14:01:59.975483Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:a6b3e79cf2c14cbe1ea999ec03d890f791daa9df2e908d995d42321dd123c470","observation_id":"cc1f2ea9-6283-4560-9b31-dfbea8634f0f","resolution":{"observed_at":"2026-08-05T14:02:14.361736Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:14.139312Z","title":"Enlarging instance-specific and class-specific information for open-set action recognition","venue":null,"work_id":"dd9e15a3-e759-49e7-a5f2-266cb6114254","year":2023},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:00.109712Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:3028a7ebef88e96281dd458c636155927863d232d3803c52ca24e000222334b1","observation_id":"500f5cfb-9b58-42ce-a045-cd8a50f365ec","resolution":{"observed_at":"2026-08-05T14:02:14.188134Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:00.192110Z","title":"Elaborative rehearsal for zero-shot action recognition","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:00.192110Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:bf10bbaade784297c8906a8d9a64179b597676e430042f315e1e63f192cf1be0","observation_id":"3a761539-6bbd-443a-baae-8ed237cb4a19","resolution":{"observed_at":"2026-08-05T14:02:00.192110Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:13.993825Z","title":"Wdiscood: Out-of- distribution detection via whitened linear discriminant analysis","venue":null,"work_id":"a57e6282-bf9c-448e-a099-ad91e591b64d","year":2023},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:00.348788Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:e13d24111c5f1ba3bc4a94d57c73fd0cd83849e5db6416b768697c1230ceb2ef","observation_id":"ce7c0c73-129b-4b78-8fb7-fb430e31df85","resolution":{"observed_at":"2026-08-05T14:02:14.024524Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:13.836386Z","title":"Haa500: Human-centric atomic action dataset with curated videos","venue":null,"work_id":"d3e1b2e0-02b3-44f4-88ec-1610cf2ab0d9","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:00.439028Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:adfcee48446a859f525ba58898bf32ab18d0939438492347f51770039d41314e","observation_id":"1c77dc5e-7b2b-4290-8ced-d051ba3fbbdc","resolution":{"observed_at":"2026-08-05T14:02:13.940558Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:13.481648Z","title":"Towards unknown-aware learning with virtual outlier synthesis","venue":null,"work_id":"f4cfd443-06fc-496e-940f-599e60c4217d","year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:00.578152Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:1d92c6aba82e2aea367dafe16f7295f492748de382c550a10d6c4c7f61e8868e","observation_id":"5a4adc85-ee08-4796-916e-c98c8e7905db","resolution":{"observed_at":"2026-08-05T14:02:13.644210Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:13.233276Z","title":"Oops! predicting unintentional ac- tion in video","venue":null,"work_id":"db09e466-2130-42e9-b431-943a5af5efa9","year":2020},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:00.696918Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:ff58ac166a00d22731f43f22985083ecd6fe6c238d6f7ad08df04995d3ec9703","observation_id":"b5dffb49-5c8c-4155-b694-976d266a0012","resolution":{"observed_at":"2026-08-05T14:02:13.343116Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.11736","last_updated":"2023-03-28T18:27:24Z","snapshot_observed_at":"2026-08-07T10:53:37.587234Z","submitted_at":"2022-06-23T14:31:33Z","title":"NovelCraft: A Dataset for Novelty Detection and Discovery in Open Worlds","version":3},"cited_work":{"arxiv_id":"2206.11736","doi":null,"metadata_source":"pith","pith_arxiv_id":"2206.11736","snapshot_observed_at":"2026-08-05T14:02:06.394120Z","title":"NovelCraft: A Dataset for Novelty Detection and Discovery in Open Worlds","venue":"cs.CV","work_id":"5fa8a380-af3c-4648-9983-3369a9ee0ce0","year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:00.857266Z"},"links":{"cited_paper":"/paper/2206.11736","citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:229c10009321aeafb591a695f22a7a5be22f7e227111da8e175e37ce7f566f53","observation_id":"4b29c84f-1298-4166-b031-c841b3872aa5","resolution":{"observed_at":"2026-08-05T14:02:06.442587Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1803.07728","last_updated":"2018-03-21T03:21:14Z","snapshot_observed_at":"2026-08-10T04:03:30.714422Z","submitted_at":"2018-03-21T03:21:14Z","title":"Unsupervised Representation Learning by Predicting Image Rotations","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1803.07728","snapshot_observed_at":"2026-08-05T14:02:00.976755Z","title":"Unsupervised representation learning by predicting image rotations.arXiv preprint arXiv:1803.07728, 2018","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:00.976755Z"},"links":{"cited_paper":"/paper/1803.07728","citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:249cc710f4a06ffb7915e90a5336f8a2cb723ef398f01c28724266dd5807a9d9","observation_id":"f335c71d-ca77-4eab-8624-febaea0ce827","resolution":{"observed_at":"2026-08-05T14:02:00.976755Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:12.981696Z","title":"Dense open-set recognition with syn- thetic outliers generated by real nvp","venue":null,"work_id":"18b613c2-6d36-411c-af3a-ac5bba61dec9","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:01.111370Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:ccc73e6dccafa5bac3534da2ed6ae983e1cbeea53b4ba02c4edfb211c4083481","observation_id":"a6b816dc-6ac6-4e48-a175-e2b833a6c4dd","resolution":{"observed_at":"2026-08-05T14:02:13.105602Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:12.875659Z","title":"Learning to discover novel visual categories via deep transfer clustering","venue":null,"work_id":"0fdc18ff-85d8-4856-8977-29a8deeca8ab","year":2019},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:01.186296Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:95a0e6333ebe8a5420c64de5214608d79478e2842fbb5a82d550aa25a3f3d898","observation_id":"3221c6d5-d2aa-47d4-b4f8-ecda163b8460","resolution":{"observed_at":"2026-08-05T14:02:12.930902Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:12.697963Z","title":"Automatically discovering and learning new visual categories with rank- ing statistics","venue":null,"work_id":"a1701605-8b55-4986-b269-c689742baa7c","year":2020},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:01.306892Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:fddc924204d7de6ea23c2909b654c8e9e3590a4ccf224b5873924c1a47928452","observation_id":"1bdba9f6-4930-4281-bcb3-8e5bfe56a1a1","resolution":{"observed_at":"2026-08-05T14:02:12.767073Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:12.544003Z","title":"Autonovel: Automatically discovering and learning novel visual cate- gories.IEEE Transactions on Pattern Analysis and Machine Intelligence, 44(10): 6767–6781, 2021","venue":null,"work_id":"5c00361e-2506-489d-9383-8f34fffc10dd","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:01.420258Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:254df0dab0aca3e9dfe867de4c9738adbcd6dc55f26c6bf7d5c12be3ea503f5a","observation_id":"76da8525-94fa-4fce-9c5a-82b214854fe7","resolution":{"observed_at":"2026-08-05T14:02:12.611309Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:12.391135Z","title":"Can spatiotemporal 3d cnns retrace the history of 2d cnns and imagenet? InProceedings of the IEEE conference on Computer Vision and Pattern Recognition, pages 6546–6555, 2018","venue":null,"work_id":"e75d3fa2-7e28-412a-8796-90ef8d5a69f4","year":2018},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:01.568984Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:fb44755204219f83653de10a65eced02517127600651a26f27903a45c4886b31","observation_id":"f4e04a7b-f4d4-4f3d-a2c1-be977b73d837","resolution":{"observed_at":"2026-08-05T14:02:12.465106Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:12.199479Z","title":"Video owl-vit: Temporally-consistent open-world localization in video","venue":null,"work_id":"01c38d32-8d64-40a6-9c1e-ad90723f7ee7","year":2023},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:01.711904Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:f961024d4ef99ad661a2eebb740aed8353100e3c08ff65fbb529dfd57f2facd6","observation_id":"2fc44e2b-0ff3-4f97-a762-be465276dca6","resolution":{"observed_at":"2026-08-05T14:02:12.268186Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:12.018722Z","title":"A baseline for detecting misclassified and out- of-distribution examples in neural networks.International Conference on Learning Representations, 2017","venue":null,"work_id":"5ae33304-649a-499b-b8cb-36fea005355f","year":2017},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:01.859472Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:9bd55c4b8b476821d8e34c21cfb35de5790fe1e985329b2d0927a7719a58642b","observation_id":"dbfb2f24-459b-4e98-935d-61a31caf0d61","resolution":{"observed_at":"2026-08-05T14:02:12.100141Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:11.872303Z","title":"Deep anomaly detection with outlier exposure.International Conference on Learning Representations, 2019","venue":null,"work_id":"eb42356c-72b3-4a4c-904d-0d8430cfc144","year":2019},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:02.004624Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:cc6621eca560ff7f667beb81d1b8712d5b2b8103b739d52f16e203a332172843","observation_id":"eefb184b-9bcd-49a5-9baa-0fa8f33eea86","resolution":{"observed_at":"2026-08-05T14:02:11.936421Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:11.717364Z","title":"Generalized odin: De- tecting out-of-distribution image without learning from out-of-distribution data","venue":null,"work_id":"87e79c76-1034-4b9b-8b7a-0b3933782412","year":2020},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:02.125505Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:1bc1d84784861dbf03825c8bd198c9528f1a9013c1d03125bcc63133147f2262","observation_id":"86fe09e9-ec1f-4529-a859-ab67b124e03b","resolution":{"observed_at":"2026-08-05T14:02:11.790544Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1705.06950","last_updated":"2017-05-19T12:07:01Z","snapshot_observed_at":"2026-08-08T17:46:50.107463Z","submitted_at":"2017-05-19T12:07:01Z","title":"The Kinetics Human Action Video Dataset","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1705.06950","snapshot_observed_at":"2026-08-05T14:02:02.276463Z","title":"The ki- netics human action video dataset.arXiv preprint arXiv:1705.06950, 2017","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:02.276463Z"},"links":{"cited_paper":"/paper/1705.06950","citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:601ca433f7c50551ad152a0ab75f1f6d1c7602fedddfd7a8ea7ce6a88508bb12","observation_id":"cad9770e-b561-4603-8198-34cf301899d0","resolution":{"observed_at":"2026-08-05T14:02:02.276463Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:11.597749Z","title":"Chal- lenges, evaluation and opportunities for open-world learning.Nature Machine Intelli- gence, 6(6):580–588, 2024","venue":null,"work_id":"42552d8e-bca1-48b7-8bb2-b9ab4ff0d1b2","year":2024},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:02.440137Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:f9186fac41d780c68118bc197ffbffe048a45fd023a82084dfdb981fbbcd8f26","observation_id":"c5621063-d162-46f4-ba36-664fea4ad350","resolution":{"observed_at":"2026-08-05T14:02:11.655563Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:11.460188Z","title":"Opengan: Open-set recognition via open data genera- tion","venue":null,"work_id":"5cd0c672-6604-4ebd-8026-c2faa37b918f","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:02.623178Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:05a20b74fd3f68461537c3e7470e70a7bef3df512d1ea4700161e7301744da8b","observation_id":"8e4b0570-c80e-4a39-a6c1-a4438d3c6fd9","resolution":{"observed_at":"2026-08-05T14:02:11.515219Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:11.303215Z","title":"Human action recognition and prediction: A survey.Interna- tional Journal of Computer Vision, 130(5):1366–1401, 2022","venue":null,"work_id":"dc49a8ec-b2f3-4748-8353-219c630048fa","year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:02.758084Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:5b6980121982bf13993b3ca1a5039a5851320090450c56aeb0fc0c3e2eab46ec","observation_id":"22802874-4449-4951-8f6e-2790280e88bd","resolution":{"observed_at":"2026-08-05T14:02:11.384104Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:11.148662Z","title":"Hmdb: a large video database for human motion recognition","venue":null,"work_id":"9886281c-b650-41aa-93cf-27ce30ccddd6","year":2011},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:02.899424Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:1b07824c23f88ce5b233ac9bb7f18a5f5c71c9fc2a53cca1a184c0e8118c42d3","observation_id":"fbd90b24-03dd-4766-b5d8-e0a45039c2de","resolution":{"observed_at":"2026-08-05T14:02:11.220752Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:10.990380Z","title":"Resound: Towards action recognition without representation bias","venue":null,"work_id":"f17260ac-8794-467c-ad61-0acd9df1a000","year":2018},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:02.986988Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:b7023efa879b3839795f11aeac2022d1a7ec546e99b01b8e4b983ca3d2308731","observation_id":"ad014836-6e04-471f-862e-4581f5b01974","resolution":{"observed_at":"2026-08-05T14:02:11.081851Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:10.845329Z","title":"Resource-rational analysis: Understanding hu- man cognition as the optimal use of limited computational resources.Behavioral and brain sciences, 43:e1, 2020","venue":null,"work_id":"f85a3a04-0511-4262-bce2-5eb60975bfee","year":2020},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.039359Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:7a05501f579348f7540476b8ba6e47201a6d43e7743e14955f73936f4bc0e12d","observation_id":"a0c76e42-ebf0-4991-a94e-cad8cda903de","resolution":{"observed_at":"2026-08-05T14:02:10.916377Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:10.726175Z","title":"Abnormal event detection at 150 fps in matlab","venue":null,"work_id":"17efd3d0-bfb0-4fe0-896a-5db663fc8bf2","year":2013},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.094171Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:02ab63a219f560d3ddb00731aa1b10ceb4fe055b33d3ca71df2b78bbbf9ee259","observation_id":"613db08a-00f3-4c9e-b2cd-feed15629500","resolution":{"observed_at":"2026-08-05T14:02:10.763141Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2010.55398","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:06.175144Z","title":"Anomaly de- tection in crowded scenes","venue":null,"work_id":"b9eaac4a-ee89-451b-bad4-df27b15d1bdb","year":1975},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.136950Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:96b61118d67e6e40560c89984d129593bb2685cf033a3c458ebb49f349e3a5ee","observation_id":"0d46ae1c-2184-4541-86e7-a7acb2973982","resolution":{"observed_at":"2026-08-05T14:02:06.229836Z","resolver_source":"raw_fallback","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:10.637310Z","title":"HowTo100M: Learning a Text-Video Embedding by Watch- ing Hundred Million Narrated Video Clips","venue":null,"work_id":"b5b24b5a-3f04-46f9-a01a-f26ca5354fe2","year":2019},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.316713Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:e065aac6083b4a50ba2f2a6e881a5f4429b809a9fee8efb1ad998f3405031c18","observation_id":"1618ee6f-091a-4d8e-bfd4-0d1ecedda25f","resolution":{"observed_at":"2026-08-05T14:02:10.680199Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:10.478005Z","title":"Poem: Out-of-distribution detection with pos- terior sampling","venue":null,"work_id":"198c0c11-ca33-4004-8061-0074a1187438","year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.507533Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:fe178222a170965dd3af3f8b58e28ffd022310f731df5cc8c3c8944c931adf16","observation_id":"8af1eb24-6367-4916-be5b-33cfbbe09b3c","resolution":{"observed_at":"2026-08-05T14:02:10.551577Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:10.332285Z","title":"Mo- ments in time dataset: one million videos for event understanding.IEEE transactions on pattern analysis and machine intelligence, 42(2):502–508, 2019","venue":null,"work_id":"f971dadd-9329-4c35-a541-5f0c4c10f57d","year":2019},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.553269Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:eab5b216bc15e5c1659f4c28b2db34cb0db39ae08ac440ccd10ddb9d8b479278","observation_id":"4ab5466a-f3e6-4ce8-9f66-58b1a945f861","resolution":{"observed_at":"2026-08-05T14:02:10.381973Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:03.594472Z","title":"Expanding language-image pretrained models for general video recognition","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.594472Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:b50423d9bcab8e37ee622bcdb1a727ae3ef5d35d71fbe57f585759c2baaa1318","observation_id":"8673dd7b-28cf-4cb7-a980-3e24a9fa7645","resolution":{"observed_at":"2026-08-05T14:02:03.594472Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:10.190137Z","title":"Outlier exposure with confidence control for out-of-distribution detec- tion.Neurocomputing, 441:138–150, 2021","venue":null,"work_id":"5d87baa2-1cb6-4fb7-b18d-10bae442bb9c","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.663952Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:52cb6d6b31099796b8db5fb3553536dca3336e6aef4499d8c62583f68d5bd859","observation_id":"5965cb6c-c003-492c-b3eb-b552385ce969","resolution":{"observed_at":"2026-08-05T14:02:10.250264Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:10.043526Z","title":"A survey on vision-based human action recognition.Image and vision computing, 28(6):976–990, 2010","venue":null,"work_id":"cb8d93bd-b46d-4c61-87f5-3c7d048a654a","year":2010},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.762859Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:9ec07e9d113d7a691e798227d8b4a32080a3c7f3a3b9bf81d6047e016cc8be86","observation_id":"43b0e379-fa27-46aa-946b-b97cb0ae7e7f","resolution":{"observed_at":"2026-08-05T14:02:10.120128Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:09.964165Z","title":"Learning transferable visual models from natural language supervision","venue":null,"work_id":"a2f95382-ad3f-4605-979a-833380704c30","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.856227Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:7d34c6b693836d926c631291ba4ca13fd16e2db3ade77414daba8970e9b85820","observation_id":"0c88274e-2914-4fa7-b1ab-7574ef637269","resolution":{"observed_at":"2026-08-05T14:02:09.995951Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:09.860676Z","title":"Fishr: Invariant gradient variances for out-of-distribution generalization","venue":null,"work_id":"5433444b-08c3-4482-90ba-1374eb51d845","year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:03.957892Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:82267e5a5a32fe5b3d98e7973c2491d18887f650bc08a2543c045400ee25973a","observation_id":"15d0e835-873b-42e7-88cc-876999ec6855","resolution":{"observed_at":"2026-08-05T14:02:09.893998Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:04.067249Z","title":"If deep learning is the answer, what is the question?Nature Reviews Neuroscience, 22(1):55–67, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.067249Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:fc1a0434dcb725fd1165ee75240fbf0b8bd9f715b7f65952d7e338b1dfbe27e6","observation_id":"2a510c09-25be-41ba-ad56-6e5e07bf51cb","resolution":{"observed_at":"2026-08-05T14:02:04.067249Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:09.670821Z","title":"Meta- recognition: The theory and practice of recognition score analysis.IEEE transactions on pattern analysis and machine intelligence, 33(8):1689–1695, 2011","venue":null,"work_id":"f174890c-d49b-40b2-9298-3b47f0e82c50","year":2011},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.147325Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:239eec2c0f9e9847395bc717cde3872c1c6bcdda17d14d863d61bda4e63cc206","observation_id":"7f0b9564-2abd-4f87-ab31-6325e71330ce","resolution":{"observed_at":"2026-08-05T14:02:09.751133Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1212.0402","last_updated":"2012-12-03T14:45:31Z","snapshot_observed_at":"2026-07-06T03:01:10.229407Z","submitted_at":"2012-12-03T14:45:31Z","title":"UCF101: A Dataset of 101 Human Actions Classes From Videos in The Wild","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1212.0402","snapshot_observed_at":"2026-08-05T14:02:04.226399Z","title":"Ucf101: A dataset of 101 human actions classes from videos in the wild","venue":null,"work_id":null,"year":2012},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.226399Z"},"links":{"cited_paper":"/paper/1212.0402","citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:28168ccb79962e2a8b7fad74c1b71ea0fe187fa4e14475b9bb477be0c0615c2d","observation_id":"4d9c788d-165c-433b-af06-00aba146c9e3","resolution":{"observed_at":"2026-08-05T14:02:04.226399Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:09.481172Z","title":"Real-world anomaly detection in surveillance videos","venue":null,"work_id":"7c88b126-054f-4e36-826a-83c876b693da","year":2018},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.324052Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:5200d6d1dfbf1a65f3bf32528d9438976014dce423a91f5b61aa4a3f84193892","observation_id":"c700c04d-6bc3-40c1-b525-fbf7175aa152","resolution":{"observed_at":"2026-08-05T14:02:09.577098Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:09.286183Z","title":"Human action recognition from various data modalities: A review.IEEE transactions on pattern analysis and machine intelligence, 45(3):3200–3225, 2022","venue":null,"work_id":"691a7c3a-e2cf-44a2-b502-5b7d402748e6","year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.402105Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:d152d0d0eb2f01e89d946ffb6cc9fa890df2d36b782b586b3f59f3ceb9af6619","observation_id":"e856d248-1424-47e0-b9f9-f84507e96a64","resolution":{"observed_at":"2026-08-05T14:02:09.378542Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:09.092110Z","title":"Black, Ivan Laptev, and Cordelia Schmid","venue":null,"work_id":"24e8014e-5193-45e9-952e-a031469880b3","year":2017},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.464878Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:77bb5b69285554a87998a59227bcf71fec39b0f22a018187657a14780f6cf834","observation_id":"e71b4814-8007-4ffb-a313-a972afcad7fe","resolution":{"observed_at":"2026-08-05T14:02:09.193093Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:08.956029Z","title":"Open-set recognition: A good closed-set classifier is all you need? 2021","venue":null,"work_id":"5cbce019-aca9-4d72-bf2b-5cf73b59288d","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.557425Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:c5386cd30b29e8257f067302166221201a63c0af17a11690984d63c9edaa2695","observation_id":"4afc6e3f-d632-420d-babe-95261d84a10e","resolution":{"observed_at":"2026-08-05T14:02:08.991793Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:08.745086Z","title":"Generalized category discovery","venue":null,"work_id":"0104727d-1eff-4fdc-991b-d1be0fb00a7f","year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.650860Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:f1e7cf547fe6c4e48f73f3d664811d2f53e988ac625ffbc8f60f9a5458cae50c","observation_id":"bedd9d56-d3cb-436c-90a5-96bb781ca5fc","resolution":{"observed_at":"2026-08-05T14:02:08.881327Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:08.563969Z","title":"No representation rules them all in category discovery","venue":null,"work_id":"c3e57b95-7278-4246-a357-315e327786d1","year":2023},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.718218Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:fee896bdf5475949f78e47bad5df39d8bb07245caf82f4c365e0b28b46bdab7a","observation_id":"fec39636-34ab-453b-9c37-6666dc260ca5","resolution":{"observed_at":"2026-08-05T14:02:08.640276Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:08.404555Z","title":"Self-supervised video representation learning by pace prediction","venue":null,"work_id":"d0271324-488e-4a40-a498-744efa9f2adb","year":2020},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.787888Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:84e78468bb0f160d0ae5a37dfb69a12cb0147a87a8194841b9b243f1c1154fb3","observation_id":"5bac9760-7640-4d5d-8880-f111f096f965","resolution":{"observed_at":"2026-08-05T14:02:08.480574Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:08.190374Z","title":"Action recognition and detection by com- bining motion and appearance features.THUMOS14 Action Recognition Challenge, 1 (2):2, 2014","venue":null,"work_id":"f52e60c6-f241-480a-995b-bf044c073590","year":2014},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.881590Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:7f9ae4cd4e4ea36bd4fd2b6f2a26cdc6e0d5c0c1adc932070c3ab2b1ea5fb953","observation_id":"ac66d5d3-b773-4a9e-b045-3b934958f55b","resolution":{"observed_at":"2026-08-05T14:02:08.289468Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2109.08472","last_updated":"2021-09-17T11:21:34Z","snapshot_observed_at":"2026-07-06T11:48:42.000083Z","submitted_at":"2021-09-17T11:21:34Z","title":"ActionCLIP: A New Paradigm for Video Action Recognition","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2109.08472","snapshot_observed_at":"2026-08-05T14:02:04.935174Z","title":"Actionclip: A new paradigm for video action recognition.arXiv preprint arXiv:2109.08472, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:04.935174Z"},"links":{"cited_paper":"/paper/2109.08472","citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:2afc99a7c5ccf93bbec9b2f89f3e8f300081bd0474f0a0b3e421b520df4f6255","observation_id":"3d76a287-af6c-469c-bde0-b99ccc3cfd3f","resolution":{"observed_at":"2026-08-05T14:02:04.935174Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:07.971277Z","title":"Openood: Benchmark- ing generalized out-of-distribution detection.Advances in Neural Information Pro- cessing Systems, 35:32598–32611, 2022","venue":null,"work_id":"0d082054-aabf-4d75-a230-0f56f2024098","year":2022},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.014884Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:64247de10ccec19a78b3409cae2e21368ddfa4c17cf0efc8548f2669925db8a9","observation_id":"3a16aaa4-b86f-4306-9abd-019f955ab3f0","resolution":{"observed_at":"2026-08-05T14:02:08.072192Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:07.744214Z","title":"Generalized out-of- distribution detection: A survey.International Journal of Computer Vision, pages 1–28, 2024","venue":null,"work_id":"9a9d406b-a09b-4868-9801-8b7605939932","year":2024},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.103596Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:78eeb9d65b6d05895e77618b11846667b4d1e626998e8dbe8af1d7b4ae7f8523","observation_id":"5e5073cc-84dd-4569-b06c-64f73cb7855a","resolution":{"observed_at":"2026-08-05T14:02:07.865968Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:07.548914Z","title":"Un- derstanding deep learning (still) requires rethinking generalization.Communications of the ACM, 64(3):107–115, 2021","venue":null,"work_id":"3d9ca34d-57b4-4088-9722-26ef7488d09c","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.216117Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:7aeba468c619f61a5911ca19eb47e3430e1866c19d53f52ecd1850ed5ca7c845","observation_id":"9b7c7aa8-c1e7-479a-a0ba-930dae3c597b","resolution":{"observed_at":"2026-08-05T14:02:07.645931Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:07.291430Z","title":"Mixture outlier exposure: Towards out-of-distribution detection in fine-grained en- vironments","venue":null,"work_id":"b291d56d-a4d3-48ad-8e50-e7132cef9f8a","year":2023},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.274180Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:1a4ada85c7c0a5e5fe0e6e70ba106e4f5b9f06e0ce3f8e30730bd6db6ec53717","observation_id":"c095afd6-0bd8-45f0-b4fa-af9e02b90dfe","resolution":{"observed_at":"2026-08-05T14:02:07.407597Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:07.138236Z","title":"Open- mix: Reviving known knowledge for discovering novel visual categories in an open world","venue":null,"work_id":"c85f5d75-8862-4ca4-b5db-5638ecaf0de9","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.381623Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:a330448b429f1133657ad5abd08d9acefac9999d71922cfbe99ee5a912a219e9","observation_id":"47a6f305-ac19-41bb-9aaf-b037d21b52b0","resolution":{"observed_at":"2026-08-05T14:02:07.199110Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:07.033945Z","title":"Learning placeholders for open-set recognition","venue":null,"work_id":"72e6f60e-48f2-4303-8512-10c0999117e1","year":2021},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.466295Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:743bb3cf325a09a71dd13a3053194a4fb7deff0d7d4b18d3d7fb4d5e0aa5cf47","observation_id":"c24e56f8-18e8-42a7-9e9e-934521c96303","resolution":{"observed_at":"2026-08-05T14:02:07.062700Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:06.875909Z","title":"Diversified outlier exposure for out-of-distribution detection via in- formative extrapolation.Advances in Neural Information Processing Systems, 36: 22702–22734, 2023","venue":null,"work_id":"abf00742-d2ab-437a-bde6-7da9b9990d50","year":2023},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.542651Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:268d4505be63231a9be572cacf32b10503328e42485b4ea6a2f4947e345b4ed4","observation_id":"aac02e99-d669-4272-bca6-ee4f99b20c7c","resolution":{"observed_at":"2026-08-05T14:02:06.914683Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:06.747330Z","title":"Unleashing mask: Explore the intrinsic out-of-distribution detection capa- bility","venue":null,"work_id":"1cffd789-b024-4ff2-a32d-11168d88da73","year":2023},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.634486Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:4a68a953a6512fcc0b65e1d24b58d3ec0b0633e8527a12fac925312be5e58734","observation_id":"3c238e71-90f7-44f6-bf6e-4e3031c6bbd7","resolution":{"observed_at":"2026-08-05T14:02:06.803921Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:05.763997Z","title":"Towards universal representation for unseen action recognition","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.763997Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:a42d6cf605295b67532113b5d4870da14b8b59aa968b10bfca0426f959965234","observation_id":"33fcc031-32cb-4f0a-afcd-e5f7ede42976","resolution":{"observed_at":"2026-08-05T14:02:05.763997Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T14:02:06.548016Z","title":"Towards open set video anomaly detec- tion","venue":null,"work_id":"1b9515f6-0cb3-4902-9e91-21332b97c59c","year":null},"citing_paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos","version":2},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-05T14:02:05.872722Z"},"links":{"citing_paper":"/paper/2508.21770"},"observation_digest":"sha256:501090cf72d14c80687b574de07c6a7feb06dd4e5e7b98b33136be51d43ce0eb","observation_id":"2722e854-314d-452b-90d6-c81958eb97ce","resolution":{"observed_at":"2026-08-05T14:02:06.642411Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2508.21770","last_updated":"2025-09-08T12:51:08Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-10T10:16:32.475606Z","submitted_at":"2025-08-29T16:43:19Z","title":"What Can We Learn from Harry Potter? An Exploratory Study of Visual Representation Learning from Atypical Videos"},"reference_resolution":{"displayed":59,"state_counts":{"malformed_identifier":0,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":8,"verified_exact":1,"verified_fuzzy":49},"total_outbound_references":59},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 59 of 59 outbound references and 0 inbound Pith citation observations for arXiv:2508.21770."}