{"as_of":"2026-08-10T00:47:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:6eafe9304a3d6f7c71f68e562790968d2b60ff21782434e3c7d6e783f308e470","coverage":[{"denominator":89,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":89,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-09T12:01:50.216585Z","state":"measured"},{"denominator":89,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":89,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2502.02487/citation-record","integrity":"/paper/2502.02487/integrity","json":"/paper/2502.02487/citation-record.json","paper":"/paper/2502.02487"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:49.706189Z","title":"Multiview transformers for video recognition,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.706189Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:30baeed45837cd13bfd48a4911af8843fa66d96d70660b634c6e7795b7ea5618","observation_id":"3caa468a-a115-4f68-9170-ed71df112f25","resolution":{"observed_at":"2026-08-09T12:01:49.706189Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:49.712239Z","title":"Anticipative feature fusion transformer for multi-modal action anticipation,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.712239Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:a1dfd3051a2c8b48c60c812e2cc280f5f2103724b1966c444ea890b345ea5113","observation_id":"9645a5bc-6da9-4528-99e5-c4a556daec64","resolution":{"observed_at":"2026-08-09T12:01:49.712239Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:49.717702Z","title":"Actionformer: Localizing moments of actions with transformers,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.717702Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:288804ecf4283f147cf8233f4a3fc0f5717d26045e69a44e4ae4c0edaacb34ea","observation_id":"f6fef922-79f2-412a-be62-e8fdb11e1b68","resolution":{"observed_at":"2026-08-09T12:01:49.717702Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:49.723287Z","title":"Ubernet: Training a universal convolutional neu- ral network for low-, mid-, and high-level vision using diverse datasets and limited memory,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.723287Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:def9e244e00b243fc39a3437e434198ac94d4059bd4676a746cb52c2ecb5dcf4","observation_id":"dff68709-c8c3-4f6c-a2b8-cdd0a72decb6","resolution":{"observed_at":"2026-08-09T12:01:49.723287Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:49.728350Z","title":"Egocentric video task translation,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.728350Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:74c556a637ab482621593831148282b7cd2ba0894a9066d82474ae671afdcb75","observation_id":"ea43b290-dd35-49f9-9656-cc7dd1cf0c61","resolution":{"observed_at":"2026-08-09T12:01:49.728350Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:49.733541Z","title":"A backpack full of skills: Egocentric video understanding with diverse task perspectives,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.733541Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:a628646f53307d664aad89712f987b82622614beb9111bf5190ca3b621610678","observation_id":"97e488e9-d37f-4bce-a4f5-1741fc6c5f3e","resolution":{"observed_at":"2026-08-09T12:01:49.733541Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:49.739421Z","title":"Test of time: Instilling video-language models with a sense of time,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.739421Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:14e6b2cddd36e033aaf3c1f8e6c176e2e4c4115583e49c90e281acb474ebc0cc","observation_id":"c9329895-a67c-4efa-8573-ba687a5dc66b","resolution":{"observed_at":"2026-08-09T12:01:49.739421Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:49.744269Z","title":"Ego4d: Around the world in 3,000 hours of egocentric video,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.744269Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:44c236ea7eb218a8780a7932c2189afb349704507a0b30a1c45879ae153d8032","observation_id":"be0e6d5a-8e9d-47ae-8b1f-798648619d6d","resolution":{"observed_at":"2026-08-09T12:01:49.744269Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:49.749355Z","title":"The evolution of first person vision methods: A survey,","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.749355Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:89f3b5eb6b76c612b986d967b1a342cf8d07eda409453148141e5edb68e62040","observation_id":"b7c41c5e-04e7-4691-9e58-bbe23448d3d8","resolution":{"observed_at":"2026-08-09T12:01:49.749355Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.811356Z","title":"An outlook into the future of egocentric vision,","venue":null,"work_id":"ac15c591-4ec5-4a52-9196-5ebc130db02e","year":2024},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.754314Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:fc74ed66f433e7ce12d640cb0333eb35b1cd7f8dc3b44e1d24349211c848df58","observation_id":"94fdf060-d883-424f-bbbf-df65aabae76f","resolution":{"observed_at":"2026-08-09T12:01:51.816711Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.785631Z","title":"The epic-kitchens dataset: Collection, challenges and baselines,","venue":null,"work_id":"22e59227-dca4-478e-8e15-76ef8fa84829","year":2021},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.759504Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:26b47e9b5dc96d47a62dc783c94d5ce2a81ae71bf8bf8c69a65d32044b5a5ba7","observation_id":"fad889b9-ed2f-422e-bc58-566b6ba2e100","resolution":{"observed_at":"2026-08-09T12:01:51.790887Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.768594Z","title":"In the eye of the beholder: A survey of models for eyes and gaze,","venue":null,"work_id":"70f8262d-7404-4a78-a9f5-c61e651a7463","year":2009},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.765230Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:ed719a91bc1bce9d1f80fa8649ece9fe065fc0a3c03285ddf16308b02b8258cf","observation_id":"fb704c68-afa9-41ff-8ab0-8f3f9200bbdc","resolution":{"observed_at":"2026-08-09T12:01:51.773722Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.753818Z","title":"Epic-tent: An egocentric video dataset for camping tent assembly,","venue":null,"work_id":"d084e65c-5e3f-4564-9968-fb8ef25a81b3","year":2019},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.770281Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:a31ef974e0b67773370eae894d758e1b4c972bb8c7b4d7681d4109c41a52f49b","observation_id":"0ff8f57e-263c-4ca2-8a31-799bfd92fd0b","resolution":{"observed_at":"2026-08-09T12:01:51.758623Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.738126Z","title":"Rescaling egocentric vision: Collection, pipeline and challenges for epic-kitchens-100,","venue":null,"work_id":"cff43d3f-2dc1-461d-9d51-a0e44dfe6904","year":2022},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.775706Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:d6ce388960971d9886a6e874a3a0796a46502fd67d61ef9af93ea20061fa3b4a","observation_id":"73685f1e-ec80-48f7-b497-e24af85cf386","resolution":{"observed_at":"2026-08-09T12:01:51.743363Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.720852Z","title":"Assembly101: A large-scale multi-view video dataset for understanding procedural activities,","venue":null,"work_id":"7dbe3a5e-1e94-44bf-9a61-582501e3aceb","year":2022},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.780565Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:358ff05fc43cbaca5e7c08d63cd54a97659dd554a4a74f1adf314fcce3ccdbfa","observation_id":"d0d8a1e3-eefa-41c3-850b-ae5ecc6dc363","resolution":{"observed_at":"2026-08-09T12:01:51.726076Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.703279Z","title":"Egocen- tric vision-based action recognition: A survey,","venue":null,"work_id":"3993a018-b345-4845-9d85-8345e7b73117","year":2022},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.785763Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:7a42a2c4d73b2d42b0f4af77ab9a133318acd091ec638802676be59556cdff51","observation_id":"1b5a975e-f8ac-49c3-a56b-274fe4a79821","resolution":{"observed_at":"2026-08-09T12:01:51.709070Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.687246Z","title":"Rolling-unrolling lstms for action anticipation from first-person video,","venue":null,"work_id":"108e7891-f104-4f7a-b05e-69c66693d178","year":2020},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.790342Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:11c77270a515964650823dc6940c288d637f571a9b0afab31e9769bf31882bd3","observation_id":"6c447f57-24d7-4ff8-bc92-2f4f10d6ed56","resolution":{"observed_at":"2026-08-09T12:01:51.692584Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.670356Z","title":"Anticipative video transformer,","venue":null,"work_id":"11e8cfac-8447-42ab-b97b-c7437e68bc32","year":2021},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.795065Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:01a5f7672e9c1feddba703d4628bd229f5349539bb6615c893e3a9c5db9f8c20","observation_id":"7b10e402-c3f0-4bdb-8c10-350bcd3e4328","resolution":{"observed_at":"2026-08-09T12:01:51.675567Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.649228Z","title":"Next- active-object prediction from egocentric videos,","venue":null,"work_id":"88e70db0-b775-40e8-80df-5b0cb6d5001d","year":2017},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.800059Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:53c3fe07dd700d9eaa4547e17f6d4e092d4ba02ac76a53742d5c1ee5d581dc31","observation_id":"2277d730-539e-4f3f-8de4-8f52e6479cc8","resolution":{"observed_at":"2026-08-09T12:01:51.655283Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.629785Z","title":"Improving action segmentation via graph-based temporal reasoning,","venue":null,"work_id":"9ec73f6f-ba7d-4651-b370-16579184062b","year":2020},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.805693Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:fd5234d6d29a7d8d3065d3fe1ba9a3b18cbb2638163cfaa22f1916f60929e3c5","observation_id":"e96d56e5-6d95-400c-8437-be5437e7884b","resolution":{"observed_at":"2026-08-09T12:01:51.635863Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.612157Z","title":"Spotem: Efficient video search for episodic memory,","venue":null,"work_id":"89179548-3af4-4ea2-a868-6039e3417b7c","year":2023},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.810747Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:2bb8755dca8ca473455f566df627836ea670fec513ac32b715874e2765e91c0c","observation_id":"24c265d3-9954-4859-ab72-4793bb7f0054","resolution":{"observed_at":"2026-08-09T12:01:51.617849Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.593981Z","title":"Amego: Active memory from long egocentric videos,","venue":null,"work_id":"db7ec5e4-a7ac-4d0e-86f8-7a7bd91608c4","year":2024},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.817163Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:874fc288762805dcb480fd848d3bf9105a2aee57030c8c14a81d5f41d3466e98","observation_id":"5d1235e1-74d2-4c34-8258-ba6f45ae9ad3","resolution":{"observed_at":"2026-08-09T12:01:51.599722Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.574883Z","title":"Egoschema: A diag- nostic benchmark for very long-form video language understand- ing,","venue":null,"work_id":"34ca5454-ac3b-4366-933d-ae68af2d8724","year":2023},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.822678Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:ba4fd072f476f1135c0a0deb8fdfbf86935632abcd1f15e4eca14e17364e2e31","observation_id":"6c79da88-ce94-4b9f-953b-9b3bdf8145f5","resolution":{"observed_at":"2026-08-09T12:01:51.580832Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.552993Z","title":"Egotaskqa: Understanding human tasks in egocentric videos,","venue":null,"work_id":"f34fab51-a908-4c27-a8b5-a084ef616ced","year":2022},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.829504Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:4b27fc253a245d71f3ee22c9c47a827c4ca38a24294eb9c2a00b43edf4fc4765","observation_id":"40f72a97-5e22-4899-a9ff-b2dd963baa76","resolution":{"observed_at":"2026-08-09T12:01:51.561696Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.533980Z","title":"Multi-modal domain adaptation for fine-grained action recognition,","venue":null,"work_id":"846594e5-16b9-41d5-8270-da2a041732f3","year":2020},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.836085Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:ecf7a5f83a651e657048ecdc7d2872d4fc5daa313af0836121f3467b3c0cc1de","observation_id":"532c44a5-e95a-4f6d-9efd-bb1ef6223ba0","resolution":{"observed_at":"2026-08-09T12:01:51.539011Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.516419Z","title":"Interact before align: Leveraging cross-modal knowledge for domain adaptive action recognition,","venue":null,"work_id":"a6f8643a-7821-4ded-942a-ea6dbf893a2d","year":2022},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.842030Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:a59c879b84e5c7e028312df43e64b1f7227f0a8e4e7d34a5fe9a63901ea40c0d","observation_id":"190aea6f-e7bf-48ee-8b17-3702515805d4","resolution":{"observed_at":"2026-08-09T12:01:51.521749Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.499177Z","title":"Temporal attentive alignment for large-scale video domain adap- tation,","venue":null,"work_id":"14379289-005a-4409-98fc-7cba7eba35e4","year":2019},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.846766Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:d6114d37731c4ba620d4e74609acbde03c66f94fc0532c285a22dc59eb5b4598","observation_id":"ffbc4dfc-b703-4ed6-a16b-e2c77b80bbd9","resolution":{"observed_at":"2026-08-09T12:01:51.504392Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:49.853187Z","title":"What can a cook in italy teach a mechanic in india? action recognition generalisation over scenarios and locations,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.853187Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:896b1f80c92ebd0560f7630cff72849f47119857f5f9a7dba498003cf1200660","observation_id":"343235c2-3095-4750-8702-9e807720f625","resolution":{"observed_at":"2026-08-09T12:01:49.853187Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.470972Z","title":"Relative norm alignment for tackling domain shift in deep multi-modal classification,","venue":null,"work_id":"c79678ab-78f9-4634-80ff-2d44eaf2e60e","year":2024},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.858238Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:3a6d80256cc8a88889d454f925562f7329f93abab2d7134b4ac17a267c927133","observation_id":"c3c37fcf-efda-4357-aa76-51c0492bfdc0","resolution":{"observed_at":"2026-08-09T12:01:51.476089Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.455151Z","title":"Human action recognition from various data modalities: A re- view,","venue":null,"work_id":"429427f3-ff83-4aa1-b8a0-377293e0fc9b","year":2023},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.863243Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:3231128a300d140a55b4b2185c5ae4ed7078c340cf135741f279d1b6b43227fc","observation_id":"d7f0ebe0-cc2d-434e-9f89-bdb26e3db7b2","resolution":{"observed_at":"2026-08-09T12:01:51.460399Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.439423Z","title":"Listen to look: Action recognition by previewing audio,","venue":null,"work_id":"925a8e41-ffea-476f-a077-1aac232c9baa","year":2020},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.868285Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:efbbbff8cbe8415b061193c3321033b5503edb0cb8b12bd226aef710d925a531","observation_id":"1291fc3b-b995-4757-8a63-fc66ae8539bb","resolution":{"observed_at":"2026-08-09T12:01:51.444596Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.422126Z","title":"Egocentric video-language pretraining,","venue":null,"work_id":"364c5417-a9fb-4e36-8e0c-51a0c19b48ef","year":2022},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.873440Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:f8313e511292ac2859ba014dca57b253ba5f8d0bc6d141173d36ac5072a714d9","observation_id":"b82a78f1-4c27-4638-b513-96b4aab4bba1","resolution":{"observed_at":"2026-08-09T12:01:51.428149Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.404777Z","title":"Egovlpv2: Egocentric video- language pre-training with fusion in the backbone,","venue":null,"work_id":"60260cbe-8d36-48b3-9911-b60b2756781a","year":2023},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.878300Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:643e22dc2705c681316b6c356c06577e68c3ca2a45838e7439c7f46477451f08","observation_id":"1aeb0c33-d4b7-4172-aedc-1ca369cf51c5","resolution":{"observed_at":"2026-08-09T12:01:51.410296Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.389210Z","title":"Hiervl: Learning hierarchical video-language embeddings,","venue":null,"work_id":"88fb5d43-ba0e-46fc-a86b-c8ba4eb6d71c","year":2023},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.883565Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:3b58253563aaef9ed449e0750fa224f2f2c4a3e06293a9682dca4a996f01d139","observation_id":"84f7cb13-883f-435b-a2f2-7917253784db","resolution":{"observed_at":"2026-08-09T12:01:51.393908Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.372078Z","title":"Learning video representations from large language models,","venue":null,"work_id":"0ff93f3e-1380-44f6-b058-8ac4a75182d2","year":2023},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.889213Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:aed269b397242c43644cd4577b9f84da41b958a26ee375b767502d7343b171c3","observation_id":"a153fc13-5c44-45da-95a9-a6068e1cfb14","resolution":{"observed_at":"2026-08-09T12:01:51.377294Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.355275Z","title":"A survey of convo- lutional neural networks: analysis, applications, and prospects,","venue":null,"work_id":"4dfd525d-2495-4a0b-babe-b46df835f819","year":2021},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.894501Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:a2a710bfc09dc03c45c2587ab075ac75021405b70be7153ecf1369106003c50d","observation_id":"25453ddc-ab7b-4438-95ab-9b761a79089d","resolution":{"observed_at":"2026-08-09T12:01:51.360337Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.338521Z","title":"A survey of the recent architectures of deep convolutional neural networks,","venue":null,"work_id":"dde9fbed-fcad-4c79-938e-499fd1cb99de","year":2020},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.899776Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:b343d6f7e6be4dbf8466108fee9d88290cc2f9c6f4e68f619cdd7c423f1c8b4b","observation_id":"bd175f98-156c-4379-b9a2-39c742e19e72","resolution":{"observed_at":"2026-08-09T12:01:51.343839Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.322678Z","title":"Recent advances in convolutional neural networks,","venue":null,"work_id":"8ef7f94a-dd38-429f-b4b2-9d34cbc2481d","year":2018},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.905061Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:2bdd8bd5d6a5ba7aa15fc672af12e1ff26febaf0804e86ebe51510951a5d7418","observation_id":"e198deb8-756c-4e7f-95d5-a6e3dc52c761","resolution":{"observed_at":"2026-08-09T12:01:51.327685Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.305428Z","title":"Dynamic edge-conditioned filters in convolutional neural networks on graphs,","venue":null,"work_id":"3ddc6049-f542-4f30-8197-3595d3fe401a","year":2017},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.910064Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:39a80c57e5f6a07f16891e15adef87ab3a02ec8f4aaea48ac1a90bf9d5c29c41","observation_id":"85d9cbec-75f1-4c5f-badc-e92bc28598a8","resolution":{"observed_at":"2026-08-09T12:01:51.310386Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.289114Z","title":"Dynamic graph cnn for learning on point clouds,","venue":null,"work_id":"cbd928d9-2a65-49a2-9796-26a30b037f8c","year":2019},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.915165Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:0f4edc36c199107ad8472ba192ec0c36fedd6d5e57834feeaceb2c698fd93cbd","observation_id":"abb06cf8-8a2a-4d29-8bb8-01dfa77aa45b","resolution":{"observed_at":"2026-08-09T12:01:51.294317Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.272431Z","title":"Graph learning in robotics: a survey,","venue":null,"work_id":"fd36d51c-4da7-4170-ad32-1380c08fb38e","year":2023},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.920755Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:7c00f5e582e806159ef6eb8776dd2c534b5f253e679cacc9f2d7342785e48f72","observation_id":"a1bee137-b291-4876-a738-fc05d247e296","resolution":{"observed_at":"2026-08-09T12:01:51.277669Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.254951Z","title":"Molecular graph convolutions: moving beyond fingerprints,","venue":null,"work_id":"3834b4dc-cfe8-4c94-bd20-89e953860833","year":2016},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.925776Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:cbb964153502958b9ea76688a784285e2fe2aa5aa8d18575814b63a55b3197f5","observation_id":"fcb54564-35df-41c6-8a0c-7fd3dbf067e1","resolution":{"observed_at":"2026-08-09T12:01:51.260764Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.233406Z","title":"Graph neural networks for social recommendation,","venue":null,"work_id":"59ebcff7-3da1-43b7-a9f8-6b3b92b108e7","year":2019},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.931226Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:c55ba727d5219efeb3679614593d6a9578d31426368042fcd570a4dd4145b807","observation_id":"eb36343a-2dd5-4213-9b11-b178093e1b1f","resolution":{"observed_at":"2026-08-09T12:01:51.238535Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.217050Z","title":"Learning to simulate complex physics with graph networks,","venue":null,"work_id":"a1e19576-6200-4294-b89b-e327e585c77d","year":2020},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.937116Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:51343fbaaca3e0346534ed3bf8291bfd3fc5ad3a1feeb9fd0ae500546b18546b","observation_id":"46f885fe-a410-498e-813b-b001804f5cfb","resolution":{"observed_at":"2026-08-09T12:01:51.222236Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.199457Z","title":"Graph convolutional networks for temporal action lo- calization,","venue":null,"work_id":"096b1bcd-e4f4-475f-b97f-998bf83223c7","year":2019},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.942376Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:11d580b71f33a887c7caab1e0a6ae4fb7c820af4ea76e3b649e3285c24a624fb","observation_id":"06829240-ed80-4cc2-86ce-5727fda4c978","resolution":{"observed_at":"2026-08-09T12:01:51.205882Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.183526Z","title":"Stacked spatio- temporal graph convolutional networks for action segmentation,","venue":null,"work_id":"0178300c-b1a1-4a15-994d-40f92dca377d","year":2020},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.947377Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:4d520fa404c050fe0636e1a733cbb2426f317eb16ea7fd46476d6970b376df10","observation_id":"16e8a73f-8ff1-4e47-83ac-5b845aa580ae","resolution":{"observed_at":"2026-08-09T12:01:51.188528Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.165896Z","title":"Action graphs: Weakly- supervised action localization with graph convolution networks,","venue":null,"work_id":"d1950a1a-1562-4538-aa15-6ddc5d40e784","year":2020},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.952478Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:6b044f1d9e648775afba5914e7e530c2090863465d02af51461096aff969c942","observation_id":"d1c24731-673a-4e23-96bf-f6c1a60b3419","resolution":{"observed_at":"2026-08-09T12:01:51.171228Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2008.12432","last_updated":"2020-08-28T01:44:01Z","snapshot_observed_at":"2026-08-08T23:11:53.542523Z","submitted_at":"2020-08-28T01:44:01Z","title":"All About Knowledge Graphs for Actions","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2008.12432","snapshot_observed_at":"2026-08-09T12:01:49.957438Z","title":"All about knowledge graphs for actions,","venue":null,"work_id":null,"year":2008},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.957438Z"},"links":{"cited_paper":"/paper/2008.12432","citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:77c79bc9f3755293ec25fea5e40c197dca502e41bcdbdc8969dd6de727577b18","observation_id":"bd095810-06f4-422a-a383-008c72308a78","resolution":{"observed_at":"2026-08-09T12:01:49.957438Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2006.03201","last_updated":"2020-06-05T02:03:25Z","snapshot_observed_at":"2026-07-06T09:26:14.466190Z","submitted_at":"2020-06-05T02:03:25Z","title":"Egocentric Object Manipulation Graphs","version":1},"cited_work":{"arxiv_id":"2006.03201","doi":null,"metadata_source":"pith","pith_arxiv_id":"2006.03201","snapshot_observed_at":"2026-08-09T12:01:50.425162Z","title":"Egocentric Object Manipulation Graphs","venue":"cs.CV","work_id":"02fc2ab2-922d-49dd-ac39-0a5c1bb63635","year":2020},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.963264Z"},"links":{"cited_paper":"/paper/2006.03201","citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:729170c1934329632dcf50a39fbf012294920d2f0e0a416e4eb626d01b020bb4","observation_id":"fe729620-fe3e-4fa9-8a42-380ed886160f","resolution":{"observed_at":"2026-08-09T12:01:50.432125Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.147073Z","title":"Forecasting action through contact representations from first person video,","venue":null,"work_id":"365f0c0f-c774-46d1-8063-7d25b9bb8e3d","year":2021},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.969265Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:5dd63b724d5d12df8d68de09b4834fffbb323d7faef4b81fe10adb01f6a5c7d0","observation_id":"fa14f356-2abb-44fa-87d0-b8128b2986ff","resolution":{"observed_at":"2026-08-09T12:01:51.154140Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.128171Z","title":"Ego-topo: Environment affordances from egocentric video,","venue":null,"work_id":"e43a2b94-7d0e-4060-8f0e-659d1f33bfa4","year":2020},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.974480Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:d8b48d1e4eb01b6c0e0ba9d52fe168049b0a142c5b3377e0ee4098a18fa45cc6","observation_id":"eac3d443-1261-43af-876e-32f26dca614f","resolution":{"observed_at":"2026-08-09T12:01:51.135266Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.102884Z","title":"Multitask learning,","venue":null,"work_id":"c5a58ce8-b5df-4cbc-b19a-55e91cb7d41e","year":1997},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.979837Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:af9ff5c32fe53ef472cfb5eb5aaef445e1f9310a02f593f0d3df308423137d41","observation_id":"7611d6a4-bdb1-4ad0-a78a-44368fdc3c54","resolution":{"observed_at":"2026-08-09T12:01:51.108312Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.084302Z","title":"A survey on multi-task learning,","venue":null,"work_id":"3c954cee-8e2f-42f8-9b92-79f73fbee390","year":2021},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.984524Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:ee43c966dc6a3ddc1c0d913f6897b68fc797c1627d66d71ff89404b278579784","observation_id":"80214ce6-2cd5-4943-999e-8e5ebe3ad368","resolution":{"observed_at":"2026-08-09T12:01:51.089378Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.063554Z","title":"Video task decathlon: Unifying image and video tasks in autonomous driving,","venue":null,"work_id":"abc64f6a-9cb8-4b1c-94a3-f6ef51669539","year":2023},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.989623Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:6e5e85494f52be0c65557bc494b3c52eb7c4cb7bc143c7d4ec2d6640265ecd7e","observation_id":"9a30971a-dc72-4e43-9182-fdeb3cbe67ea","resolution":{"observed_at":"2026-08-09T12:01:51.069968Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.045345Z","title":"Mutual context network for jointly estimating egocentric gaze and action,","venue":null,"work_id":"897cf774-d1bf-4788-95e3-0789d43bf90f","year":2020},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:49.994522Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:2f8c188326451277fb0b10a8730262172b5a18ba8339a82b53e246afcc7d23ad","observation_id":"b7fe998e-abbf-4c78-81a9-d6c0dc845665","resolution":{"observed_at":"2026-08-09T12:01:51.051360Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.020853Z","title":"Efficiently identifying task groupings for multi-task learning,","venue":null,"work_id":"9f2b5a29-0876-4c5e-8c35-99955119dee6","year":2021},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.003877Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:d9e6a9949139b03c4d356c21a77d496bf1a434dfd928526acba5252203d68ace","observation_id":"60739f7b-3079-4589-86ce-44e402a42197","resolution":{"observed_at":"2026-08-09T12:01:51.030665Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:51.001330Z","title":"A unified sequence interface for vision tasks,","venue":null,"work_id":"cd3e58cb-804a-43ba-8552-0203e59c5811","year":2022},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.008996Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:0f75ce0da3854add38f661b504918fad4b0dd29ee82760c4056b22b9d5c18e0d","observation_id":"1ef3d6db-950a-4001-beda-03ecd68430ae","resolution":{"observed_at":"2026-08-09T12:01:51.007232Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.984089Z","title":"Adamv-moe: Adaptive multi-task vision mixture-of- experts,","venue":null,"work_id":"e1714cc5-74da-4ade-8613-169813283f38","year":2023},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.013743Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:a516a6b5dcc1570cac196ff9d816efb9d074c5faae8168587c66b1aa1ea7d8bb","observation_id":"20a29420-20b1-48f5-907c-1afabbd8fbdd","resolution":{"observed_at":"2026-08-09T12:01:50.989817Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.965837Z","title":"Deep multitask learning with progressive parameter sharing,","venue":null,"work_id":"cd675369-f451-4d7b-a7b7-0b1bb4512c3e","year":2023},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.018555Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:2684296682616b1b477d65b4bce7139b0c9a88f66aed84b356ab19e47e322342","observation_id":"6d41b72b-9453-4faa-8127-afaae57a5090","resolution":{"observed_at":"2026-08-09T12:01:50.971263Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.950461Z","title":"Unihcp: A unified model for human- centric perceptions,","venue":null,"work_id":"247a580a-eb43-4b7b-a466-abf17c51bf64","year":2023},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.028777Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:ffa0e316ea959bee43e26e24a584f3e53cc6c48f9e4a58118cdaf4fc2113eb65","observation_id":"b481df6c-721d-415f-ac1c-9afcd2c08440","resolution":{"observed_at":"2026-08-09T12:01:50.955488Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.932621Z","title":"Learning with whom to share in multi-task feature learning,","venue":null,"work_id":"dd214499-3bba-4598-963e-8e44185dca06","year":2011},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.033983Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:dd6145c707e4b9181f7b6c5e761838bd1b8804f5df2567cc7d2543fde3a6451a","observation_id":"9b37a0eb-4d79-421a-a3e6-75449f432d2e","resolution":{"observed_at":"2026-08-09T12:01:50.938399Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.914151Z","title":"Learning to branch for multi- task learning,","venue":null,"work_id":"3672d245-ed61-4c00-b2e8-7dd6996e1814","year":2020},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.039397Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:50bdf44e10f45404e5cde9066719f7524252bdc11c1e49443fd2fd3a0f223be3","observation_id":"a359dd15-e73e-483b-a470-e36430acd828","resolution":{"observed_at":"2026-08-09T12:01:50.922029Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.896647Z","title":"Which tasks should be learned together in multi-task learning?","venue":null,"work_id":"fc60beae-fa0f-4399-aa49-30a855eb97b2","year":2020},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.044824Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:84e574b0b4485f33fe1b2b51fe27ded234f388c6d9ded5f7673d37fae9e57ca0","observation_id":"b20e91b4-e4e2-4be3-98c1-cfd65dcdc293","resolution":{"observed_at":"2026-08-09T12:01:50.902446Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.878220Z","title":"Adashare: Learning what to share for efficient deep multi-task learning,","venue":null,"work_id":"1829c6b6-9140-49f9-9464-1a16b4095d68","year":2020},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.054304Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:231e4f7a0f488f5215c060a42f93aa54106e637a1a7eaa0fbaffc0e5e2e0c5e4","observation_id":"472b1db2-730f-45d9-a588-8c511ce16f35","resolution":{"observed_at":"2026-08-09T12:01:50.884676Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.861959Z","title":"Multitask learning to improve egocentric action recognition,","venue":null,"work_id":"eb82e48f-aa6b-44fd-b1d3-3521d7a42bc7","year":2019},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.060163Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:f120a0ec07b2a96727bf61d5054933f59942d1669d42597a7ebf03411076d75d","observation_id":"017e116f-bf9c-4b27-b4d0-231fb67e6e82","resolution":{"observed_at":"2026-08-09T12:01:50.867346Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.845504Z","title":"Interactive prototype learning for egocentric action recognition,","venue":null,"work_id":"c6bacacf-6642-4497-96b6-75ec51198683","year":2021},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.065273Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:245da5346881fa5daf387286d2debdbc67b0a653b90f11338bcb1f2671ab4e54","observation_id":"1f789a35-1dc5-43a3-83f1-3658e37dd0a0","resolution":{"observed_at":"2026-08-09T12:01:50.851004Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.828541Z","title":"Multi-task learning using uncertainty to weigh losses for scene geometry and semantics,","venue":null,"work_id":"78734f70-d001-4a13-adbe-f8aa358efa68","year":2018},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.071595Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:568235e2f9ed7c08c004a99d7bf1e64969a6b5b0bf58cf9bcef37a60bd623173","observation_id":"2e525a13-7b24-402a-b0b3-08377f807ec2","resolution":{"observed_at":"2026-08-09T12:01:50.834020Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.811052Z","title":"Grad- norm: Gradient normalization for adaptive loss balancing in deep multitask networks,","venue":null,"work_id":"a2088343-8b96-4597-ac88-b09871e9aaf1","year":2018},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.076425Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:268dd4aa7ee312c75c20aef6e750214de0a0b8fb9f511b384f05a9a9bdd064a4","observation_id":"278799a8-327f-4e7c-aca0-3ea7870681bb","resolution":{"observed_at":"2026-08-09T12:01:50.816635Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1806.08028","last_updated":"2018-06-21T00:54:07Z","snapshot_observed_at":"2026-07-06T06:46:01.245927Z","submitted_at":"2018-06-21T00:54:07Z","title":"Gradient Adversarial Training of Neural Networks","version":1},"cited_work":{"arxiv_id":"1806.08028","doi":null,"metadata_source":"pith","pith_arxiv_id":"1806.08028","snapshot_observed_at":"2026-08-09T12:01:50.370014Z","title":"Gradient Adversarial Training of Neural Networks","venue":"cs.LG","work_id":"f38a7464-7949-4d29-a501-829e49a90406","year":2018},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.081280Z"},"links":{"cited_paper":"/paper/1806.08028","citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:5ac1c200df762dffca9dda5d38191168d271b5d6a4f83b977eac0afa83ae737d","observation_id":"c22936a9-c624-4bc0-abf6-a6c14a1cf32d","resolution":{"observed_at":"2026-08-09T12:01:50.386648Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.792488Z","title":"Dy- namic task prioritization for multitask learning,","venue":null,"work_id":"496890ac-adac-4b3b-ab76-6765ce2ea295","year":2018},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.086523Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:10bfff9b3e8d859a54074b9876cfdb90480b39949979429ce43343f3acf23795","observation_id":"d742a745-59de-46d2-9384-3899589b283e","resolution":{"observed_at":"2026-08-09T12:01:50.798917Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.776442Z","title":"Gradient surgery for multi-task learning,","venue":null,"work_id":"9863dadd-ae7d-4742-8c4b-ba9ed55d864e","year":2020},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.091504Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:cdabaf45c1b196f7d7179ef9f57a1be3b96d1f01146989f87cb999498ff646b3","observation_id":"b2a94dee-325a-4a04-997b-a9307c1af404","resolution":{"observed_at":"2026-08-09T12:01:50.781728Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.755300Z","title":"Mti-net: Multi- scale task interaction networks for multi-task learning,","venue":null,"work_id":"208aad8f-9d27-40b4-9b98-da4f73ce39e7","year":2020},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.096501Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:9fb99a40b7c88defed907089aa840e50216b8066cb041bb12ec2842f32a08710","observation_id":"fdda7ede-12e9-4625-bf80-80bfc7b86a6e","resolution":{"observed_at":"2026-08-09T12:01:50.761425Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1706.05098","last_updated":"2017-06-15T21:38:12Z","snapshot_observed_at":"2026-08-04T14:30:44.904839Z","submitted_at":"2017-06-15T21:38:12Z","title":"An Overview of Multi-Task Learning in Deep Neural Networks","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1706.05098","snapshot_observed_at":"2026-08-09T12:01:50.104093Z","title":"An overview of multi-task learning in deep neural networks,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.104093Z"},"links":{"cited_paper":"/paper/1706.05098","citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:debc887859f5bc0bc027005c4088e1f1147cb732d034cada6276299d65efd0bd","observation_id":"a3688891-c277-42a2-9ed6-c0aa61654ecd","resolution":{"observed_at":"2026-08-09T12:01:50.104093Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.736997Z","title":"Video self-stitching graph network for temporal action localization,","venue":null,"work_id":"a57aafb7-48d7-4264-bcaa-05d7d5442e87","year":2021},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.110327Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:6e59b4cfbb252cf9e80f8c8515b492a6c81d98c3927c1f370648511851ddebeb","observation_id":"c9598c8b-b24b-4651-b27e-ff8e761196c6","resolution":{"observed_at":"2026-08-09T12:01:50.742683Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.717863Z","title":"Omnivore: A single model for many visual modalities,","venue":null,"work_id":"bb498a58-8141-4791-a614-fdfbf9655818","year":2022},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.116893Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:d9aeb7cebc423ce92d0ddb7b24e7ab104d62bf46b39b9f2f4440d233e40bdf4b","observation_id":"5eecfdf4-211b-4b95-8e5d-c2320f9a8115","resolution":{"observed_at":"2026-08-09T12:01:50.724250Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.02025","last_updated":"2023-07-05T05:23:49Z","snapshot_observed_at":"2026-07-06T15:50:22.337115Z","submitted_at":"2023-07-05T05:23:49Z","title":"NMS Threshold matters for Ego4D Moment Queries -- 2nd place solution to the Ego4D Moment Queries Challenge 2023","version":1},"cited_work":{"arxiv_id":"2307.02025","doi":null,"metadata_source":"pith","pith_arxiv_id":"2307.02025","snapshot_observed_at":"2026-08-09T12:01:50.299813Z","title":"NMS Threshold matters for Ego4D Moment Queries -- 2nd place solution to the Ego4D Moment Queries Challenge 2023","venue":"cs.CV","work_id":"15d675ff-f73a-4713-b9a4-24a8176609f3","year":2023},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.129383Z"},"links":{"cited_paper":"/paper/2307.02025","citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:47999891ffa0c42d8b5991ebf0fe03d81369936af89bfd8f2492b5b50a2a1716","observation_id":"b6af6acd-0164-4777-986b-a94e1a0c6502","resolution":{"observed_at":"2026-08-09T12:01:50.310756Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.679614Z","title":"Focal loss for dense object detection,","venue":null,"work_id":"0cfef711-8df6-471f-9af9-4b8acdc29171","year":2017},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.135670Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:361bca89dbf20e622d67e568bc00507d4d8494bda86b46efd9d630eb63b22bd3","observation_id":"df14ecf5-7372-4274-ae9a-76be28be5e34","resolution":{"observed_at":"2026-08-09T12:01:50.685476Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.661435Z","title":"Distance-iou loss: Faster and better learning for bounding box regression,","venue":null,"work_id":"31f067d7-4fac-45f8-9a22-d1f165c002ef","year":2020},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.141355Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:1f9764c67a30950cafbbc6dcf06a7fb2de084c3b74bca9a7251a9917522b8201","observation_id":"bc961733-06a6-449e-a432-9a79e435103e","resolution":{"observed_at":"2026-08-09T12:01:50.667285Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.147886Z","title":"Slowfast networks for video recognition,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.147886Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:78154b15a76ceec49829fe14d694443d8da2d0838d157e30754c6900ff5009e6","observation_id":"d6746634-6ef6-45ef-91b8-0725b385697d","resolution":{"observed_at":"2026-08-09T12:01:50.147886Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.697298Z","title":"Quo vadis, action recognition? a new model and the kinetics dataset,","venue":null,"work_id":"37739ad2-7205-480f-ba28-7627e1986ce5","year":2017},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.153568Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:ba2196880840ff6c8c20a0c57bfc6bded84e0a743c848cf04fa5bde80da0630b","observation_id":"996e68e5-a2f5-4222-8de8-9f502293c610","resolution":{"observed_at":"2026-08-09T12:01:50.704073Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.167322Z","title":"Semi-supervised classification with graph convolutional networks,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.167322Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:9e9350b3a299b3c21a2a7a05937d465fd39c12b42886c908ef123fefc9fdb8e9","observation_id":"671a4d75-0f68-46ac-a511-3a52177deeba","resolution":{"observed_at":"2026-08-09T12:01:50.167322Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.614168Z","title":"Graph attention networks,","venue":null,"work_id":"4c17349e-3acd-40cd-8667-f5842ad1a2b0","year":2018},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.173369Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:9fc0fcf654ca95111c590df58de2d6143d57b9f40d898dd125dfd25f673e8d36","observation_id":"2620d383-1f50-4acd-a356-4fec5353e6d6","resolution":{"observed_at":"2026-08-09T12:01:50.621159Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.594383Z","title":"Inductive representation learning on large graphs,","venue":null,"work_id":"bcacdf51-82d8-4140-b669-c160695c37fc","year":2017},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.179072Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:af9c985ea4f9be4f1d7ebeeb243f906fbdb1423dc6fd355ca8c1f10a93a023f2","observation_id":"81f0712f-4459-4bfe-86b0-805b8aa46f8a","resolution":{"observed_at":"2026-08-09T12:01:50.600454Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.567665Z","title":"Signed graph convolutional net- works,","venue":null,"work_id":"dcdce262-01e8-49af-9cd2-87a49a802522","year":2018},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.185825Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:a1342408a669fbe93117bed4c091990dde4b3e1e108c483b00d2e7953782aa6c","observation_id":"10b99aaf-fdb7-402f-b470-9d009f50d247","resolution":{"observed_at":"2026-08-09T12:01:50.575088Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.543821Z","title":"Action sensitivity learning for temporal action localization,","venue":null,"work_id":"a8ea9912-4833-43fe-bb9f-2a0eb13359d0","year":2023},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.192193Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:b3f0494f08ae5b249e0f224944546ef0eeed0097834128da8981794a1972e655","observation_id":"e25abf86-9944-4d4e-a85a-de524051b8c2","resolution":{"observed_at":"2026-08-09T12:01:50.550186Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.09172","last_updated":"2023-09-25T12:11:43Z","snapshot_observed_at":"2026-07-06T15:43:00.250504Z","submitted_at":"2023-06-15T14:50:17Z","title":"Action Sensitivity Learning for the Ego4D Episodic Memory Challenge 2023","version":2},"cited_work":{"arxiv_id":"2306.09172","doi":null,"metadata_source":"pith","pith_arxiv_id":"2306.09172","snapshot_observed_at":"2026-08-09T12:01:50.266965Z","title":"Action Sensitivity Learning for the Ego4D Episodic Memory Challenge 2023","venue":"cs.CV","work_id":"57e58b85-56dd-425c-9714-74f8a8aeae19","year":2023},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":87,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.198165Z"},"links":{"cited_paper":"/paper/2306.09172","citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:1a33e6acc5e97833d5f4a1c03eb381fdc700e6d92deb8e2ef1e85bd37a64294f","observation_id":"e6b6011b-9c3a-465b-836c-c0a5c0f4c1f4","resolution":{"observed_at":"2026-08-09T12:01:50.277472Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.519532Z","title":"Intention-conditioned long- term human egocentric action anticipation,","venue":null,"work_id":"e74d522e-ce7d-45c5-9845-f32c0c565149","year":2023},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":88,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.204840Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:e96cf83c5b691af44c70af7795cb5650642e6d25bd591909177716f3cb5627c7","observation_id":"f18fadf7-050f-4448-97c0-6f8ed0600c0e","resolution":{"observed_at":"2026-08-09T12:01:50.526257Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.498017Z","title":"Antgpt: Can large language models help long-term action anticipation from videos?","venue":null,"work_id":"60889086-0122-498c-b95f-8dfe83f235b7","year":2024},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":89,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.210111Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:382eb62ee31bf4ba85bf6f6fa0e1d81a9b6164c7b1ea9718a04286b8d65a7dc5","observation_id":"9a8604a6-e7f8-42e5-b8fd-80dc4d4a0398","resolution":{"observed_at":"2026-08-09T12:01:50.504623Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T12:01:50.478082Z","title":"Palm: Predicting actions through language models,","venue":null,"work_id":"3e7a2ef0-d562-424a-b37c-19372edf6c97","year":2024},"citing_paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives","version":1},"reference_index":90,"source":"pdf_text","source_observed_at":"2026-08-09T12:01:50.216585Z"},"links":{"citing_paper":"/paper/2502.02487"},"observation_digest":"sha256:f7f6775af7fceea77e0c356e8104c4d7a342f8cbd63d1961a3a99ed2bb0054c7","observation_id":"5ede94de-8cf9-4360-a926-349b6dffa1a4","resolution":{"observed_at":"2026-08-09T12:01:50.484823Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2502.02487","last_updated":"2025-02-04T17:03:49Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-09T11:55:57.133701Z","submitted_at":"2025-02-04T17:03:49Z","title":"Hier-EgoPack: Hierarchical Egocentric Video Understanding with Diverse Task Perspectives"},"reference_resolution":{"displayed":89,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":14,"verified_exact":4,"verified_fuzzy":71},"total_outbound_references":89},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 89 of 89 outbound references and 0 inbound Pith citation observations for arXiv:2502.02487."}