{"as_of":"2026-08-10T03:22:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:0a36aeb623e473929062c800edf011df9388498d0b754ea2d0f58af5ce59276c","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":8,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":8,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":8,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":8,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-09T00:23:08.106301Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T06:39:37.586156Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2403.01422","last_updated":"2025-08-11T12:47:49Z","snapshot_observed_at":"2026-08-09T20:54:23.915134Z","submitted_at":"2024-03-03T07:43:39Z","title":"DreamFrame: Enhancing Video Understanding via Automatically Generated QA and Style-Consistent Keyframes","version":5},"cited_work":{"arxiv_id":"2403.01422","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.01422","snapshot_observed_at":"2026-07-04T06:39:37.586156Z","title":"Moviellm: Enhancing long video understanding with ai-generated movies","venue":null,"work_id":"e83e5156-4bb5-40b3-82ad-ab54d0a66dac","year":2024},"citing_paper":{"arxiv_id":"2406.04264","last_updated":"2025-01-01T15:53:58Z","snapshot_observed_at":"2026-08-03T20:38:36.602554Z","submitted_at":"2024-06-06T17:09:32Z","title":"MLVU: Benchmarking Multi-task Long Video Understanding","version":3},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-05-14T19:55:26.333923Z"},"links":{"cited_paper":"/paper/2403.01422","citing_paper":"/paper/2406.04264"},"observation_digest":"sha256:70b97600f479713cacfa95e812bc96228879d37ba8b0daf27f084cc56b4305c7","observation_id":"eec8055e-32e7-4197-bc63-f3613be186f8","resolution":{"observed_at":"2026-05-14T19:55:26.445584Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.01422","last_updated":"2025-08-11T12:47:49Z","snapshot_observed_at":"2026-08-09T20:54:23.915134Z","submitted_at":"2024-03-03T07:43:39Z","title":"DreamFrame: Enhancing Video Understanding via Automatically Generated QA and Style-Consistent Keyframes","version":5},"cited_work":{"arxiv_id":"2403.01422","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.01422","snapshot_observed_at":"2026-07-04T06:39:37.586156Z","title":"Moviellm: Enhancing long video understanding with ai-generated movies","venue":null,"work_id":"e83e5156-4bb5-40b3-82ad-ab54d0a66dac","year":2024},"citing_paper":{"arxiv_id":"2501.13106","last_updated":"2025-06-03T03:33:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-22T18:59:46Z","title":"VideoLLaMA 3: Frontier Multimodal Foundation Models for Image and Video Understanding","version":4},"reference_index":149,"source":"pdf_text","source_observed_at":"2026-05-11T01:19:59.603343Z"},"links":{"cited_paper":"/paper/2403.01422","citing_paper":"/paper/2501.13106"},"observation_digest":"sha256:d9200c1e0ff2be855e493844ac3030ca0398e550a9c2f80ec47bac1e8040d6a5","observation_id":"a3bf9cf7-8906-4a9c-acc7-b93a7426b70e","resolution":{"observed_at":"2026-05-11T01:20:00.076582Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.01422","last_updated":"2025-08-11T12:47:49Z","snapshot_observed_at":"2026-08-09T20:54:23.915134Z","submitted_at":"2024-03-03T07:43:39Z","title":"DreamFrame: Enhancing Video Understanding via Automatically Generated QA and Style-Consistent Keyframes","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.01422","snapshot_observed_at":"2026-08-09T00:23:08.106301Z","title":"Moviellm: Enhancing long video understanding with ai-generated movies,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.03897","last_updated":"2025-07-07T10:40:56Z","snapshot_observed_at":"2026-08-09T00:16:57.664764Z","submitted_at":"2025-02-06T09:18:30Z","title":"UniForm: A Unified Multi-Task Diffusion Transformer for Audio-Video Generation","version":5},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-09T00:23:08.106301Z"},"links":{"cited_paper":"/paper/2403.01422","citing_paper":"/paper/2502.03897"},"observation_digest":"sha256:2f14cf6a2913905b536a6c49c8bf0dbe6b278391c78a4d959bda7025c81fdd70","observation_id":"3f50a790-a285-458a-a68b-6cf1940de7c4","resolution":{"observed_at":"2026-08-09T00:23:08.106301Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.01422","last_updated":"2025-08-11T12:47:49Z","snapshot_observed_at":"2026-08-09T20:54:23.915134Z","submitted_at":"2024-03-03T07:43:39Z","title":"DreamFrame: Enhancing Video Understanding via Automatically Generated QA and Style-Consistent Keyframes","version":5},"cited_work":{"arxiv_id":"2403.01422","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.01422","snapshot_observed_at":"2026-07-04T06:39:37.586156Z","title":"Moviellm: Enhancing long video understanding with ai-generated movies","venue":null,"work_id":"e83e5156-4bb5-40b3-82ad-ab54d0a66dac","year":2024},"citing_paper":{"arxiv_id":"2604.11283","last_updated":"2026-06-01T08:50:38Z","snapshot_observed_at":"2026-08-02T05:26:21.584719Z","submitted_at":"2026-04-13T10:42:31Z","title":"Multimodal Large Language Model-Enabled Video Translation: A Role-Oriented Survey","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-10T16:36:33.264166Z"},"links":{"cited_paper":"/paper/2403.01422","citing_paper":"/paper/2604.11283"},"observation_digest":"sha256:53cf0eb23c26ebb81b64c804b2d273d58b92058594cf1de4f8921c40d3ac6f43","observation_id":"2abf08f3-9364-4881-ab45-3467b07caf6a","resolution":{"observed_at":"2026-05-11T08:30:57.084788Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.01422","last_updated":"2025-08-11T12:47:49Z","snapshot_observed_at":"2026-08-09T20:54:23.915134Z","submitted_at":"2024-03-03T07:43:39Z","title":"DreamFrame: Enhancing Video Understanding via Automatically Generated QA and Style-Consistent Keyframes","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.01422","snapshot_observed_at":"2026-07-12T22:04:31.302192Z","title":"Moviellm: Enhancing long video understanding with ai-generated movies,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2604.11283","last_updated":"2026-06-01T08:50:38Z","snapshot_observed_at":"2026-08-02T05:26:21.584719Z","submitted_at":"2026-04-13T10:42:31Z","title":"Multimodal Large Language Model-Enabled Video Translation: A Role-Oriented Survey","version":2},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-07-12T22:04:31.302192Z"},"links":{"cited_paper":"/paper/2403.01422","citing_paper":"/paper/2604.11283"},"observation_digest":"sha256:5febe6ff83ea78d869c59c13b18bd1e7a51da1a15e48be9160dcc41269a6a816","observation_id":"5dde9bc3-9171-4c61-bc3a-f7ffe8c86a11","resolution":{"observed_at":"2026-07-12T22:04:31.302192Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.01422","last_updated":"2025-08-11T12:47:49Z","snapshot_observed_at":"2026-08-09T20:54:23.915134Z","submitted_at":"2024-03-03T07:43:39Z","title":"DreamFrame: Enhancing Video Understanding via Automatically Generated QA and Style-Consistent Keyframes","version":5},"cited_work":{"arxiv_id":"2403.01422","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.01422","snapshot_observed_at":"2026-07-04T06:39:37.586156Z","title":"Moviellm: Enhancing long video understanding with ai-generated movies","venue":null,"work_id":"e83e5156-4bb5-40b3-82ad-ab54d0a66dac","year":2024},"citing_paper":{"arxiv_id":"2605.31069","last_updated":"2026-05-29T09:38:43Z","snapshot_observed_at":"2026-08-07T07:32:37.435310Z","submitted_at":"2026-05-29T09:38:43Z","title":"Towards Effective Long-Video Event Prediction via Multi-Level Event Semantics Mining","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-06-28T23:16:42.001361Z"},"links":{"cited_paper":"/paper/2403.01422","citing_paper":"/paper/2605.31069"},"observation_digest":"sha256:72712b788430dd9c4d1d199f94432e009d20cda057627066fa0eee8b5cc4033a","observation_id":"ce518f8b-60c1-4ac8-9543-4306ad6383f2","resolution":{"observed_at":"2026-06-29T00:02:50.348153Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.01422","last_updated":"2025-08-11T12:47:49Z","snapshot_observed_at":"2026-08-09T20:54:23.915134Z","submitted_at":"2024-03-03T07:43:39Z","title":"DreamFrame: Enhancing Video Understanding via Automatically Generated QA and Style-Consistent Keyframes","version":5},"cited_work":{"arxiv_id":"2403.01422","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.01422","snapshot_observed_at":"2026-07-04T06:39:37.586156Z","title":"Moviellm: Enhancing long video understanding with ai-generated movies","venue":null,"work_id":"e83e5156-4bb5-40b3-82ad-ab54d0a66dac","year":2024},"citing_paper":{"arxiv_id":"2606.21734","last_updated":"2026-06-19T20:43:49Z","snapshot_observed_at":"2026-08-05T18:05:51.515234Z","submitted_at":"2026-06-19T20:43:49Z","title":"HPP: Hierarchical Programmatic Probing for Long Video Understanding by Decoupling Perception and Reasoning","version":1},"reference_index":152,"source":"arxiv_source","source_observed_at":"2026-06-26T14:19:53.450263Z"},"links":{"cited_paper":"/paper/2403.01422","citing_paper":"/paper/2606.21734"},"observation_digest":"sha256:d9eb88e0749bd256b24f4a93bd891d4754ac7762d4226a00288f302c4496985a","observation_id":"7c586a91-264b-4629-b4b8-0e6c2741cbfc","resolution":{"observed_at":"2026-07-04T06:39:37.587883Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.01422","last_updated":"2025-08-11T12:47:49Z","snapshot_observed_at":"2026-08-09T20:54:23.915134Z","submitted_at":"2024-03-03T07:43:39Z","title":"DreamFrame: Enhancing Video Understanding via Automatically Generated QA and Style-Consistent Keyframes","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.01422","snapshot_observed_at":"2026-08-01T10:59:17.682021Z","title":"Moviellm: Enhancing long video understanding with ai-generated movies.arXiv preprint arXiv:2403.01422, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.20057","last_updated":"2026-07-22T11:59:38Z","snapshot_observed_at":"2026-08-05T01:59:04.033220Z","submitted_at":"2026-07-22T11:59:38Z","title":"Antigen-specific Antibody Multi-modal Foundation Model for Functional Antibody Design","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-01T10:59:17.682021Z"},"links":{"cited_paper":"/paper/2403.01422","citing_paper":"/paper/2607.20057"},"observation_digest":"sha256:2a277166facad645abb295ffcc71057822c44e84ab144d560f8482cfd69c555d","observation_id":"0dcbb900-7f09-48eb-9572-e28596bcc381","resolution":{"observed_at":"2026-08-01T10:59:17.682021Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2403.01422/citation-record","integrity":"/paper/2403.01422/integrity","json":"/paper/2403.01422/citation-record.json","paper":"/paper/2403.01422"},"outbound":[],"paper":{"arxiv_id":"2403.01422","last_updated":"2025-08-11T12:47:49Z","latest_version":5,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-09T20:54:23.915134Z","submitted_at":"2024-03-03T07:43:39Z","title":"DreamFrame: Enhancing Video Understanding via Automatically Generated QA and Style-Consistent Keyframes"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 8 inbound Pith citation observations for arXiv:2403.01422."}