{"as_of":"2026-08-10T11:28:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:2501dab56ee88f6fa72f1421eb82a27cebe70e071513c07ffacf435d188f55b6","coverage":[{"denominator":12,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":12,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T16:24:41.740062Z","state":"measured"},{"denominator":13,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":13,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-10T17:45:51.528645Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-11T06:11:01.758109Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2507.13609","last_updated":"2025-07-18T02:29:19Z","snapshot_observed_at":"2026-08-07T13:58:45.437787Z","submitted_at":"2025-07-18T02:29:19Z","title":"CoTasks: Chain-of-Thought based Video Instruction Tuning Tasks","version":1},"cited_work":{"arxiv_id":"2507.13609","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.13609","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"arXiv preprint arXiv:2507.13609 (2025)","venue":null,"work_id":"59a719db-ff4a-4eb0-9fe8-e7b1bba5ce03","year":2025},"citing_paper":{"arxiv_id":"2604.08209","last_updated":"2026-04-09T13:09:40Z","snapshot_observed_at":"2026-07-06T22:57:22.475588Z","submitted_at":"2026-04-09T13:09:40Z","title":"OmniJigsaw: Enhancing Omni-Modal Reasoning via Modality-Orchestrated Reordering","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-10T17:45:51.528645Z"},"links":{"cited_paper":"/paper/2507.13609","citing_paper":"/paper/2604.08209"},"observation_digest":"sha256:0ad54dfc4c7a913bfa1622abfd4bd519b0a23ee83ddb8a7085e394af968224ce","observation_id":"a6afcdd0-8530-4bb8-864b-d8f053310c6c","resolution":{"observed_at":"2026-05-11T06:11:01.761451Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2507.13609/citation-record","integrity":"/paper/2507.13609/integrity","json":"/paper/2507.13609/citation-record.json","paper":"/paper/2507.13609"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T16:24:42.244344Z","title":"Ming Nie, Dan Ding, Chunwei Wang, Yuanfan Guo, Jianhua Han, Hang Xu, and Li Zhang","venue":null,"work_id":"5fce81ff-971d-4b27-b099-4c683a5dbf2d","year":2025},"citing_paper":{"arxiv_id":"2507.13609","last_updated":"2025-07-18T02:29:19Z","snapshot_observed_at":"2026-08-07T13:58:45.437787Z","submitted_at":"2025-07-18T02:29:19Z","title":"CoTasks: Chain-of-Thought based Video Instruction Tuning Tasks","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T16:24:41.412118Z"},"links":{"citing_paper":"/paper/2507.13609"},"observation_digest":"sha256:6df7f3a8d661adb5caca07366864c469e9571ba64af38871aa348d44436915cb","observation_id":"37732643-e2fc-40a0-99dc-eb54753a8084","resolution":{"observed_at":"2026-08-06T16:24:42.316157Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.08195","last_updated":"2024-03-13T02:41:39Z","snapshot_observed_at":"2026-07-06T17:43:44.153710Z","submitted_at":"2024-03-13T02:41:39Z","title":"Efficiently verifiable quantum advantage on near-term analog quantum simulators","version":1},"cited_work":{"arxiv_id":"2403.08195","doi":null,"metadata_source":"pith","pith_arxiv_id":"2403.08195","snapshot_observed_at":"2026-08-06T16:24:42.020745Z","title":"Efficiently verifiable quantum advantage on near-term analog quantum simulators","venue":"quant-ph","work_id":"09605899-0a34-4225-a65a-af93c54d6357","year":2024},"citing_paper":{"arxiv_id":"2507.13609","last_updated":"2025-07-18T02:29:19Z","snapshot_observed_at":"2026-08-07T13:58:45.437787Z","submitted_at":"2025-07-18T02:29:19Z","title":"CoTasks: Chain-of-Thought based Video Instruction Tuning Tasks","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T16:24:41.228235Z"},"links":{"cited_paper":"/paper/2403.08195","citing_paper":"/paper/2507.13609"},"observation_digest":"sha256:b2246931cf30ebb72283a5254d283ed7d40cdc27a7de1ffbd0b52d22f801a79d","observation_id":"be343b91-acae-498a-835f-0d07ea9f6e8a","resolution":{"observed_at":"2026-08-06T16:24:42.092205Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.08485","last_updated":"2023-12-11T17:46:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-04-17T17:59:25Z","title":"Visual Instruction Tuning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.08485","snapshot_observed_at":"2026-08-06T16:24:41.268571Z","title":"arXiv preprint arXiv:2304.08485","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.13609","last_updated":"2025-07-18T02:29:19Z","snapshot_observed_at":"2026-08-07T13:58:45.437787Z","submitted_at":"2025-07-18T02:29:19Z","title":"CoTasks: Chain-of-Thought based Video Instruction Tuning Tasks","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T16:24:41.268571Z"},"links":{"cited_paper":"/paper/2304.08485","citing_paper":"/paper/2507.13609"},"observation_digest":"sha256:a6fe10b75fed87106004943729cb05c37d101175c28598ca660c766c4734eb1a","observation_id":"7b3bea5f-d3e0-4a6a-be3f-d6f46e219ab0","resolution":{"observed_at":"2026-08-06T16:24:41.268571Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.15343","last_updated":"2023-09-27T12:05:41Z","snapshot_observed_at":"2026-07-06T15:08:30.190912Z","submitted_at":"2023-03-27T15:53:01Z","title":"Sigmoid Loss for Language Image Pre-Training","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.15343","snapshot_observed_at":"2026-08-06T16:24:41.451116Z","title":"arXiv preprint arXiv:2303.15343","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.13609","last_updated":"2025-07-18T02:29:19Z","snapshot_observed_at":"2026-08-07T13:58:45.437787Z","submitted_at":"2025-07-18T02:29:19Z","title":"CoTasks: Chain-of-Thought based Video Instruction Tuning Tasks","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T16:24:41.451116Z"},"links":{"cited_paper":"/paper/2303.15343","citing_paper":"/paper/2507.13609"},"observation_digest":"sha256:20f915df228fe7366dd2d6e2ac36411ae9a727a0013f8abb1c691ba2f5217b29","observation_id":"54389f8f-62b4-4380-a9ee-6608ebb42840","resolution":{"observed_at":"2026-08-06T16:24:41.451116Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.12191","last_updated":"2024-10-03T15:54:49Z","snapshot_observed_at":"2026-08-06T05:35:29.109022Z","submitted_at":"2024-09-18T17:59:32Z","title":"Qwen2-VL: Enhancing Vision-Language Model's Perception of the World at Any Resolution","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.12191","snapshot_observed_at":"2026-08-06T16:24:41.515619Z","title":"arXiv preprint arXiv:2409.12191","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.13609","last_updated":"2025-07-18T02:29:19Z","snapshot_observed_at":"2026-08-07T13:58:45.437787Z","submitted_at":"2025-07-18T02:29:19Z","title":"CoTasks: Chain-of-Thought based Video Instruction Tuning Tasks","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T16:24:41.515619Z"},"links":{"cited_paper":"/paper/2409.12191","citing_paper":"/paper/2507.13609"},"observation_digest":"sha256:5e392715ce67574cabc3f4dfe0c740234a02c14665094dbfc551cd3344cf8aeb","observation_id":"ca354f7c-7c9a-428b-b1af-114d4308b64a","resolution":{"observed_at":"2026-08-06T16:24:41.515619Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.12559","last_updated":"2025-06-08T05:57:12Z","snapshot_observed_at":"2026-08-07T17:00:27.328884Z","submitted_at":"2025-03-16T16:14:52Z","title":"AdaReTaKe: Adaptive Redundancy Reduction to Perceive Longer for Video-language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.12559","snapshot_observed_at":"2026-08-06T16:24:41.591387Z","title":"arXiv preprint arXiv:2503.12559","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.13609","last_updated":"2025-07-18T02:29:19Z","snapshot_observed_at":"2026-08-07T13:58:45.437787Z","submitted_at":"2025-07-18T02:29:19Z","title":"CoTasks: Chain-of-Thought based Video Instruction Tuning Tasks","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T16:24:41.591387Z"},"links":{"cited_paper":"/paper/2503.12559","citing_paper":"/paper/2507.13609"},"observation_digest":"sha256:7738c388d26d78ca61377b692a8a12c570da3420b1082a59f6cc2d2398c8f646","observation_id":"ca1a1c3b-6481-4ff9-a327-12cf066bf2a2","resolution":{"observed_at":"2026-08-06T16:24:41.591387Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.15402","last_updated":"2024-06-06T01:58:54Z","snapshot_observed_at":"2026-07-06T16:24:10.053246Z","submitted_at":"2023-09-27T04:53:10Z","title":"Navigate through Enigmatic Labyrinth A Survey of Chain of Thought Reasoning: Advances, Frontiers and Future","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.15402","snapshot_observed_at":"2026-08-06T16:24:41.740062Z","title":"arXiv preprint arXiv:2309.15402","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.13609","last_updated":"2025-07-18T02:29:19Z","snapshot_observed_at":"2026-08-07T13:58:45.437787Z","submitted_at":"2025-07-18T02:29:19Z","title":"CoTasks: Chain-of-Thought based Video Instruction Tuning Tasks","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T16:24:41.740062Z"},"links":{"cited_paper":"/paper/2309.15402","citing_paper":"/paper/2507.13609"},"observation_digest":"sha256:71fd9a86efbcd3c025f1f197cb3954aecdac6935d13b56aac21466fb3ef7952b","observation_id":"637a9373-8e2c-4136-9393-eac8b48b9f35","resolution":{"observed_at":"2026-08-06T16:24:41.740062Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1907.06987","last_updated":"2022-10-17T19:58:40Z","snapshot_observed_at":"2026-07-06T08:08:03.081236Z","submitted_at":"2019-07-15T12:58:21Z","title":"A Short Note on the Kinetics-700 Human Action Dataset","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1907.06987","snapshot_observed_at":"2026-08-06T16:24:41.105486Z","title":"arXiv preprint arXiv:1907.06987","venue":null,"work_id":null,"year":1907},"citing_paper":{"arxiv_id":"2507.13609","last_updated":"2025-07-18T02:29:19Z","snapshot_observed_at":"2026-08-07T13:58:45.437787Z","submitted_at":"2025-07-18T02:29:19Z","title":"CoTasks: Chain-of-Thought based Video Instruction Tuning Tasks","version":1},"reference_index":2019,"source":"pdf_text","source_observed_at":"2026-08-06T16:24:41.105486Z"},"links":{"cited_paper":"/paper/1907.06987","citing_paper":"/paper/2507.13609"},"observation_digest":"sha256:5e450c64b2697067b9834ba94afcfd1a0a9bcd07b0fa5fe9e840d9a677fb1c66","observation_id":"90edebba-c5ba-42d7-95f5-bb9ba9faaf5c","resolution":{"observed_at":"2026-08-06T16:24:41.105486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2201.11903","last_updated":"2023-01-10T23:07:57Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-01-28T02:33:07Z","title":"Chain-of-Thought Prompting Elicits Reasoning in Large Language Models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2201.11903","snapshot_observed_at":"2026-08-06T16:24:41.655241Z","title":"arXiv preprint arXiv:2201.11903","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.13609","last_updated":"2025-07-18T02:29:19Z","snapshot_observed_at":"2026-08-07T13:58:45.437787Z","submitted_at":"2025-07-18T02:29:19Z","title":"CoTasks: Chain-of-Thought based Video Instruction Tuning Tasks","version":1},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-06T16:24:41.655241Z"},"links":{"cited_paper":"/paper/2201.11903","citing_paper":"/paper/2507.13609"},"observation_digest":"sha256:0df1c0d5dc4d0c5dd4d7eef2ba790b6e7abbd5ea45a3a287aa4963705af5a563","observation_id":"a9e17f96-67d1-4a74-9b81-b6ff2fa31646","resolution":{"observed_at":"2026-08-06T16:24:41.655241Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.12966","last_updated":"2023-10-13T02:41:28Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-08-24T17:59:17Z","title":"Qwen-VL: A Versatile Vision-Language Model for Understanding, Localization, Text Reading, and Beyond","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.12966","snapshot_observed_at":"2026-08-06T16:24:40.966042Z","title":"arXiv preprint arXiv:2308.12966","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.13609","last_updated":"2025-07-18T02:29:19Z","snapshot_observed_at":"2026-08-07T13:58:45.437787Z","submitted_at":"2025-07-18T02:29:19Z","title":"CoTasks: Chain-of-Thought based Video Instruction Tuning Tasks","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-06T16:24:40.966042Z"},"links":{"cited_paper":"/paper/2308.12966","citing_paper":"/paper/2507.13609"},"observation_digest":"sha256:ae88b4750aa2672bc72d0a95d5bd9b8e758e1234d0c9fa606577bac5b5288c79","observation_id":"bd7a54d2-4f38-43c4-bc32-37150467047f","resolution":{"observed_at":"2026-08-06T16:24:40.966042Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T16:24:42.438942Z","title":"In Proceedings of the 62nd Annual Meeting of the Association for Computational Linguistics (ACL 2024)","venue":null,"work_id":"aed04c19-411b-4f45-a266-34417b8533c2","year":2024},"citing_paper":{"arxiv_id":"2507.13609","last_updated":"2025-07-18T02:29:19Z","snapshot_observed_at":"2026-08-07T13:58:45.437787Z","submitted_at":"2025-07-18T02:29:19Z","title":"CoTasks: Chain-of-Thought based Video Instruction Tuning Tasks","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-06T16:24:41.352400Z"},"links":{"citing_paper":"/paper/2507.13609"},"observation_digest":"sha256:63122c8b80b498b7a358ea0181d4f69d7b6820b32be594c4c699f0f6a1dc9482","observation_id":"5b389022-fb87-4a7f-bc65-299628f2506d","resolution":{"observed_at":"2026-08-06T16:24:42.508347Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.13923","last_updated":"2025-02-19T18:00:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-02-19T18:00:14Z","title":"Qwen2.5-VL Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.13923","snapshot_observed_at":"2026-08-06T16:24:41.078263Z","title":"arXiv preprint arXiv:2502.13923","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.13609","last_updated":"2025-07-18T02:29:19Z","snapshot_observed_at":"2026-08-07T13:58:45.437787Z","submitted_at":"2025-07-18T02:29:19Z","title":"CoTasks: Chain-of-Thought based Video Instruction Tuning Tasks","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-06T16:24:41.078263Z"},"links":{"cited_paper":"/paper/2502.13923","citing_paper":"/paper/2507.13609"},"observation_digest":"sha256:9b2cce91cf2ec31aa487f5e83a06d7397b84c6d93f23f9a6e66d2a40b7c50d5e","observation_id":"93e158aa-1e0c-446c-b124-0d0dc9a25ede","resolution":{"observed_at":"2026-08-06T16:24:41.078263Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2507.13609","last_updated":"2025-07-18T02:29:19Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-07T13:58:45.437787Z","submitted_at":"2025-07-18T02:29:19Z","title":"CoTasks: Chain-of-Thought based Video Instruction Tuning Tasks"},"reference_resolution":{"displayed":12,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":9,"verified_exact":1,"verified_fuzzy":2},"total_outbound_references":12},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 12 of 12 outbound references and 1 inbound Pith citation observation for arXiv:2507.13609."}