{"as_of":"2026-08-20T08:35:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:d365a6a754aa1ccd36f0cfbf896b7796f1ef3ccae2ebc952f7dabff125aa4b8c","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":13,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":13,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-20T06:33:59.587034+00:00","state":"measured"},{"denominator":13,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":13,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T14:32:47.279314Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-06-29T22:44:02.042014Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2403.13447","last_updated":"2024-03-20T09:42:43Z","snapshot_observed_at":"2026-08-16T14:08:06.107215Z","submitted_at":"2024-03-20T09:42:43Z","title":"HyperLLaVA: Dynamic Visual and Language Expert Tuning for Multimodal Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.13447","snapshot_observed_at":"2026-08-11T00:08:16.299905Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.19684","last_updated":"2025-02-17T15:32:50Z","snapshot_observed_at":"2026-08-18T09:40:01.388208Z","submitted_at":"2024-12-27T15:21:17Z","title":"Boosting Private Domain Understanding of Efficient MLLMs: A Tuning-free, Adaptive, Universal Prompt Optimization Framework","version":2},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-11T00:08:16.299905Z"},"links":{"cited_paper":"/paper/2403.13447","citing_paper":"/paper/2412.19684"},"observation_digest":"sha256:a6866ead4c205a54469d3e9a21150efd00172f6f428d7924481105e5c888680e","observation_id":"9fd2e07e-d1f5-4c00-ba9b-0677e76443ac","resolution":{"observed_at":"2026-08-11T00:08:16.299905Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.13447","last_updated":"2024-03-20T09:42:43Z","snapshot_observed_at":"2026-08-16T14:08:06.107215Z","submitted_at":"2024-03-20T09:42:43Z","title":"HyperLLaVA: Dynamic Visual and Language Expert Tuning for Multimodal Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.13447","snapshot_observed_at":"2026-08-10T23:45:55.174446Z","title":"Hyperllava: Dynamic visual and language expert tuning for multimodal large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.19978","last_updated":"2024-12-28T02:36:51Z","snapshot_observed_at":"2026-08-18T20:49:58.164223Z","submitted_at":"2024-12-28T02:36:51Z","title":"MAKIMA: Tuning-free Multi-Attribute Open-domain Video Editing via Mask-Guided Attention Modulation","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-10T23:45:55.174446Z"},"links":{"cited_paper":"/paper/2403.13447","citing_paper":"/paper/2412.19978"},"observation_digest":"sha256:eca6a42a3fb487d792f7add72c7d55ebe72b7fbf30b94acb0bc99741f11a0007","observation_id":"6c70dfd5-0457-4155-8c9e-585af2f8449a","resolution":{"observed_at":"2026-08-10T23:45:55.174446Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.13447","last_updated":"2024-03-20T09:42:43Z","snapshot_observed_at":"2026-08-16T14:08:06.107215Z","submitted_at":"2024-03-20T09:42:43Z","title":"HyperLLaVA: Dynamic Visual and Language Expert Tuning for Multimodal Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.13447","snapshot_observed_at":"2026-08-10T22:52:30.354020Z","title":"Hyperllava: Dynamic visual and language expert tuning for multimodal large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.00599","last_updated":"2025-03-25T08:10:15Z","snapshot_observed_at":"2026-08-13T01:52:55.412320Z","submitted_at":"2024-12-31T18:56:46Z","title":"VideoRefer Suite: Advancing Spatial-Temporal Object Understanding with Video LLM","version":3},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-10T22:52:30.354020Z"},"links":{"cited_paper":"/paper/2403.13447","citing_paper":"/paper/2501.00599"},"observation_digest":"sha256:e7d7816bcda1be602497a31e664f89d669a95080ac70007bc0d2981323e55a13","observation_id":"bd1e044f-f6ce-4512-99eb-ad6fab58f73a","resolution":{"observed_at":"2026-08-10T22:52:30.354020Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.13447","last_updated":"2024-03-20T09:42:43Z","snapshot_observed_at":"2026-08-16T14:08:06.107215Z","submitted_at":"2024-03-20T09:42:43Z","title":"HyperLLaVA: Dynamic Visual and Language Expert Tuning for Multimodal Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.13447","snapshot_observed_at":"2026-08-07T13:20:33.922548Z","title":"arXiv preprint arXiv:2403.13447 (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.23830","last_updated":"2025-05-28T08:38:39Z","snapshot_observed_at":"2026-08-17T18:11:33.461839Z","submitted_at":"2025-05-28T08:38:39Z","title":"EvoMoE: Expert Evolution in Mixture of Experts for Multimodal Large Language Models","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T13:20:33.922548Z"},"links":{"cited_paper":"/paper/2403.13447","citing_paper":"/paper/2505.23830"},"observation_digest":"sha256:bd15742bdd3cdba1eedb4d32fc3701b13bffbf8c36794e8475c467b13ad3733f","observation_id":"ff969bf6-1f68-412c-b487-3d4e152c9be7","resolution":{"observed_at":"2026-08-07T13:20:33.922548Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.13447","last_updated":"2024-03-20T09:42:43Z","snapshot_observed_at":"2026-08-16T14:08:06.107215Z","submitted_at":"2024-03-20T09:42:43Z","title":"HyperLLaVA: Dynamic Visual and Language Expert Tuning for Multimodal Large Language Models","version":1},"cited_work":{"arxiv_id":"2403.13447","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.13447","snapshot_observed_at":"2026-06-29T22:44:02.042014Z","title":"arXiv preprint arXiv:2403.13447 , year=","venue":null,"work_id":"27dcf7ea-7ed5-478f-be11-26e21324465d","year":2024},"citing_paper":{"arxiv_id":"2506.05831","last_updated":"2026-04-07T10:06:53Z","snapshot_observed_at":"2026-08-11T06:23:45.741698Z","submitted_at":"2025-06-06T07:56:41Z","title":"HeartcareGPT: A Unified Multimodal ECG Suite for Dual Signal-Image Modeling and Understanding","version":4},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-19T10:44:01.880405Z"},"links":{"cited_paper":"/paper/2403.13447","citing_paper":"/paper/2506.05831"},"observation_digest":"sha256:2b779440e7aa7940cefdf3432e493fd28b5f2d0ea94a7009a9a594d4d9da33ea","observation_id":"506d5558-cab4-4d7a-9890-c4cdea11294b","resolution":{"observed_at":"2026-05-19T10:47:15.111720Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.13447","last_updated":"2024-03-20T09:42:43Z","snapshot_observed_at":"2026-08-16T14:08:06.107215Z","submitted_at":"2024-03-20T09:42:43Z","title":"HyperLLaVA: Dynamic Visual and Language Expert Tuning for Multimodal Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.13447","snapshot_observed_at":"2026-08-06T17:47:23.729881Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.10015","last_updated":"2025-07-17T11:10:58Z","snapshot_observed_at":"2026-08-18T14:42:18.270344Z","submitted_at":"2025-07-14T07:51:01Z","title":"(Almost) Free Modality Stitching of Foundation Models","version":3},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-06T17:47:23.729881Z"},"links":{"cited_paper":"/paper/2403.13447","citing_paper":"/paper/2507.10015"},"observation_digest":"sha256:5b5a3f81d232eab66b2d701e8a8c7235fd912ef6e7a202dea04a5477e9ea49fa","observation_id":"8def57ec-fe95-4cf0-a1ba-3bf30e0819d6","resolution":{"observed_at":"2026-08-06T17:47:23.729881Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.13447","last_updated":"2024-03-20T09:42:43Z","snapshot_observed_at":"2026-08-16T14:08:06.107215Z","submitted_at":"2024-03-20T09:42:43Z","title":"HyperLLaVA: Dynamic Visual and Language Expert Tuning for Multimodal Large Language Models","version":1},"cited_work":{"arxiv_id":"2403.13447","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.13447","snapshot_observed_at":"2026-06-29T22:44:02.042014Z","title":"arXiv preprint arXiv:2403.13447 , year=","venue":null,"work_id":"27dcf7ea-7ed5-478f-be11-26e21324465d","year":2024},"citing_paper":{"arxiv_id":"2604.11789","last_updated":"2026-04-20T14:38:53Z","snapshot_observed_at":"2026-08-09T05:10:13.009841Z","submitted_at":"2026-04-13T17:55:02Z","title":"LMMs Meet Object-Centric Vision: Understanding, Segmentation, Editing and Generation","version":2},"reference_index":229,"source":"pdf_text","source_observed_at":"2026-05-10T15:35:37.095627Z"},"links":{"cited_paper":"/paper/2403.13447","citing_paper":"/paper/2604.11789"},"observation_digest":"sha256:838219f962b94ccae9d1bd8b62ab909129cf2be01d626982379b29e5585f52f0","observation_id":"0e73cde8-548d-448d-9ecf-2855e7eb5475","resolution":{"observed_at":"2026-05-11T10:11:08.011056Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.13447","last_updated":"2024-03-20T09:42:43Z","snapshot_observed_at":"2026-08-16T14:08:06.107215Z","submitted_at":"2024-03-20T09:42:43Z","title":"HyperLLaVA: Dynamic Visual and Language Expert Tuning for Multimodal Large Language Models","version":1},"cited_work":{"arxiv_id":"2403.13447","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.13447","snapshot_observed_at":"2026-06-29T22:44:02.042014Z","title":"arXiv preprint arXiv:2403.13447 , year=","venue":null,"work_id":"27dcf7ea-7ed5-478f-be11-26e21324465d","year":2024},"citing_paper":{"arxiv_id":"2605.18621","last_updated":"2026-05-18T16:31:31Z","snapshot_observed_at":"2026-08-18T03:43:49.050514Z","submitted_at":"2026-05-18T16:31:31Z","title":"CrossView Suite: Harnessing Cross-view Spatial Intelligence of MLLMs with Dataset, Model and Benchmark","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-05-20T10:37:27.926364Z"},"links":{"cited_paper":"/paper/2403.13447","citing_paper":"/paper/2605.18621"},"observation_digest":"sha256:fc907a0a0a9797d95409fc36dae005029ea296489660bd18c1c79e6b17d3b14f","observation_id":"fcc84857-c64d-40d4-b97a-b21dd9fa491d","resolution":{"observed_at":"2026-05-20T10:38:12.296649Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.13447","last_updated":"2024-03-20T09:42:43Z","snapshot_observed_at":"2026-08-16T14:08:06.107215Z","submitted_at":"2024-03-20T09:42:43Z","title":"HyperLLaVA: Dynamic Visual and Language Expert Tuning for Multimodal Large Language Models","version":1},"cited_work":{"arxiv_id":"2403.13447","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.13447","snapshot_observed_at":"2026-06-29T22:44:02.042014Z","title":"arXiv preprint arXiv:2403.13447 , year=","venue":null,"work_id":"27dcf7ea-7ed5-478f-be11-26e21324465d","year":2024},"citing_paper":{"arxiv_id":"2605.19559","last_updated":"2026-05-19T09:02:20Z","snapshot_observed_at":"2026-08-16T01:19:28.511820Z","submitted_at":"2026-05-19T09:02:20Z","title":"EgoCoT-Bench: Benchmarking Grounded and Verifiable Operation-Centric Chain of Thought Reasoning for MLLMs","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-05-20T05:53:05.450946Z"},"links":{"cited_paper":"/paper/2403.13447","citing_paper":"/paper/2605.19559"},"observation_digest":"sha256:e90b365ddca5eba925c31366b8e37e6f0bad86268ea5dd83714cd8af3953ad8b","observation_id":"b99a429e-11ea-4bb5-b4cf-7b8d25016fed","resolution":{"observed_at":"2026-05-20T05:53:22.300589Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.13447","last_updated":"2024-03-20T09:42:43Z","snapshot_observed_at":"2026-08-16T14:08:06.107215Z","submitted_at":"2024-03-20T09:42:43Z","title":"HyperLLaVA: Dynamic Visual and Language Expert Tuning for Multimodal Large Language Models","version":1},"cited_work":{"arxiv_id":"2403.13447","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.13447","snapshot_observed_at":"2026-06-29T22:44:02.042014Z","title":"arXiv preprint arXiv:2403.13447 , year=","venue":null,"work_id":"27dcf7ea-7ed5-478f-be11-26e21324465d","year":2024},"citing_paper":{"arxiv_id":"2605.20277","last_updated":"2026-05-19T04:33:27Z","snapshot_observed_at":"2026-08-16T00:30:16.458616Z","submitted_at":"2026-05-19T04:33:27Z","title":"Regulating Anatomy-Aware Rewards via Trajectory-Integral Feedback for Volumetric Computed Tomography Analysis","version":1},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-05-21T08:06:02.176934Z"},"links":{"cited_paper":"/paper/2403.13447","citing_paper":"/paper/2605.20277"},"observation_digest":"sha256:2f16776e4d2e7d2d82a28cccc2656ebbb023e9f0e875832ae2fb11d683100eaf","observation_id":"406d8b95-a340-458f-b6c6-527213da1d84","resolution":{"observed_at":"2026-05-21T08:09:51.732871Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.13447","last_updated":"2024-03-20T09:42:43Z","snapshot_observed_at":"2026-08-16T14:08:06.107215Z","submitted_at":"2024-03-20T09:42:43Z","title":"HyperLLaVA: Dynamic Visual and Language Expert Tuning for Multimodal Large Language Models","version":1},"cited_work":{"arxiv_id":"2403.13447","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.13447","snapshot_observed_at":"2026-06-29T22:44:02.042014Z","title":"arXiv preprint arXiv:2403.13447 , year=","venue":null,"work_id":"27dcf7ea-7ed5-478f-be11-26e21324465d","year":2024},"citing_paper":{"arxiv_id":"2605.26102","last_updated":"2026-05-31T07:20:32Z","snapshot_observed_at":"2026-08-15T01:55:04.903017Z","submitted_at":"2026-05-25T17:58:03Z","title":"InstructSAM: Segment Any Instance with Any Instructions","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-06-29T22:34:00.442420Z"},"links":{"cited_paper":"/paper/2403.13447","citing_paper":"/paper/2605.26102"},"observation_digest":"sha256:b2ce22cda2e6a45e4c5b8b248e44d609dabd7455e50903e6d81a9f459e061313","observation_id":"8cf38b1b-8881-4395-b705-e38bf45b5cdd","resolution":{"observed_at":"2026-06-29T22:44:02.043609Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.13447","last_updated":"2024-03-20T09:42:43Z","snapshot_observed_at":"2026-08-16T14:08:06.107215Z","submitted_at":"2024-03-20T09:42:43Z","title":"HyperLLaVA: Dynamic Visual and Language Expert Tuning for Multimodal Large Language Models","version":1},"cited_work":{"arxiv_id":"2403.13447","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.13447","snapshot_observed_at":"2026-06-29T22:44:02.042014Z","title":"arXiv preprint arXiv:2403.13447 , year=","venue":null,"work_id":"27dcf7ea-7ed5-478f-be11-26e21324465d","year":2024},"citing_paper":{"arxiv_id":"2605.30011","last_updated":"2026-05-28T14:36:53Z","snapshot_observed_at":"2026-08-14T22:48:59.292649Z","submitted_at":"2026-05-28T14:36:53Z","title":"VisualThink-VLA: Visual Intermediate Reasoning for Effective and Low-Latency Vision-Language-Action Policies","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-06-29T08:12:24.023638Z"},"links":{"cited_paper":"/paper/2403.13447","citing_paper":"/paper/2605.30011"},"observation_digest":"sha256:58bd37580f2c57def6d27f007c4517e95c12fcea3f29eab8026221c620e2cb00","observation_id":"18329300-22f7-42bb-85e4-0dc8030b575e","resolution":{"observed_at":"2026-06-29T08:13:14.673968Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.13447","last_updated":"2024-03-20T09:42:43Z","snapshot_observed_at":"2026-08-16T14:08:06.107215Z","submitted_at":"2024-03-20T09:42:43Z","title":"HyperLLaVA: Dynamic Visual and Language Expert Tuning for Multimodal Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.13447","snapshot_observed_at":"2026-08-15T14:32:47.279314Z","title":"2403.13447 , archivePrefix =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.07371","last_updated":"2026-08-07T16:12:58Z","snapshot_observed_at":"2026-08-20T02:29:30.779952Z","submitted_at":"2026-08-07T16:12:58Z","title":"Trajectory-Relative Hindsight Distillation for Agentic Reinforcement Learning","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-15T14:32:47.279314Z"},"links":{"cited_paper":"/paper/2403.13447","citing_paper":"/paper/2608.07371"},"observation_digest":"sha256:004ce2b689f307cb2338b7994f90e638778e7a01b4fad905d52f64eb3ac8b8af","observation_id":"e6cf9bf9-8a6c-4b66-8880-d69a209ba3a9","resolution":{"observed_at":"2026-08-15T14:32:47.279314Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2403.13447/citation-record","integrity":"/paper/2403.13447/integrity","json":"/paper/2403.13447/citation-record.json","paper":"/paper/2403.13447"},"outbound":[],"paper":{"arxiv_id":"2403.13447","last_updated":"2024-03-20T09:42:43Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-16T14:08:06.107215Z","submitted_at":"2024-03-20T09:42:43Z","title":"HyperLLaVA: Dynamic Visual and Language Expert Tuning for Multimodal Large Language Models"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"thesis":"As of 20 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 13 inbound Pith citation observations for arXiv:2403.13447."}