{"as_of":"2026-08-10T19:12:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:869dd778c0d4b91420b0a839596ae3fa5dbd7c6ecba5fab0dca439edd2dc181b","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":8,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":8,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":8,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":8,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T14:14:41.573083Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T17:29:59.942508Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2010.03743","last_updated":"2021-09-13T18:53:35Z","snapshot_observed_at":"2026-07-06T10:02:30.383493Z","submitted_at":"2020-10-08T03:07:00Z","title":"Visual News: Benchmark and Challenges in News Image Captioning","version":3},"cited_work":{"arxiv_id":"2010.03743","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2010.03743","snapshot_observed_at":"2026-07-04T17:29:59.942508Z","title":"Visual news: Benchmark and challenges in news image captioning","venue":null,"work_id":"330a683a-591c-4216-a638-9dd74b7c70aa","year":2010},"citing_paper":{"arxiv_id":"2306.14565","last_updated":"2024-03-19T22:53:25Z","snapshot_observed_at":"2026-08-06T22:33:34.254048Z","submitted_at":"2023-06-26T10:26:33Z","title":"Mitigating Hallucination in Large Multi-Modal Models via Robust Instruction Tuning","version":4},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-14T17:34:56.836034Z"},"links":{"cited_paper":"/paper/2010.03743","citing_paper":"/paper/2306.14565"},"observation_digest":"sha256:4746c2c846197906490f1533961d94cf4ff61508e7fa6ada76ccd46470a1fab7","observation_id":"7e312fc7-f0ce-4e86-afba-61b28b3ff7ab","resolution":{"observed_at":"2026-05-14T17:34:56.959055Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2010.03743","last_updated":"2021-09-13T18:53:35Z","snapshot_observed_at":"2026-07-06T10:02:30.383493Z","submitted_at":"2020-10-08T03:07:00Z","title":"Visual News: Benchmark and Challenges in News Image Captioning","version":3},"cited_work":{"arxiv_id":"2010.03743","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2010.03743","snapshot_observed_at":"2026-07-04T17:29:59.942508Z","title":"Visual news: Benchmark and challenges in news image captioning","venue":null,"work_id":"330a683a-591c-4216-a638-9dd74b7c70aa","year":2010},"citing_paper":{"arxiv_id":"2310.14566","last_updated":"2024-03-25T06:05:24Z","snapshot_observed_at":"2026-08-02T21:20:28.501823Z","submitted_at":"2023-10-23T04:49:09Z","title":"HallusionBench: An Advanced Diagnostic Suite for Entangled Language Hallucination and Visual Illusion in Large Vision-Language Models","version":5},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-17T01:22:04.035994Z"},"links":{"cited_paper":"/paper/2010.03743","citing_paper":"/paper/2310.14566"},"observation_digest":"sha256:c50b4b4e26385c5f076ac1ad0e5a97252eeed560babd71140a680c7a87201beb","observation_id":"c55f76d9-ecdf-422b-a5bf-01e6f9f4e171","resolution":{"observed_at":"2026-05-17T01:22:04.097788Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2010.03743","last_updated":"2021-09-13T18:53:35Z","snapshot_observed_at":"2026-07-06T10:02:30.383493Z","submitted_at":"2020-10-08T03:07:00Z","title":"Visual News: Benchmark and Challenges in News Image Captioning","version":3},"cited_work":{"arxiv_id":"2010.03743","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2010.03743","snapshot_observed_at":"2026-07-04T17:29:59.942508Z","title":"Visual news: Benchmark and challenges in news image captioning","venue":null,"work_id":"330a683a-591c-4216-a638-9dd74b7c70aa","year":2010},"citing_paper":{"arxiv_id":"2410.05160","last_updated":"2025-01-02T05:26:47Z","snapshot_observed_at":"2026-07-06T19:29:05.492290Z","submitted_at":"2024-10-07T16:14:05Z","title":"VLM2Vec: Training Vision-Language Models for Massive Multimodal Embedding Tasks","version":3},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-17T21:19:43.882232Z"},"links":{"cited_paper":"/paper/2010.03743","citing_paper":"/paper/2410.05160"},"observation_digest":"sha256:204ee15cad799170de38e07717cc85cb42298cb5db7a46f8cb21ef4e8e10aa86","observation_id":"fb20b0ee-17f2-454c-bded-af0540c217a6","resolution":{"observed_at":"2026-05-17T21:19:43.993625Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2010.03743","last_updated":"2021-09-13T18:53:35Z","snapshot_observed_at":"2026-07-06T10:02:30.383493Z","submitted_at":"2020-10-08T03:07:00Z","title":"Visual News: Benchmark and Challenges in News Image Captioning","version":3},"cited_work":{"arxiv_id":"2010.03743","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2010.03743","snapshot_observed_at":"2026-07-04T17:29:59.942508Z","title":"Visual news: Benchmark and challenges in news image captioning","venue":null,"work_id":"330a683a-591c-4216-a638-9dd74b7c70aa","year":2010},"citing_paper":{"arxiv_id":"2504.18361","last_updated":"2026-05-15T09:37:39Z","snapshot_observed_at":"2026-07-06T21:14:48.026751Z","submitted_at":"2025-04-25T14:04:36Z","title":"COCO-Inpaint: A Benchmark for Detecting and Localizing Inpainting-Based Image Manipulations","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-22T17:51:40.120648Z"},"links":{"cited_paper":"/paper/2010.03743","citing_paper":"/paper/2504.18361"},"observation_digest":"sha256:6f6c361809043c2acb8d0cb3bd13275c9117c8f48854461d922eb3177ff3d220","observation_id":"f862ca8b-5817-4984-b96e-318cfddf26b8","resolution":{"observed_at":"2026-05-22T17:51:54.515678Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2010.03743","last_updated":"2021-09-13T18:53:35Z","snapshot_observed_at":"2026-07-06T10:02:30.383493Z","submitted_at":"2020-10-08T03:07:00Z","title":"Visual News: Benchmark and Challenges in News Image Captioning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.03743","snapshot_observed_at":"2026-08-07T14:14:41.573083Z","title":"Visual news: Benchmark and challenges in news image captioning.arXiv preprint arXiv:2010.03743, 2020","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2505.19650","last_updated":"2025-05-27T11:57:17Z","snapshot_observed_at":"2026-08-09T09:30:25.697403Z","submitted_at":"2025-05-26T08:09:44Z","title":"Modality Curation: Building Universal Embeddings for Advanced Multimodal Information Retrieval","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T14:14:41.573083Z"},"links":{"cited_paper":"/paper/2010.03743","citing_paper":"/paper/2505.19650"},"observation_digest":"sha256:f68285e6c49ac0a3dfcb953ad582fef3f2cfb34d1d814a30608775f52a834caf","observation_id":"e83cd7ee-0e9e-4b0d-a869-07c9dd7e76c4","resolution":{"observed_at":"2026-08-07T14:14:41.573083Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2010.03743","last_updated":"2021-09-13T18:53:35Z","snapshot_observed_at":"2026-07-06T10:02:30.383493Z","submitted_at":"2020-10-08T03:07:00Z","title":"Visual News: Benchmark and Challenges in News Image Captioning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.03743","snapshot_observed_at":"2026-08-05T14:57:50.911361Z","title":"Visual news: Benchmark and challenges in news image captioning.arXiv preprint arXiv:2010.03743, 2020","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2508.20670","last_updated":"2025-09-09T06:47:10Z","snapshot_observed_at":"2026-08-05T14:57:50.334172Z","submitted_at":"2025-08-28T11:22:15Z","title":"\"Humor, Art, or Misinformation?\": A Multimodal Dataset for Intent-Aware Synthetic Image Detection","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-05T14:57:50.911361Z"},"links":{"cited_paper":"/paper/2010.03743","citing_paper":"/paper/2508.20670"},"observation_digest":"sha256:7cae57c693e1871e6c240e3a81bb70860b3634a6e63f381d1188e9fc12c841fc","observation_id":"ddd491fb-3710-40f5-bf59-0558232b45e7","resolution":{"observed_at":"2026-08-05T14:57:50.911361Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2010.03743","last_updated":"2021-09-13T18:53:35Z","snapshot_observed_at":"2026-07-06T10:02:30.383493Z","submitted_at":"2020-10-08T03:07:00Z","title":"Visual News: Benchmark and Challenges in News Image Captioning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.03743","snapshot_observed_at":"2026-08-05T12:46:43.300999Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2509.01259","last_updated":"2025-09-01T08:48:33Z","snapshot_observed_at":"2026-08-10T12:28:46.448747Z","submitted_at":"2025-09-01T08:48:33Z","title":"ReCap: Event-Aware Image Captioning with Article Retrieval and Semantic Gaussian Normalization","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-05T12:46:43.300999Z"},"links":{"cited_paper":"/paper/2010.03743","citing_paper":"/paper/2509.01259"},"observation_digest":"sha256:6598e9ee0dc68f71042b8332b25bd0680912dd1543131f97dd9d4e7902c240d9","observation_id":"b585c92a-61e7-46ee-9e52-f5a27fe01c27","resolution":{"observed_at":"2026-08-05T12:46:43.300999Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2010.03743","last_updated":"2021-09-13T18:53:35Z","snapshot_observed_at":"2026-07-06T10:02:30.383493Z","submitted_at":"2020-10-08T03:07:00Z","title":"Visual News: Benchmark and Challenges in News Image Captioning","version":3},"cited_work":{"arxiv_id":"2010.03743","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2010.03743","snapshot_observed_at":"2026-07-04T17:29:59.942508Z","title":"Visual news: Benchmark and challenges in news image captioning","venue":null,"work_id":"330a683a-591c-4216-a638-9dd74b7c70aa","year":2010},"citing_paper":{"arxiv_id":"2606.24627","last_updated":"2026-06-23T14:23:56Z","snapshot_observed_at":"2026-08-02T19:23:15.341233Z","submitted_at":"2026-06-23T14:23:56Z","title":"The Warrant Gap: Claim-Conditioned Re-scoring for Fact-Checking","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-06-25T23:42:53.179881Z"},"links":{"cited_paper":"/paper/2010.03743","citing_paper":"/paper/2606.24627"},"observation_digest":"sha256:92e335abd8c34db3aff48bd89f79b973f8b515ba63115c082945a3900d10dac1","observation_id":"c3e2ff3a-461c-428a-bc1c-d5044039da15","resolution":{"observed_at":"2026-07-04T17:29:59.943845Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2010.03743/citation-record","integrity":"/paper/2010.03743/integrity","json":"/paper/2010.03743/citation-record.json","paper":"/paper/2010.03743"},"outbound":[],"paper":{"arxiv_id":"2010.03743","last_updated":"2021-09-13T18:53:35Z","latest_version":3,"primary_category":"cs.CV","snapshot_observed_at":"2026-07-06T10:02:30.383493Z","submitted_at":"2020-10-08T03:07:00Z","title":"Visual News: Benchmark and Challenges in News Image Captioning"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 8 inbound Pith citation observations for arXiv:2010.03743."}