{"as_of":"2026-08-10T09:46:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:17d8a856b59d0d7c2c5e801f47f9db005b44524a1f6995ef76b8719df29b1653","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":19,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":19,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":19,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":19,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T15:29:56.768023Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T15:19:55.949718Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":"2107.12666","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-07-04T15:19:55.949718Z","title":"Semantically self-aligned network for text-to- image part-aware person re-identification","venue":null,"work_id":"b08047d7-907a-4bb0-9f46-a38834c7d788","year":2021},"citing_paper":{"arxiv_id":"2503.22171","last_updated":"2026-04-17T04:53:40Z","snapshot_observed_at":"2026-08-06T10:42:36.270805Z","submitted_at":"2025-03-28T06:18:15Z","title":"An Empirical Study of Validating Synthetic Data for Text-Based Person Retrieval","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-22T22:59:36.542774Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2503.22171"},"observation_digest":"sha256:09bc3839449787f1b8521fbbbd3d648596e58b7ba340cf36d1470485d857f655","observation_id":"dc5143cb-705e-44d0-b891-f92ddbb3ee2b","resolution":{"observed_at":"2026-05-22T23:02:13.680384Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-08-07T15:29:56.768023Z","title":"Semantically self-aligned network for text-to- image part-aware person re-identification","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.11036","last_updated":"2025-05-21T02:26:17Z","snapshot_observed_at":"2026-08-07T15:23:09.503029Z","submitted_at":"2025-05-21T02:26:17Z","title":"Human-centered Interactive Learning via MLLMs for Text-to-Image Person Re-identification","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T15:29:56.768023Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2506.11036"},"observation_digest":"sha256:68ec5946b3af75e4c869986a15164b58cf7e9df592fea828029d1fa7912df481","observation_id":"8b716916-f028-4b17-aa23-3cdb0742ed1a","resolution":{"observed_at":"2026-08-07T15:29:56.768023Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-08-06T18:59:47.088064Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.06744","last_updated":"2025-07-09T10:59:13Z","snapshot_observed_at":"2026-08-06T18:53:07.845028Z","submitted_at":"2025-07-09T10:59:13Z","title":"Dual-Granularity Cross-Modal Identity Association for Weakly-Supervised Text-to-Person Image Matching","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T18:59:47.088064Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2507.06744"},"observation_digest":"sha256:73bbf0e5c752f05612a07814215b805aa87f2f6a531646367e8aabaaa362c98f","observation_id":"2d6c15b2-9baf-4c3c-af26-e1a2b791fbeb","resolution":{"observed_at":"2026-08-06T18:59:47.088064Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-08-05T23:17:12.403417Z","title":"Semantically self-aligned network for text-to-image part-aware person re-identification,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2508.05547","last_updated":"2025-08-07T16:27:37Z","snapshot_observed_at":"2026-08-06T15:48:30.505708Z","submitted_at":"2025-08-07T16:27:37Z","title":"Adapting Vision-Language Models Without Labels: A Comprehensive Survey","version":1},"reference_index":275,"source":"pdf_text","source_observed_at":"2026-08-05T23:17:12.403417Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2508.05547"},"observation_digest":"sha256:cbd7422c87d669c303b4158d22ccf0657f6edd8893a6ef149b6b1d123e238887","observation_id":"c3a9d907-74d8-46c4-a4df-1a3199f5f089","resolution":{"observed_at":"2026-08-05T23:17:12.403417Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-08-05T22:22:58.625918Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2508.07144","last_updated":"2025-08-10T02:27:06Z","snapshot_observed_at":"2026-08-10T05:03:55.942200Z","submitted_at":"2025-08-10T02:27:06Z","title":"Dynamic Pattern Alignment Learning for Pretraining Lightweight Human-Centric Vision Models","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-05T22:22:58.625918Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2508.07144"},"observation_digest":"sha256:f436b1aedb7186931816c1df3857b2aa0a3af5cbfe7c5c52daf3101713117c0c","observation_id":"226d3e76-7cbc-41b5-8836-f8d4e9305d57","resolution":{"observed_at":"2026-08-05T22:22:58.625918Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-08-04T19:45:23.803453Z","title":"https: //github.com/kakaobrain/coyo-dataset","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.09118","last_updated":"2025-09-11T03:06:22Z","snapshot_observed_at":"2026-08-09T21:28:41.255043Z","submitted_at":"2025-09-11T03:06:22Z","title":"Gradient-Attention Guided Dual-Masking Synergetic Framework for Robust Text-based Person Retrieval","version":1},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-04T19:45:23.803453Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2509.09118"},"observation_digest":"sha256:921e5d8735dbb456529f5929fb34b1cd1f45b53b2e19451ef2066e1200ffd95e","observation_id":"096afdae-fe43-44b1-a99a-b757c8203d4e","resolution":{"observed_at":"2026-08-04T19:45:23.803453Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":"2107.12666","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-07-04T15:19:55.949718Z","title":"Semantically self-aligned network for text-to- image part-aware person re-identification","venue":null,"work_id":"b08047d7-907a-4bb0-9f46-a38834c7d788","year":2021},"citing_paper":{"arxiv_id":"2604.05504","last_updated":"2026-04-07T06:58:12Z","snapshot_observed_at":"2026-07-06T22:54:12.531121Z","submitted_at":"2026-04-07T06:58:12Z","title":"Semantic Communication with an LLM-enabled Knowledge Base","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-05-10T19:37:27.193183Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2604.05504"},"observation_digest":"sha256:39b2cfbc975b01a534b6110be6147c61f339e15969addae0d5d25a530c64aaf8","observation_id":"d65f39f4-f2b2-426a-82b3-f8b667356c72","resolution":{"observed_at":"2026-05-10T22:40:52.669941Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":"2107.12666","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-07-04T15:19:55.949718Z","title":"Semantically self-aligned network for text-to- image part-aware person re-identification","venue":null,"work_id":"b08047d7-907a-4bb0-9f46-a38834c7d788","year":2021},"citing_paper":{"arxiv_id":"2604.08598","last_updated":"2026-04-27T09:35:39Z","snapshot_observed_at":"2026-08-04T04:38:47.127542Z","submitted_at":"2026-04-07T07:48:15Z","title":"Pretrain-then-Adapt: Uncertainty-Aware Test-Time Adaptation for Text-based Person Search","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-10T19:14:34.478933Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2604.08598"},"observation_digest":"sha256:6f468e156267e10de3a6520c67c216215bfcdb619b02a2512e87065810c37be0","observation_id":"32b65575-a9e9-45a0-ae48-9c2f8711865e","resolution":{"observed_at":"2026-05-10T23:20:49.703340Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":"2107.12666","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-07-04T15:19:55.949718Z","title":"Semantically self-aligned network for text-to- image part-aware person re-identification","venue":null,"work_id":"b08047d7-907a-4bb0-9f46-a38834c7d788","year":2021},"citing_paper":{"arxiv_id":"2604.18376","last_updated":"2026-04-20T15:03:01Z","snapshot_observed_at":"2026-08-01T17:30:57.593119Z","submitted_at":"2026-04-20T15:03:01Z","title":"Towards Robust Text-to-Image Person Retrieval: Multi-View Reformulation for Semantic Compensation","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-10T04:46:47.870281Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2604.18376"},"observation_digest":"sha256:c3e8df26b77a496639bee31e82e1e6cb5a83d75233fc3c279ac13e0e0b1c9737","observation_id":"707ff598-d1bc-4ea3-9da1-27d25cd0861c","resolution":{"observed_at":"2026-05-10T11:40:19.419531Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":"2107.12666","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-07-04T15:19:55.949718Z","title":"Semantically self-aligned network for text-to- image part-aware person re-identification","venue":null,"work_id":"b08047d7-907a-4bb0-9f46-a38834c7d788","year":2021},"citing_paper":{"arxiv_id":"2604.23282","last_updated":"2026-05-27T12:43:38Z","snapshot_observed_at":"2026-08-06T18:29:07.423664Z","submitted_at":"2026-04-25T12:53:15Z","title":"Bridging the Pose-Semantic Gap: A Cascade Framework for Text-Based Person Anomaly Search","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-05-08T08:27:45.596741Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2604.23282"},"observation_digest":"sha256:59a9d0ac7b00bb39d209f0b22add3a1bde1455ba70f5532baae75c72595994a9","observation_id":"e994f592-9c8c-413f-b4bd-2ce461ab6d74","resolution":{"observed_at":"2026-05-11T20:36:11.119781Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":"2107.12666","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-07-04T15:19:55.949718Z","title":"Semantically self-aligned network for text-to- image part-aware person re-identification","venue":null,"work_id":"b08047d7-907a-4bb0-9f46-a38834c7d788","year":2021},"citing_paper":{"arxiv_id":"2604.23282","last_updated":"2026-05-27T12:43:38Z","snapshot_observed_at":"2026-08-06T18:29:07.423664Z","submitted_at":"2026-04-25T12:53:15Z","title":"Bridging the Pose-Semantic Gap: A Cascade Framework for Text-Based Person Anomaly Search","version":2},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-07-04T15:13:26.781464Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2604.23282"},"observation_digest":"sha256:d27d7675dda1c2411409ae611248060262007864107ccbb3442e92a7624b896f","observation_id":"e2fb091c-1568-4afe-9984-331ef1ba8d9d","resolution":{"observed_at":"2026-07-04T15:19:55.951642Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":"2107.12666","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-07-04T15:19:55.949718Z","title":"Semantically self-aligned network for text-to- image part-aware person re-identification","venue":null,"work_id":"b08047d7-907a-4bb0-9f46-a38834c7d788","year":2021},"citing_paper":{"arxiv_id":"2604.27122","last_updated":"2026-06-27T17:42:07Z","snapshot_observed_at":"2026-08-08T17:21:14.017881Z","submitted_at":"2026-04-29T19:18:36Z","title":"InterPartAbility: Phrase-Region Grounding for Interpretable Text-to-Image Person Re-Identification","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-07T09:29:09.212522Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2604.27122"},"observation_digest":"sha256:594371fabe6efec90b66add48470c8ace4e5a52ae4d507d51733e0ba4f5b479a","observation_id":"9a5347b8-7e39-4b8a-a202-2f5cf0ac6283","resolution":{"observed_at":"2026-05-12T09:46:26.300090Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":"2107.12666","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-07-04T15:19:55.949718Z","title":"Semantically self-aligned network for text-to- image part-aware person re-identification","venue":null,"work_id":"b08047d7-907a-4bb0-9f46-a38834c7d788","year":2021},"citing_paper":{"arxiv_id":"2604.27122","last_updated":"2026-06-27T17:42:07Z","snapshot_observed_at":"2026-08-08T17:21:14.017881Z","submitted_at":"2026-04-29T19:18:36Z","title":"InterPartAbility: Phrase-Region Grounding for Interpretable Text-to-Image Person Re-Identification","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-07-01T08:23:22.559471Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2604.27122"},"observation_digest":"sha256:4c69791e1627df3b9c391c8a8ea588aff2e150bd14f85a55fa9eee5abd6761d4","observation_id":"b7936741-81c3-4136-9cae-56f7e3036bb3","resolution":{"observed_at":"2026-07-01T08:25:33.093659Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":"2107.12666","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-07-04T15:19:55.949718Z","title":"Semantically self-aligned network for text-to- image part-aware person re-identification","venue":null,"work_id":"b08047d7-907a-4bb0-9f46-a38834c7d788","year":2021},"citing_paper":{"arxiv_id":"2606.01825","last_updated":"2026-07-01T05:50:06Z","snapshot_observed_at":"2026-08-05T10:25:09.105468Z","submitted_at":"2026-06-01T07:41:44Z","title":"ROGLE: Robust Global-Local Alignment with Automated Region Supervision for Text-Based Person Search","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-06-28T15:18:17.427707Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2606.01825"},"observation_digest":"sha256:e5863b54bf58d4e6e28bf8b1750d7f3aa9b22b9abb38102c8ca8582bd1312417","observation_id":"f09ad2ac-2527-4262-be5f-b5e3392de30a","resolution":{"observed_at":"2026-07-01T22:36:16.990911Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":"2107.12666","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-07-04T15:19:55.949718Z","title":"Semantically self-aligned network for text-to- image part-aware person re-identification","venue":null,"work_id":"b08047d7-907a-4bb0-9f46-a38834c7d788","year":2021},"citing_paper":{"arxiv_id":"2606.01825","last_updated":"2026-07-01T05:50:06Z","snapshot_observed_at":"2026-08-05T10:25:09.105468Z","submitted_at":"2026-06-01T07:41:44Z","title":"ROGLE: Robust Global-Local Alignment with Automated Region Supervision for Text-Based Person Search","version":2},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-07-02T23:17:03.456746Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2606.01825"},"observation_digest":"sha256:4ec6cd12dd2f85e4bd99883c34e555d71284fcdf092f43ce1c27cfb2d0b89372","observation_id":"d6427b9e-89cd-4f39-be55-c0f8d696103f","resolution":{"observed_at":"2026-07-02T23:17:28.901638Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":"2107.12666","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-07-04T15:19:55.949718Z","title":"Semantically self-aligned network for text-to- image part-aware person re-identification","venue":null,"work_id":"b08047d7-907a-4bb0-9f46-a38834c7d788","year":2021},"citing_paper":{"arxiv_id":"2606.02242","last_updated":"2026-06-01T13:37:55Z","snapshot_observed_at":"2026-08-07T10:32:57.890413Z","submitted_at":"2026-06-01T13:37:55Z","title":"Towards Resolving Optimization Conflicts Between Image- and Text-Based Person Re-Identification","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-06-28T15:36:19.675877Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2606.02242"},"observation_digest":"sha256:04e9f4dd3513c1e080ea5c2debee359c681c6fd618889b8df399630ad47c74f6","observation_id":"43e207be-1edc-4a82-ac80-abca59234059","resolution":{"observed_at":"2026-07-01T22:16:16.252018Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":"2107.12666","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-07-04T15:19:55.949718Z","title":"Semantically self-aligned network for text-to- image part-aware person re-identification","venue":null,"work_id":"b08047d7-907a-4bb0-9f46-a38834c7d788","year":2021},"citing_paper":{"arxiv_id":"2606.30458","last_updated":"2026-06-29T15:24:03Z","snapshot_observed_at":"2026-08-07T19:25:03.772420Z","submitted_at":"2026-06-29T15:24:03Z","title":"Cross-Resolution Semantic Transfer for Robust Text-to-Image Retrieval in Low-Resolution Surveillance","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-06-30T06:52:54.640706Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2606.30458"},"observation_digest":"sha256:22732e0f8211f2f4b5cc7028b8c544e4b8260e0464195852ebfaea99b3a6a7a1","observation_id":"e4e0a5c0-8fa9-422d-9329-05f874cae751","resolution":{"observed_at":"2026-06-30T06:54:20.310169Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-08-02T01:01:00.841746Z","title":"arXiv preprint arXiv:2107.12666","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.14821","last_updated":"2026-07-16T10:39:22Z","snapshot_observed_at":"2026-08-07T23:46:00.212599Z","submitted_at":"2026-07-16T10:39:22Z","title":"Blurring Modal Boundaries: A Unified Survey from Single- to Multi-Modal Person Re-ldentification","version":1},"reference_index":130,"source":"pdf_text","source_observed_at":"2026-08-02T01:01:00.841746Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2607.14821"},"observation_digest":"sha256:e4bd3cf0166b2476272dfc7d00d607589456a0215a882cae081738d7e7122509","observation_id":"b8cc3981-4ee6-4879-b171-8394ee790cc3","resolution":{"observed_at":"2026-08-02T01:01:00.841746Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-08-01T08:39:52.070351Z","title":"Semantically self-aligned network for text-to-image part-aware person re-identification,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.21057","last_updated":"2026-07-23T08:42:25Z","snapshot_observed_at":"2026-08-06T17:46:04.605724Z","submitted_at":"2026-07-23T08:42:25Z","title":"Achieving Text-based Person Retrieval with Any Granularity","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-01T08:39:52.070351Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2607.21057"},"observation_digest":"sha256:befd3e0bf4481dc3ba79a32a0a38999e8370134f4bff5283bfb83446c4360194","observation_id":"ecf7bc91-b023-4c86-b00a-b64413d76e35","resolution":{"observed_at":"2026-08-01T08:39:52.070351Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2107.12666/citation-record","integrity":"/paper/2107.12666/integrity","json":"/paper/2107.12666/citation-record.json","paper":"/paper/2107.12666"},"outbound":[],"paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-09T23:19:07.484167Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 19 inbound Pith citation observations for arXiv:2107.12666."}