{"as_of":"2026-08-21T22:39:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:0e271dd98eb11e9d071112d2d7a3ba24d40b6d59bb84059574a51a5c7a09adf2","coverage":[{"denominator":50,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":50,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T05:43:07.337382Z","state":"measured"},{"denominator":50,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":50,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-21T06:32:19.484+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2506.13774/citation-record","integrity":"/paper/2506.13774/integrity","json":"/paper/2506.13774/citation-record.json","paper":"/paper/2506.13774"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.257158Z","title":"Artificial Intelligence, Values, and Alignment","venue":null,"work_id":"7712a691-d171-4a90-82a5-77521c0d31f6","year":2020},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.102830Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:d530905387fbd26d4aea31f563a9f3bf0c213a96693c3cf58310955faac52feb","observation_id":"03ffc49a-1a7b-429c-8020-54a96502feb0","resolution":{"observed_at":"2026-08-07T05:43:08.261527Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.243311Z","title":"Artificial Morality: Top -down, Bottom-up, and Hybrid Approaches","venue":null,"work_id":"3383f635-6bd5-4feb-91f4-9c7e3516488e","year":2005},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.113520Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:4b57e9f91773ef394ec996093770915ac67b70dbf633919a717ac6012684b8be","observation_id":"93026cbf-d874-413a-9bd0-b2651e5c8a3b","resolution":{"observed_at":"2026-08-07T05:43:08.247705Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.228150Z","title":"Translating Principles into Practices of Digital Ethics: Five Risks of Being Unethical","venue":null,"work_id":"1b4a8211-5c62-45bf-b96c-c438d5de7c77","year":2019},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.119100Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:1e89e64fda13f81d598db88084c89b976201d46d204f9802705754363cc61f71","observation_id":"e4ec19f7-68a5-4123-90ff-cb5f9267b597","resolution":{"observed_at":"2026-08-07T05:43:08.232427Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.15217","last_updated":"2023-09-11T17:25:24Z","snapshot_observed_at":"2026-08-17T11:20:55.974248Z","submitted_at":"2023-07-27T22:29:25Z","title":"Open Problems and Fundamental Limitations of Reinforcement Learning from Human Feedback","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.15217","snapshot_observed_at":"2026-08-07T05:43:07.123823Z","title":"Open Problems and Fundamental Limitations of Reinforcement Learning from Human Feedback","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.123823Z"},"links":{"cited_paper":"/paper/2307.15217","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:6bc7188c4a5ca6250aa50fa85c4d446f7191900b14b16431ab4ecc6cdba5a7b2","observation_id":"cab35d18-a447-4352-ac03-fd37e5153ac7","resolution":{"observed_at":"2026-08-07T05:43:07.123823Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.09269","last_updated":"2024-11-07T16:43:01Z","snapshot_observed_at":"2026-08-16T14:18:36.432085Z","submitted_at":"2024-02-14T15:55:30Z","title":"Personalized Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.09269","snapshot_observed_at":"2026-08-07T05:43:07.128902Z","title":"Personalized Large Language Models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.128902Z"},"links":{"cited_paper":"/paper/2402.09269","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:a78dc0b7d3b55c825c51f53ce6d6934e7056b9104032dfd22b9a844ffbfe3854","observation_id":"c5354f90-fad6-4bb2-96cf-a66ca53ebf2d","resolution":{"observed_at":"2026-08-07T05:43:07.128902Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.214393Z","title":"Towards an End -to-End Personal Fine -Tuning Framework for AI Value Alignment","venue":null,"work_id":"298d7b8a-51ae-4263-8c8e-ffe4a0adaecd","year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.133884Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:143363042cbd8d5a6e2636c54e456623dd5dc10c92fabea0f02f7c68a766ef25","observation_id":"49e8007a-9ada-423f-8dbc-f1bc1335f848","resolution":{"observed_at":"2026-08-07T05:43:08.218849Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.200406Z","title":"Safer Agentic AI","venue":null,"work_id":"a47d4a23-5e7c-480c-a21c-e0e93c50984e","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.139059Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:d10c0158c858d1aaabe59f2dcc156ae9cb6f1355ccfe115799d2ce340c25e584","observation_id":"a1f78d6b-2a43-4697-b583-f52683d34c35","resolution":{"observed_at":"2026-08-07T05:43:08.204958Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.186337Z","title":"Introducing the Model Context Protocol","venue":null,"work_id":"9c952372-d8d1-49a3-8321-2ea5769348b5","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.143540Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:8ce99e420026ca2ab5f1150306e569029e7d6d722d89f7707ed3f869c441dc44","observation_id":"ab4fc287-4bcb-4053-b2df-466b9746a33d","resolution":{"observed_at":"2026-08-07T05:43:08.190940Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.172155Z","title":"Deep Reinforcement Learning from Human Preferences","venue":null,"work_id":"3319e636-628e-4f09-9e7c-ff0a59bf88c6","year":2017},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.147943Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:65156a1add6e0aba824751411853d81d5bbfd66fb2e3869aa91a6674a83beedf","observation_id":"0b0b6157-c207-4c1b-b451-66d573f3884b","resolution":{"observed_at":"2026-08-07T05:43:08.176703Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.04175","last_updated":"2024-06-25T18:37:19Z","snapshot_observed_at":"2026-08-16T13:45:28.640764Z","submitted_at":"2024-06-06T15:32:29Z","title":"Confabulation: The Surprising Value of Large Language Model Hallucinations","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.04175","snapshot_observed_at":"2026-08-07T05:43:07.152276Z","title":"Confabulation: The Surprising Value of Large Language Model Hallucinations","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.152276Z"},"links":{"cited_paper":"/paper/2406.04175","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:aae685f89bb51cb7e9b568920984520507e803967a775441c338b547b86e3382","observation_id":"c1a6a662-4abe-47f4-b593-49c0d6ede551","resolution":{"observed_at":"2026-08-07T05:43:07.152276Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.157190Z","title":"Choice Vectors: Streamlining Personal AI Alignment Through Binary Selection","venue":null,"work_id":"1be5211c-f217-4e1a-a39c-d81b9e5497ca","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.157090Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:c41f7146459fd96d92de6ada49b2b70d9f1521fda9a3826445d02cf877efb6da","observation_id":"75f881bc-d2a2-4cc7-a983-2f2387095612","resolution":{"observed_at":"2026-08-07T05:43:08.161821Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.143284Z","title":null,"venue":null,"work_id":"e7267fa3-3ec8-4ed7-be5c-c3da057702c7","year":1923},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.161465Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:70b6415887754802497891c54418deb9717fa54359acfd8be10c8b85394e5746","observation_id":"f8843945-5b73-4fb6-8fdd-6eacf9e521a7","resolution":{"observed_at":"2026-08-07T05:43:08.147667Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.165649Z","title":"Universality of Representation in Biologic al and Artificial Neural Networks","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.165649Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:d46bdd15cdb2a28d9d796e7dd5c15690da8517edf6613577077e1935e867be32","observation_id":"232768fa-1ba1-48b4-b37f-24a77dad24c0","resolution":{"observed_at":"2026-08-07T05:43:07.165649Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.129041Z","title":"The neural bases of cognitive conflict and control in mo ral judgment","venue":null,"work_id":"59b595d4-f6a1-4922-a7e4-bb7d2d224a63","year":2004},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.169921Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:c4cbfeea2c7153207d1c8be2dac7f53d4c1b0b09e12bc9cc6b2beab2d23a9e80","observation_id":"b5015ad3-26fb-48f4-8a6f-ff1489e4c7a9","resolution":{"observed_at":"2026-08-07T05:43:08.133544Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.113746Z","title":"The neural basis of human social values: Evidence from functional MRI","venue":null,"work_id":"0115e97c-bb9f-4421-aab6-1dd1cbd9de5e","year":2009},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.174334Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:cf3dd9d38bd066fed676ad32b1c96e29be6003a64837bf4541b59b65c06e7c19","observation_id":"5f808b93-0ce8-402b-a21e-fc5f0480cb73","resolution":{"observed_at":"2026-08-07T05:43:08.118460Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.098656Z","title":"A Cognitive Theory of Consciousness: The Workspace of the Mind; Cambridge University Press: Cambridge, UK, 1988","venue":null,"work_id":"eb5a53f3-d80a-4612-b3f0-3535426229c5","year":1988},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.179006Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:b78d60c52e3f79c780a2a95c982d38c505b7f611e670cc964b1916e3264922c0","observation_id":"2d831fb6-b9b9-40aa-af80-9648e02cea1c","resolution":{"observed_at":"2026-08-07T05:43:08.103481Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.082983Z","title":"Unified Theories of Cognition; Harvard University Press: Cambridge, MA, USA, 1990","venue":null,"work_id":"b2683094-f6bd-491d-b95c-631fdd1058c0","year":1990},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.183568Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:182cb9b01f5a3658ff0440f5cdaf166bd7ada83fdc7d481146b70d0132914af5","observation_id":"6886e3ee-5210-46b2-b28e-efef9e830027","resolution":{"observed_at":"2026-08-07T05:43:08.088209Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2505.08662","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.772777Z","title":"Revealing economic facts: LLMs know more than they say","venue":null,"work_id":"b29812ab-6be9-4477-95f1-2d525898babc","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.188289Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:ca4e65a00cd129ab474785d553c48ec70a07076fd76c3537d0f853da9295af59","observation_id":"bec20e91-f154-4787-a276-9d29e7bc19d6","resolution":{"observed_at":"2026-08-07T05:43:07.780589Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.01081","last_updated":"2025-04-08T18:38:04Z","snapshot_observed_at":"2026-08-21T11:35:08.605820Z","submitted_at":"2025-04-01T18:00:20Z","title":"ShieldGemma 2: Robust and Tractable Image Content Moderation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.01081","snapshot_observed_at":"2026-08-07T05:43:07.193249Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.193249Z"},"links":{"cited_paper":"/paper/2504.01081","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:0661742b1f9ce9b479c9defe588fc0f9e77e3d74614a2092e7f8826cedbc8267","observation_id":"565c71ff-9edc-481f-a63f-4b03b833470d","resolution":{"observed_at":"2026-08-07T05:43:07.193249Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.068714Z","title":"Superego -Agent LGDemo (Branch: Fastapi_Mcp)","venue":null,"work_id":"cf36d267-18f0-4732-afff-e270a768737a","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.198921Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:5962099522bf5ecac6cc390cdf5820d7998209d4178ca9b1c61f574b72e48219","observation_id":"40c700b9-bd04-4265-a0d1-417a7d2802c3","resolution":{"observed_at":"2026-08-07T05:43:08.073012Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.04249","last_updated":"2024-02-27T04:43:08Z","snapshot_observed_at":"2026-08-16T09:07:20.265665Z","submitted_at":"2024-02-06T18:59:08Z","title":"HarmBench: A Standardized Evaluation Framework for Automated Red Teaming and Robust Refusal","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.04249","snapshot_observed_at":"2026-08-07T05:43:07.203855Z","title":"HarmBench: A standardized evaluation framework for automated red teaming and robust refusal","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.203855Z"},"links":{"cited_paper":"/paper/2402.04249","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:ae36650a9ecb50c62a8688d513b07d743b2fa974bed50b403b2e88777a61e5e3","observation_id":"c55222a1-896d-459a-9f3e-dfd48bfa9033","resolution":{"observed_at":"2026-08-07T05:43:07.203855Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.053714Z","title":"AgentHarm: A benchmark for measuring harmfulness of LLM agents","venue":null,"work_id":"0dc17e2f-f693-4623-968e-0e5a27f52f23","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.209304Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:15f37adb8d9d7044429eaef91f6b6cca8366e9ec23245c64934516063477ce17","observation_id":"218dadac-60a5-4f85-909e-cf5e6a25d19e","resolution":{"observed_at":"2026-08-07T05:43:08.058305Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.038900Z","title":"Do the rewards justify the means? Measuring trade -offs between rewards and ethical behavior in the Machiavelli benchmark","venue":null,"work_id":"ca0c0869-c5b6-4046-95ff-417ba22d957e","year":2023},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.213631Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:36dab579679ead71dba220f06378c9d8bd3c30e1c1f390d4dcc08cc405c3b7bf","observation_id":"0d21dcae-dcab-4951-b7a8-6e0cf673c985","resolution":{"observed_at":"2026-08-07T05:43:08.043769Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.12272","last_updated":"2024-04-18T15:45:27Z","snapshot_observed_at":"2026-08-20T21:06:18.098294Z","submitted_at":"2024-04-18T15:45:27Z","title":"Who Validates the Validators? Aligning LLM-Assisted Evaluation of LLM Outputs with Human Preferences","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.12272","snapshot_observed_at":"2026-08-07T05:43:07.218428Z","title":"Who validates the validators? Aligning LLM-assisted evaluation of LLM outputs with human preferences","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.218428Z"},"links":{"cited_paper":"/paper/2404.12272","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:783f756885f7f9955b4cfe7b122644cae1cda8fc30291c293ab307592f37049f","observation_id":"7d230248-3583-4771-9434-10d2f6bcc2e5","resolution":{"observed_at":"2026-08-07T05:43:07.218428Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.024084Z","title":"Vijil Test Library: Evaluating LLM Trustworthiness Across Eight Dimensions","venue":null,"work_id":"6da1e6c9-9e9e-418f-8b96-d6c0462510c1","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.223161Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:7c81bd369d73c2c63ecdbe9ad3ec7f290b57fe70f56c02e99ab8be45798b70f7","observation_id":"b9a1bd80-b170-4952-ac58-bb35175f738c","resolution":{"observed_at":"2026-08-07T05:43:08.029110Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:08.009296Z","title":"INSPECT: An Extensible Toolkit for AI Behavior Evaluation","venue":null,"work_id":"a21e96d2-fc8f-490a-848c-69a4f04affb3","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.227370Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:51eab60ba3537f4e36bfa53196c17ee905eaad4e2a8c69216745a602c72ddd85","observation_id":"dd65809f-b519-4dd9-aca7-bc452cb6adae","resolution":{"observed_at":"2026-08-07T05:43:08.013655Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.994918Z","title":"Governance in Agentic Workflows: Leveraging LLMs as Oversight Agents","venue":null,"work_id":"b8b4767a-3323-41e3-99a6-276eec52e91e","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.231834Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:e787e564268272d85b378c110cc528540606acf460671a0e064946d86460bf55","observation_id":"5c5cdae8-08c2-4c38-a138-b11566074059","resolution":{"observed_at":"2026-08-07T05:43:07.999467Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.08968","last_updated":"2025-03-03T22:10:04Z","snapshot_observed_at":"2026-08-16T13:10:15.539328Z","submitted_at":"2024-10-11T16:38:01Z","title":"Controllable Safety Alignment: Inference-Time Adaptation to Diverse Safety Requirements","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.08968","snapshot_observed_at":"2026-08-07T05:43:07.235972Z","title":"Controllable Safety Alignment: Inference-Time Adaptation to Diverse Safety Requirements","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.235972Z"},"links":{"cited_paper":"/paper/2410.08968","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:86e006a2c06b0288d2452d0e9fafc5c9e990bacf0d50e9a8cf546dc972ff757a","observation_id":"f34152c1-3bdc-4796-b437-1ecc6eaf0dc8","resolution":{"observed_at":"2026-08-07T05:43:07.235972Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.00166","last_updated":"2026-04-30T20:54:21Z","snapshot_observed_at":"2026-08-17T06:31:05.924897Z","submitted_at":"2025-05-30T19:11:52Z","title":"Disentangled Safety Adapters Enable Efficient Guardrails and Flexible Inference-Time Alignment","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.00166","snapshot_observed_at":"2026-08-07T05:43:07.240615Z","title":"Disentangled Safety Adapters Enable Efficient Guardrails and Flexible Inference-Time Alignment","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.240615Z"},"links":{"cited_paper":"/paper/2506.00166","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:d03177afb54f02765d6371439dea7e43f4f05913e829ee87297fc2f767b2d91f","observation_id":"c56c63fe-b2f0-40db-8b6f-319be09db8c4","resolution":{"observed_at":"2026-08-07T05:43:07.240615Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.11206","last_updated":"2024-01-20T10:41:03Z","snapshot_observed_at":"2026-08-19T06:38:15.589047Z","submitted_at":"2024-01-20T10:41:03Z","title":"InferAligner: Inference-Time Alignment for Harmlessness through Cross-Model Guidance","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.11206","snapshot_observed_at":"2026-08-07T05:43:07.245116Z","title":"InferAligner: Inference -Time Alignment for Harmlessness through Cross-Model Guidance","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.245116Z"},"links":{"cited_paper":"/paper/2401.11206","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:a3cdfd1ce1c832e41755d20d5cbab7d392899b20a9950aa4fc46d44aaff82b4f","observation_id":"01f226e5-346c-4313-bf58-63314a39d76c","resolution":{"observed_at":"2026-08-07T05:43:07.245116Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.01208","last_updated":"2025-06-20T10:54:05Z","snapshot_observed_at":"2026-08-16T10:10:26.213081Z","submitted_at":"2025-02-03T09:59:32Z","title":"On Almost Surely Safe Alignment of Large Language Models at Inference-Time","version":3},"cited_work":{"arxiv_id":"2502.01208","doi":null,"metadata_source":"pith","pith_arxiv_id":"2502.01208","snapshot_observed_at":"2026-08-07T05:43:07.622958Z","title":"On Almost Surely Safe Alignment of Large Language Models at Inference-Time","venue":"cs.LG","work_id":"0ad26be9-3465-4b1c-8504-3dccd9876074","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.249475Z"},"links":{"cited_paper":"/paper/2502.01208","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:2541d561c658d1658d1494ff31095574768670d73639a1f21c5f4a243dadb955","observation_id":"f47fd653-d640-4acc-b755-e93a1f270232","resolution":{"observed_at":"2026-08-07T05:43:07.630186Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.02039","last_updated":"2025-06-02T22:37:06Z","snapshot_observed_at":"2026-08-19T23:29:14.488310Z","submitted_at":"2025-03-03T20:32:05Z","title":"Dynamic Search for Inference-Time Alignment in Diffusion Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.02039","snapshot_observed_at":"2026-08-07T05:43:07.253673Z","title":"Dynamic Search for Inference-Time Alignment in Diffusion Models (DSearch)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.253673Z"},"links":{"cited_paper":"/paper/2503.02039","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:e32c2e1e4ede796417990dd7aaf9a24f388fada8caf438daa58097fe9ac2e5e0","observation_id":"535df10b-43a8-4d3e-afd6-35257bf5828b","resolution":{"observed_at":"2026-08-07T05:43:07.253673Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.18837","last_updated":"2025-01-31T01:09:32Z","snapshot_observed_at":"2026-08-08T23:58:18.293467Z","submitted_at":"2025-01-31T01:09:32Z","title":"Constitutional Classifiers: Defending against Universal Jailbreaks across Thousands of Hours of Red Teaming","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.18837","snapshot_observed_at":"2026-08-07T05:43:07.258162Z","title":"Constit utional Classifiers: Defending against Universal Jailbreaks across Thousands of Hours of Red Teaming","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.258162Z"},"links":{"cited_paper":"/paper/2501.18837","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:d2e8e1b1f75f52aa81e8b056b76e17f403429c2aaefa0c0e81e612d6da38973b","observation_id":"ad083c49-e7a3-48d4-8c78-de4e9196656b","resolution":{"observed_at":"2026-08-07T05:43:07.258162Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2212.08073","last_updated":"2022-12-15T06:19:23Z","snapshot_observed_at":"2026-08-16T03:49:00.703994Z","submitted_at":"2022-12-15T06:19:23Z","title":"Constitutional AI: Harmlessness from AI Feedback","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2212.08073","snapshot_observed_at":"2026-08-07T05:43:07.262772Z","title":"Constitutional AI: Harmlessness from AI Feedback","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.262772Z"},"links":{"cited_paper":"/paper/2212.08073","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:906cffb72bc8dde7ba3321132490937618f1a339f29e2749d093c0788d9ff7f7","observation_id":"b7f7c262-ade5-4959-a0a1-84e014de2c97","resolution":{"observed_at":"2026-08-07T05:43:07.262772Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2204.05862","last_updated":"2022-04-12T15:02:38Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-04-12T15:02:38Z","title":"Training a Helpful and Harmless Assistant with Reinforcement Learning from Human Feedback","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2204.05862","snapshot_observed_at":"2026-08-07T05:43:07.267381Z","title":"T raining a Helpful and Harmless Assistant with Reinforcement Learning from Human Feedback","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.267381Z"},"links":{"cited_paper":"/paper/2204.05862","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:978976252395d53004032478e33a5d9fb59cbe387346fb35c65d01f064f79bcc","observation_id":"7ae553b9-9542-4f09-986a-e2bb36766dc6","resolution":{"observed_at":"2026-08-07T05:43:07.267381Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.980210Z","title":"Monitoring Reasoning Models for Misbehavior and the Risks of Promoting Obfuscation","venue":null,"work_id":"81d3f66c-2d50-4fc0-85ca-76f2cdf707bc","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.272296Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:fd5613d81375b0ce250c6d2f9a192fd2d6cd469897544a719942e64f37ceb9cb","observation_id":"0307194a-196c-4c9a-adca-c6c142cca4c6","resolution":{"observed_at":"2026-08-07T05:43:07.984868Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.06942","last_updated":"2024-07-23T06:47:13Z","snapshot_observed_at":"2026-08-19T14:55:04.920923Z","submitted_at":"2023-12-12T02:34:06Z","title":"AI Control: Improving Safety Despite Intentional Subversion","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.06942","snapshot_observed_at":"2026-08-07T05:43:07.278246Z","title":"AI Control: Improving Safety Despite Intentional Subversion","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.278246Z"},"links":{"cited_paper":"/paper/2312.06942","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:845fd76ca7f68886e8386c1719b31e788e4b41daf0538e818137d50c09026412","observation_id":"171bc1d9-efa4-4d45-85cc-b610962d7f10","resolution":{"observed_at":"2026-08-07T05:43:07.278246Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.965984Z","title":"OpenAI x DFT: The First Moral Graph","venue":null,"work_id":"a49e7d4e-9ee5-412b-b110-1a33995fc6d8","year":2023},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.283191Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:b15f95657307b426575041d483b582785b7d9215c991f888d00cf5ac5e513edc","observation_id":"f222863a-22fb-447c-badb-4127e9498e77","resolution":{"observed_at":"2026-08-07T05:43:07.970374Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.951344Z","title":"Model Integrity","venue":null,"work_id":"e06c8563-8feb-4a7e-a507-c8e5dc561609","year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.287785Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:696838c797422f1b997164ccdc034de02a3e409b58c8eb9af1d6e30216c68b73","observation_id":"8ff4d278-26db-4e2f-a2be-693040d26acf","resolution":{"observed_at":"2026-08-07T05:43:07.956049Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.936583Z","title":"The Global Landscape of AI Ethics Guidelines","venue":null,"work_id":"c36e5993-a8bf-47d8-90ea-2e903e4f3069","year":2019},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.292294Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:1c8c7887020b4bbadff06ac1ce762d3fe295febb682204aaae058392b8f367f3","observation_id":"27bab82c-a0dc-4c0c-ad58-3d049b54f333","resolution":{"observed_at":"2026-08-07T05:43:07.941746Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.921150Z","title":"WhatsApp MCP Exploited: Exfiltrating Your Message History via MCP","venue":null,"work_id":"12f67775-f8ae-4b0e-a319-9ca9c3cae0f0","year":null},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.296584Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:3efb92c0a930d6662c56947fec4b3a2d19d7027e1f3e50885e01eee2e33c25ea","observation_id":"5d78eeaf-2263-4a86-9a54-c9b9b492120d","resolution":{"observed_at":"2026-08-07T05:43:07.926146Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.890964Z","title":"MCP Security Notification: Tool Poisoning Attacks","venue":null,"work_id":"44f1ce5f-dd10-41da-b7f7-e1ed40a33143","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.306568Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:399e36ddcded1d8c42b382fbc613c3caa236b16dac5f39d8b5058459f5fe65d3","observation_id":"48e2fdca-9791-46d1-a005-6bdfaca00cf9","resolution":{"observed_at":"2026-08-07T05:43:07.895246Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.311471Z","title":"Emergent misalignment: Narrow finetuning can produce broadly misaligned LLMs","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.311471Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:38a445150834f788da10f9d1d33d210f6d0f76f13c2b8a92b6b12058e51349e8","observation_id":"20244967-148d-46a2-bb57-cdfbdc0be7ec","resolution":{"observed_at":"2026-08-07T05:43:07.311471Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.875882Z","title":"On Emergent Misalignment","venue":null,"work_id":"fd1df4c5-7e8b-4d6e-aa4a-9710dce04976","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.315617Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:d4f214612415d493035b609ce83d5cff6e86720886419d976c1704e7405029f2","observation_id":"31272682-141c-4702-a791-551e22302342","resolution":{"observed_at":"2026-08-07T05:43:07.880226Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.860100Z","title":"Model Plurality","venue":null,"work_id":"18089c75-6a9a-4013-842e-0cc5d2f5d45d","year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.319836Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:4b21322fbbfdee0584998ca046b248389fbf19b41c6aecc011f6bcd1a5e9c7f6","observation_id":"5f40ad55-bb11-4b30-bc91-b86935cea64a","resolution":{"observed_at":"2026-08-07T05:43:07.864809Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.844679Z","title":"Model Plurality: A Taxonomy for Pluralistic AI","venue":null,"work_id":"3806dc4b-430c-4fb0-9a78-b552894fc032","year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.324191Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:ffbbc0bb9701c97511581cbe838a4e2adc1d2735ccb00b73fe1f7ffec7adfc8b","observation_id":"37621a40-2612-4ca3-812b-e77b1cdf5562","resolution":{"observed_at":"2026-08-07T05:43:07.849749Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.17950","last_updated":"2025-04-24T21:28:16Z","snapshot_observed_at":"2026-08-21T09:04:02.494253Z","submitted_at":"2025-04-24T21:28:16Z","title":"Collaborating Action by Action: A Multi-agent LLM Framework for Embodied Reasoning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.17950","snapshot_observed_at":"2026-08-07T05:43:07.328340Z","title":"Collaborating Action by Action: A Multi-agent LLM Framework for Embodied Reasoning","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.328340Z"},"links":{"cited_paper":"/paper/2504.17950","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:f6a048069e4ba8a1d19c4d3515d438d67c4ca11e0d5ddb2e188ccb196f178c22","observation_id":"6f6e36b4-7098-4f6c-9730-c2637ec23bd0","resolution":{"observed_at":"2026-08-07T05:43:07.328340Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.11581","last_updated":"2025-03-23T13:00:03Z","snapshot_observed_at":"2026-08-20T10:46:11.994812Z","submitted_at":"2024-11-18T13:57:35Z","title":"OASIS: Open Agent Social Interaction Simulations with One Million Agents","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.11581","snapshot_observed_at":"2026-08-07T05:43:07.332394Z","title":"OASIS: Open Agent Social Interaction Simulations with One Million Agents","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.332394Z"},"links":{"cited_paper":"/paper/2411.11581","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:ff2d7c2efe7a5ca3266372c4ae497f880f0bc238e4d786722cb77154ab26a55b","observation_id":"dc7ee6ef-2dca-4126-922c-1c188a5c2c2e","resolution":{"observed_at":"2026-08-07T05:43:07.332394Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.00114","last_updated":"2024-10-31T18:11:22Z","snapshot_observed_at":"2026-08-16T13:04:01.052204Z","submitted_at":"2024-10-31T18:11:22Z","title":"Project Sid: Many-agent simulations toward AI civilization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.00114","snapshot_observed_at":"2026-08-07T05:43:07.337382Z","title":"Project Sid: Many- agent simulations toward AI civilization","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.337382Z"},"links":{"cited_paper":"/paper/2411.00114","citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:011377c28fa1ec6b8cab47a3782eaa0927f12290a75daa24f2b50192f35bbf8f","observation_id":"1b915a17-c0ea-45d4-93aa-e2fcfdfb4806","resolution":{"observed_at":"2026-08-07T05:43:07.337382Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:43:07.905785Z","title":null,"venue":null,"work_id":"43b30316-97dc-4362-a2f2-5b1c307d160d","year":2025},"citing_paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values","version":2},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-07T05:43:07.301344Z"},"links":{"citing_paper":"/paper/2506.13774"},"observation_digest":"sha256:c380ef74697f51aae98f6a81f33089ad91d7ca2d37b5a05b5cec6961f200a2f7","observation_id":"106163d8-0eb8-450e-9b6f-081131884779","resolution":{"observed_at":"2026-08-07T05:43:07.910198Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.13774","last_updated":"2025-08-08T20:29:52Z","latest_version":2,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-13T02:00:24.514117Z","submitted_at":"2025-06-08T20:31:26Z","title":"Personalized Constitutionally-Aligned Agentic Superego: Secure AI Behavior Aligned to Diverse Human Values"},"reference_resolution":{"displayed":50,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":21,"verified_exact":2,"verified_fuzzy":27},"total_outbound_references":50},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"thesis":"As of 21 August 2026, this Paper Citation Record lists 50 of 50 outbound references and 0 inbound Pith citation observations for arXiv:2506.13774."}