From 02a178bcbb1f39095c773204cb94b87a1d7d2cfa Mon Sep 17 00:00:00 2001 From: Advait Bhat Date: Wed, 9 Sep 2026 12:22:30 -0700 Subject: [PATCH 1/4] Approve first publication review batch --- content/publications.json | 456 ++++++++++++++++++++++++++++ maintenance/batch.md | 176 +---------- maintenance/recent-review.md | 222 +------------- maintenance/review.json | 209 +++++++++++-- maintenance/scholar-supplement.json | 17 +- 5 files changed, 660 insertions(+), 420 deletions(-) diff --git a/content/publications.json b/content/publications.json index 92a8393..afe3b3c 100644 --- a/content/publications.json +++ b/content/publications.json @@ -1660,5 +1660,461 @@ "Cristian Danescu-Niculescu-Mizil", "Dan Jurafsky" ] + }, + { + "id": "openalex-w7162039229", + "title": "CandorMD: An AI-Assisted Audio Simulation and Feedback System for Training Clinicians for Medical Error Disclosure", + "authors": "Inna Wanyin Lin and Sahand Sabour and Hong Sng and Maxine Chan and Minlie Huang and Andrew White and Tim Althoff", + "authorNames": [ + "Inna Wanyin Lin", + "Sahand Sabour", + "Hong Sng", + "Maxine Chan", + "Minlie Huang", + "Andrew White", + "Tim Althoff" + ], + "year": 2026, + "venue": "arXiv (Cornell University)", + "doi": "10.48550/arxiv.2605.20701", + "url": "https://doi.org/10.48550/arxiv.2605.20701", + "status": "preprint", + "type": "preprint", + "openalexId": "W7162039229", + "description": "", + "highlight": false, + "award": "", + "pdf": "", + "image": "", + "code": "", + "legacy": {}, + "personIds": [ + "innalin", + "timalthoff" + ], + "reviewedOn": "2026-09-09" + }, + { + "id": "openalex-w7170643724", + "title": "Capable language models can outgrow the benefits of collaboration", + "authors": "Yubin Kim and Ken Gu and Chanwoo Park and Chunjong Park and Samuel Schmidgall and A. Ali Heydari and Yao Yan and Zhihan Zhang and Yuchen Zhuang and Liu Y and Mark Malhotra and Paul Pu Liang and Hae Won Park and Yuzhe Yang and Xuhai Xu and Yilun Du and Shwetak Patel and Tim Althoff and Daniel McDuff and Xin Liu", + "authorNames": [ + "Yubin Kim", + "Ken Gu", + "Chanwoo Park", + "Chunjong Park", + "Samuel Schmidgall", + "A. Ali Heydari", + "Yao Yan", + "Zhihan Zhang", + "Yuchen Zhuang", + "Liu Y", + "Mark Malhotra", + "Paul Pu Liang", + "Hae Won Park", + "Yuzhe Yang", + "Xuhai Xu", + "Yilun Du", + "Shwetak Patel", + "Tim Althoff", + "Daniel McDuff", + "Xin Liu" + ], + "year": 2026, + "venue": "Nature Machine Intelligence", + "doi": "10.1038/s42256-026-01268-y", + "url": "https://doi.org/10.1038/s42256-026-01268-y", + "status": "published", + "type": "article", + "openalexId": "W7170643724", + "description": "", + "highlight": false, + "award": "", + "pdf": "", + "image": "", + "code": "", + "legacy": {}, + "personIds": [ + "kengu", + "timalthoff" + ], + "reviewedOn": "2026-09-09" + }, + { + "id": "openalex-w4416863183", + "title": "How Conversational Structure and Style Shape Online Community Experiences", + "authors": "Galen Weld and Carl Pearson and Bradley Spahn and Tim Althoff and Amy X. Zhang and Sanjay Kairam", + "authorNames": [ + "Galen Weld", + "Carl Pearson", + "Bradley Spahn", + "Tim Althoff", + "Amy X. Zhang", + "Sanjay Kairam" + ], + "year": 2026, + "venue": "Proceedings of the International AAAI Conference on Web and Social Media", + "doi": "10.1609/icwsm.v20i1.42762", + "url": "https://doi.org/10.1609/icwsm.v20i1.42762", + "status": "published", + "type": "other", + "openalexId": "W4416863183", + "description": "", + "highlight": false, + "award": "", + "pdf": "", + "image": "", + "code": "", + "legacy": {}, + "personIds": [ + "galenweld", + "timalthoff" + ], + "reviewedOn": "2026-09-09" + }, + { + "id": "openalex-w7166852917", + "title": "Inferring Events from Time Series using Language Models", + "authors": "Mingtian Tan and Mike A. Merrill and Zachary Gottesman and Tim Althoff and David Evans and Thomas Hartvigsen", + "authorNames": [ + "Mingtian Tan", + "Mike A. Merrill", + "Zachary Gottesman", + "Tim Althoff", + "David Evans", + "Thomas Hartvigsen" + ], + "year": 2026, + "venue": "Proceedings of the 64th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)", + "doi": "10.18653/v1/2026.acl-long.157", + "url": "https://aclanthology.org/2026.acl-long.157/", + "status": "published", + "type": "article", + "openalexId": "W7166852917", + "description": "", + "highlight": false, + "award": "", + "pdf": "", + "image": "", + "code": "", + "legacy": {}, + "personIds": [ + "mikemerrill", + "timalthoff" + ], + "reviewedOn": "2026-09-09" + }, + { + "id": "manual-synthworlds-2026", + "title": "SynthWorlds: Controlled Parallel Worlds for Disentangling Reasoning and Knowledge in Language Models", + "authors": "Ken Gu and Advait Bhat and Mike Merrill and Robert West and Xin Liu and Daniel McDuff and Tim Althoff", + "authorNames": [ + "Ken Gu", + "Advait Bhat", + "Mike Merrill", + "Robert West", + "Xin Liu", + "Daniel McDuff", + "Tim Althoff" + ], + "year": 2026, + "venue": "ICLR 2026", + "url": "https://proceedings.iclr.cc/paper_files/paper/2026/hash/b213d870740582dd6af77bbdaed900c9-Abstract-Conference.html", + "status": "published", + "type": "conference", + "description": "", + "highlight": false, + "award": "", + "pdf": "", + "image": "", + "code": "", + "legacy": {}, + "personIds": [ + "kengu", + "advaitbhat", + "mikemerrill", + "timalthoff" + ], + "reviewedOn": "2026-09-09" + }, + { + "id": "openalex-w7122726170", + "title": "Transforming wearable data into personal health insights using large language model agents", + "authors": "Mike A. Merrill and Akshay Paruchuri and Naghmeh Rezaei and Geza Kovacs and Javier Perez and Yun Liu and Erik Schenck and Nova Hammerquist and Jake Sunshine and Shyam Tailor and Kumar Ayush and Hao-Wei Su and Qian He and Cory Y. McLean and Mark Malhotra and Shwetak Patel and Jiening Zhan and Tim Althoff and Daniel McDuff and Yi Liu", + "authorNames": [ + "Mike A. Merrill", + "Akshay Paruchuri", + "Naghmeh Rezaei", + "Geza Kovacs", + "Javier Perez", + "Yun Liu", + "Erik Schenck", + "Nova Hammerquist", + "Jake Sunshine", + "Shyam Tailor", + "Kumar Ayush", + "Hao-Wei Su", + "Qian He", + "Cory Y. McLean", + "Mark Malhotra", + "Shwetak Patel", + "Jiening Zhan", + "Tim Althoff", + "Daniel McDuff", + "Yi Liu" + ], + "year": 2026, + "venue": "Nature Communications", + "doi": "10.1038/s41467-025-67922-y", + "url": "https://doi.org/10.1038/s41467-025-67922-y", + "status": "published", + "type": "article", + "openalexId": "W7122726170", + "description": "", + "highlight": false, + "award": "", + "pdf": "", + "image": "", + "code": "", + "legacy": {}, + "personIds": [ + "mikemerrill", + "timalthoff" + ], + "reviewedOn": "2026-09-09" + }, + { + "id": "openalex-w4407425790", + "title": "Human Decision-making is Susceptible to AI-driven Manipulation", + "authors": "Sahand Sabour and June M. Liu and Siyang Liu and Yao, Chris Z. and Shiyao Cui and Xuanming Zhang and Wen Zhang and Yaru Cao and Advait Bhat and Jinan Guan and Wei Wu and Rada Mihalcea and Wang, Hongning and Tim Althoff and Tatia M. C. Lee and Minlie Huang", + "authorNames": [ + "Sahand Sabour", + "June M. Liu", + "Siyang Liu", + "Yao, Chris Z.", + "Shiyao Cui", + "Xuanming Zhang", + "Wen Zhang", + "Yaru Cao", + "Advait Bhat", + "Jinan Guan", + "Wei Wu", + "Rada Mihalcea", + "Wang, Hongning", + "Tim Althoff", + "Tatia M. C. Lee", + "Minlie Huang" + ], + "year": 2025, + "venue": "arXiv (Cornell University)", + "doi": "10.48550/arxiv.2502.07663", + "url": "http://arxiv.org/abs/2502.07663", + "status": "preprint", + "type": "preprint", + "openalexId": "W4407425790", + "description": "", + "highlight": false, + "award": "", + "pdf": "", + "image": "", + "code": "", + "legacy": {}, + "personIds": [ + "advaitbhat", + "timalthoff" + ], + "reviewedOn": "2026-09-09" + }, + { + "id": "openalex-w4416076746", + "title": "LSM-2: Learning from Incomplete Wearable Sensor Data", + "authors": "Maxwell A. Xu and Girish Narayanswamy and Kumar Ayush and Dimitris Spathis and Shun Liao and Shyam A. Tailor and Ahmed Hosny Saleh Metwally and A. Ali Heydari and Yuwei Zhang and Jake Garrison and Samy Abdel-Ghaffar and Xuhai Xu and Ken Gu and Jacob E. Sunshine and Ming‐Zher Poh and Yun Liu and Tim Althoff and Shrikanth Narayanan and Pushmeet Kohli and Mark Malhotra and Shwetak Patel and Yuzhe Yang and James M. Rehg and Xin Liu and Daniel McDuff", + "authorNames": [ + "Maxwell A. Xu", + "Girish Narayanswamy", + "Kumar Ayush", + "Dimitris Spathis", + "Shun Liao", + "Shyam A. Tailor", + "Ahmed Hosny Saleh Metwally", + "A. Ali Heydari", + "Yuwei Zhang", + "Jake Garrison", + "Samy Abdel-Ghaffar", + "Xuhai Xu", + "Ken Gu", + "Jacob E. Sunshine", + "Ming‐Zher Poh", + "Yun Liu", + "Tim Althoff", + "Shrikanth Narayanan", + "Pushmeet Kohli", + "Mark Malhotra", + "Shwetak Patel", + "Yuzhe Yang", + "James M. Rehg", + "Xin Liu", + "Daniel McDuff" + ], + "year": 2025, + "venue": "arXiv (Cornell University)", + "doi": "10.48550/arxiv.2506.05321", + "url": "http://arxiv.org/abs/2506.05321", + "status": "preprint", + "type": "preprint", + "openalexId": "W4416076746", + "description": "", + "highlight": false, + "award": "", + "pdf": "", + "image": "", + "code": "", + "legacy": {}, + "personIds": [ + "kengu", + "timalthoff" + ], + "reviewedOn": "2026-09-09" + }, + { + "id": "openalex-w4415250179", + "title": "Perceptions of Moderators as a Large-Scale Measure of Online Community Governance", + "authors": "Galen Weld and Leon Leibmann and Amy X. Zhang and Tim Althoff", + "authorNames": [ + "Galen Weld", + "Leon Leibmann", + "Amy X. Zhang", + "Tim Althoff" + ], + "year": 2025, + "venue": "Proceedings of the ACM on Human-Computer Interaction", + "doi": "10.1145/3757644", + "url": "https://doi.org/10.1145/3757644", + "status": "published", + "type": "article", + "openalexId": "W4415250179", + "description": "", + "highlight": false, + "award": "", + "pdf": "", + "image": "", + "code": "", + "legacy": {}, + "personIds": [ + "galenweld", + "timalthoff" + ], + "reviewedOn": "2026-09-09" + }, + { + "id": "openalex-w4417255562", + "title": "RADAR: Benchmarking Language Models on Imperfect Tabular Data", + "authors": "Ken Gu and Zhihan Zhang and Kate Lin and Yuwei Zhang and Akshay Paruchuri and Hong Yu and Mehran Kazemi and Kumar Ayush and A. Ali Heydari and Maxwell A. Xu and Girish Narayanswamy and Yun Liu and Ming-Zher Poh and Yuzhe Yang and Mark Malhotra and Shwetak Patel and Hamid Palangi and Xuhai Xu and Daniel McDuff and Tim Althoff and Xin Liu", + "authorNames": [ + "Ken Gu", + "Zhihan Zhang", + "Kate Lin", + "Yuwei Zhang", + "Akshay Paruchuri", + "Hong Yu", + "Mehran Kazemi", + "Kumar Ayush", + "A. Ali Heydari", + "Maxwell A. Xu", + "Girish Narayanswamy", + "Yun Liu", + "Ming-Zher Poh", + "Yuzhe Yang", + "Mark Malhotra", + "Shwetak Patel", + "Hamid Palangi", + "Xuhai Xu", + "Daniel McDuff", + "Tim Althoff", + "Xin Liu" + ], + "year": 2025, + "venue": "NeurIPS", + "doi": "10.52202/085713-3699", + "url": "https://doi.org/10.52202/085713-3699", + "status": "published", + "type": "article", + "openalexId": "W4417255562", + "description": "", + "highlight": false, + "award": "", + "pdf": "", + "image": "", + "code": "", + "legacy": {}, + "personIds": [ + "kengu", + "timalthoff" + ], + "reviewedOn": "2026-09-09" + }, + { + "id": "openalex-w4411121032", + "title": "Reddit Rules and Rulers: Quantifying the Link Between Rules and Perceptions of Governance Across Thousands of Communities", + "authors": "Leon Leibmann and Galen Weld and Amy X. Zhang and Tim Althoff", + "authorNames": [ + "Leon Leibmann", + "Galen Weld", + "Amy X. Zhang", + "Tim Althoff" + ], + "year": 2025, + "venue": "Proceedings of the International AAAI Conference on Web and Social Media", + "doi": "10.1609/icwsm.v19i1.35863", + "url": "https://doi.org/10.1609/icwsm.v19i1.35863", + "status": "published", + "type": "other", + "openalexId": "W4411121032", + "description": "", + "highlight": false, + "award": "", + "pdf": "", + "image": "", + "code": "", + "legacy": {}, + "personIds": [ + "galenweld", + "timalthoff" + ], + "reviewedOn": "2026-09-09" + }, + { + "id": "openalex-w4417141639", + "title": "Self-Improving VLM Judges Without Human Annotations", + "authors": "Inna Wanyin Lin and Yushi Hu and Shuyue Stella Li and Scott Geng and Pang Wei Koh and Luke Zettlemoyer and Tim Althoff and Marjan Ghazvininejad", + "authorNames": [ + "Inna Wanyin Lin", + "Yushi Hu", + "Shuyue Stella Li", + "Scott Geng", + "Pang Wei Koh", + "Luke Zettlemoyer", + "Tim Althoff", + "Marjan Ghazvininejad" + ], + "year": 2025, + "venue": "arXiv (Cornell University)", + "doi": "10.48550/arxiv.2512.05145", + "url": "http://arxiv.org/abs/2512.05145", + "status": "preprint", + "type": "preprint", + "openalexId": "W4417141639", + "description": "", + "highlight": false, + "award": "", + "pdf": "", + "image": "", + "code": "", + "legacy": {}, + "personIds": [ + "innalin", + "timalthoff" + ], + "reviewedOn": "2026-09-09" } ] diff --git a/maintenance/batch.md b/maintenance/batch.md index 84be02a..dc8c09e 100644 --- a/maintenance/batch.md +++ b/maintenance/batch.md @@ -4,184 +4,10 @@ Edit decisions with `python3 scripts/review.py`; approval changes local content meets-rule: 0 -needs-membership-review: 57 +needs-membership-review: 46 does-not-meet-rule: 264 -## openalex-w4416863183 - -How Conversational Structure and Style Shape Online Community Experiences - -Source: https://openalex.org/W4416863183 - -Matched people (lab relevance still needs review): galenweld, timalthoff - -Target: new record - -Proposed fields: title, authors, authorNames, year, venue, doi, url, status, openalexId, type - -Lab relevance: needs-membership-review — Enough lab identities match, but publication-time membership needs confirmation. - -Author identity needs per-paper verification (mixed OpenAlex profile): galenweld - -## openalex-w7122726170 - -Transforming wearable data into personal health insights using large language model agents - -Source: https://openalex.org/W7122726170 - -Matched people (lab relevance still needs review): mikemerrill, timalthoff - -Target: new record - -Proposed fields: title, authors, authorNames, year, venue, doi, url, status, openalexId, type - -Lab relevance: needs-membership-review — Enough lab identities match, but publication-time membership needs confirmation. - -Other versions grouped here: openalex-w4399596965 - -## openalex-w7162039229 - -CandorMD: An AI-Assisted Audio Simulation and Feedback System for Training Clinicians for Medical Error Disclosure - -Source: https://openalex.org/W7162039229 - -Matched people (lab relevance still needs review): innalin, timalthoff - -Target: new record - -Proposed fields: title, authors, authorNames, year, venue, doi, url, status, openalexId, type - -Lab relevance: needs-membership-review — Enough lab identities match, but publication-time membership needs confirmation. - -Other versions grouped here: openalex-w7162150239 - -## openalex-w7166852917 - -Inferring Events from Time Series using Language Models - -Source: https://openalex.org/W7166852917 - -Matched people (lab relevance still needs review): mikemerrill, timalthoff - -Target: new record - -Proposed fields: title, authors, authorNames, year, venue, doi, url, status, openalexId, type - -Lab relevance: needs-membership-review — Enough lab identities match, but publication-time membership needs confirmation. - -## openalex-w7170643724 - -Capable language models can outgrow the benefits of collaboration - -Source: https://openalex.org/W7170643724 - -Matched people (lab relevance still needs review): kengu, timalthoff - -Target: new record - -Proposed fields: title, authors, authorNames, year, venue, doi, url, status, openalexId, type - -Lab relevance: needs-membership-review — Enough lab identities match, but publication-time membership needs confirmation. - -Author identity needs per-paper verification (mixed OpenAlex profile): kengu - -Other versions grouped here: openalex-w7125481708 - -## openalex-w4407425790 - -Human Decision-making is Susceptible to AI-driven Manipulation - -Source: https://openalex.org/W4407425790 - -Matched people (lab relevance still needs review): advaitbhat, timalthoff - -Target: new record - -Proposed fields: title, authors, authorNames, year, venue, doi, url, status, openalexId, type - -Lab relevance: needs-membership-review — Enough lab identities match, but publication-time membership needs confirmation. - -## openalex-w4411121032 - -Reddit Rules and Rulers: Quantifying the Link Between Rules and Perceptions of Governance Across Thousands of Communities - -Source: https://openalex.org/W4411121032 - -Matched people (lab relevance still needs review): galenweld, timalthoff - -Target: new record - -Proposed fields: title, authors, authorNames, year, venue, doi, url, status, openalexId, type - -Lab relevance: needs-membership-review — Enough lab identities match, but publication-time membership needs confirmation. - -Author identity needs per-paper verification (mixed OpenAlex profile): galenweld - -Other versions grouped here: openalex-w4406840671 - -## openalex-w4415250179 - -Perceptions of Moderators as a Large-Scale Measure of Online Community Governance - -Source: https://openalex.org/W4415250179 - -Matched people (lab relevance still needs review): galenweld, timalthoff - -Target: new record - -Proposed fields: title, authors, authorNames, year, venue, doi, url, status, openalexId, type - -Lab relevance: needs-membership-review — Enough lab identities match, but publication-time membership needs confirmation. - -Author identity needs per-paper verification (mixed OpenAlex profile): galenweld - -Other versions grouped here: openalex-w4391420916 - -## openalex-w4416076746 - -LSM-2: Learning from Incomplete Wearable Sensor Data - -Source: https://openalex.org/W4416076746 - -Matched people (lab relevance still needs review): kengu, timalthoff - -Target: new record - -Proposed fields: title, authors, authorNames, year, venue, doi, url, status, openalexId, type - -Lab relevance: needs-membership-review — Enough lab identities match, but publication-time membership needs confirmation. - -Author identity needs per-paper verification (mixed OpenAlex profile): kengu - -## openalex-w4417141639 - -Self-Improving VLM Judges Without Human Annotations - -Source: https://openalex.org/W4417141639 - -Matched people (lab relevance still needs review): innalin, timalthoff - -Target: new record - -Proposed fields: title, authors, authorNames, year, venue, doi, url, status, openalexId, type - -Lab relevance: needs-membership-review — Enough lab identities match, but publication-time membership needs confirmation. - -## openalex-w4417255562 - -RADAR: Benchmarking Language Models on Imperfect Tabular Data - -Source: https://openalex.org/W4417255562 - -Matched people (lab relevance still needs review): kengu, timalthoff - -Target: new record - -Proposed fields: title, authors, authorNames, year, venue, doi, url, status, openalexId, type - -Lab relevance: needs-membership-review — Enough lab identities match, but publication-time membership needs confirmation. - ## openalex-w3199811133 Making Online Communities ‘Better’: A Taxonomy of Community Values on Reddit diff --git a/maintenance/recent-review.md b/maintenance/recent-review.md index 25b57e0..70df16e 100644 --- a/maintenance/recent-review.md +++ b/maintenance/recent-review.md @@ -1,227 +1,7 @@ # Recent publication review -12 distinct papers from 2025 onward. Choose **include**, **exclude**, or **unsure**. Reply in chat with the numbered decisions or candidate IDs. Nothing here has been approved for the website. +0 distinct papers from 2025 onward. Choose **include**, **exclude**, or **unsure**. Reply in chat with the numbered decisions or candidate IDs. Nothing here has been approved for the website. Rule: Tim must be an author, together with at least one current or past lab member. Confirm membership at the time of the work and whether it belongs to the lab. Duplicate versions appear under one item; source records remain available. Google Scholar profiles last checked: 2026-09-05. See `_planning/SCHOLAR_LATEST_REVIEW.md` for coverage and remaining uncertainties. - -## 1. CandorMD: An AI-Assisted Audio Simulation and Feedback System for Training Clinicians for Medical Error Disclosure - -**Year:** 2026 · **Venue:** arXiv (Cornell University) - -**Matched lab coauthors:** Inna Lin, Tim Althoff - -**Full author list:** Inna Wanyin Lin; Sahand Sabour; Hong Sng; Maxine Chan; Minlie Huang; Andrew White; Tim Althoff - -[Paper](https://doi.org/10.48550/arxiv.2605.20701) · [Source record](https://openalex.org/W7162039229) - -Other version: [CandorMD: An AI-Assisted Audio Simulation and Feedback System for Training Clinicians for Medical Error Disclosure](https://arxiv.org/abs/2605.20701) (2026). - -**Needs checking:** lab attribution and membership dates. - -**Decision:** Pending - -Candidate ID: `openalex-w7162039229` - -## 2. Capable language models can outgrow the benefits of collaboration - -**Year:** 2026 · **Venue:** Nature Machine Intelligence - -**Matched lab coauthors:** Ken Gu, Tim Althoff - -**Full author list:** Yubin Kim; Ken Gu; Chanwoo Park; Chunjong Park; Samuel Schmidgall; A. Ali Heydari; Yao Yan; Zhihan Zhang; Yuchen Zhuang; Liu Y; Mark Malhotra; Paul Pu Liang; Hae Won Park; Yuzhe Yang; Xuhai Xu; Yilun Du; Shwetak Patel; Tim Althoff; Daniel McDuff; Xin Liu - -[Paper](https://doi.org/10.1038/s42256-026-01268-y) · [Source record](https://openalex.org/W7170643724) - -Other version: [Towards a Science of Scaling Agent Systems](https://doi.org/10.21203/rs.3.rs-8414536/v1) (2026). - -**Check author identity:** Ken Gu. - -**Needs checking:** lab attribution and membership dates. - -**Decision:** Pending - -Candidate ID: `openalex-w7170643724` - -## 3. How Conversational Structure and Style Shape Online Community Experiences - -**Year:** 2026 · **Venue:** Proceedings of the International AAAI Conference on Web and Social Media - -**Matched lab coauthors:** Galen Weld, Tim Althoff - -**Full author list:** Galen Weld; Carl Pearson; Bradley Spahn; Tim Althoff; Amy X. Zhang; Sanjay Kairam - -[Paper](https://doi.org/10.1609/icwsm.v20i1.42762) · [Source record](https://openalex.org/W4416863183) - -**Check author identity:** Galen Weld. - -**Needs checking:** lab attribution and membership dates. - -**Decision:** Pending - -Candidate ID: `openalex-w4416863183` - -## 4. Inferring Events from Time Series using Language Models - -**Year:** 2026 · **Venue:** Proceedings of the 64th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers) - -**Matched lab coauthors:** Mike Merrill, Tim Althoff - -**Full author list:** Mingtian Tan; Mike A. Merrill; Zachary Gottesman; Tim Althoff; David Evans; Thomas Hartvigsen - -[Paper](https://aclanthology.org/2026.acl-long.157/) · [Source record](https://openalex.org/W7166852917) - -**Reviewer-corrected metadata:** type, url, venue. - -**Needs checking:** lab attribution and membership dates. - -**Decision:** Pending - -Candidate ID: `openalex-w7166852917` - -## 5. SynthWorlds: Controlled Parallel Worlds for Disentangling Reasoning and Knowledge in Language Models - -**Year:** 2026 · **Venue:** ICLR 2026 - -**Matched lab coauthors:** Ken Gu, Advait Bhat, Mike Merrill, Tim Althoff - -**Full author list:** Ken Gu; Advait Bhat; Mike Merrill; Robert West; Xin Liu; Daniel McDuff; Tim Althoff - -[Paper](https://proceedings.iclr.cc/paper_files/paper/2026/hash/b213d870740582dd6af77bbdaed900c9-Abstract-Conference.html) · [Source record](https://proceedings.iclr.cc/paper_files/paper/2026/hash/b213d870740582dd6af77bbdaed900c9-Abstract-Conference.html) - -Found on Tim and Advait's Google Scholar profiles; title and full authors verified in ICLR proceedings. OpenAlex record not found. Review in chat; add through the manual content workflow after approval. - -**Needs checking:** lab attribution and membership dates. - -**Decision:** Pending - -Candidate ID: `manual-synthworlds-2026` - -## 6. Transforming wearable data into personal health insights using large language model agents - -**Year:** 2026 · **Venue:** Nature Communications - -**Matched lab coauthors:** Mike Merrill, Tim Althoff - -**Full author list:** Mike A. Merrill; Akshay Paruchuri; Naghmeh Rezaei; Geza Kovacs; Javier Perez; Yun Liu; Erik Schenck; Nova Hammerquist; Jake Sunshine; Shyam Tailor; Kumar Ayush; Hao-Wei Su; Qian He; Cory Y. McLean; Mark Malhotra; Shwetak Patel; Jiening Zhan; Tim Althoff; Daniel McDuff; Yi Liu - -[Paper](https://doi.org/10.1038/s41467-025-67922-y) · [Source record](https://openalex.org/W7122726170) - -Other version: [Transforming Wearable Data into Personal Health Insights using Large Language Model Agents](http://arxiv.org/abs/2406.06464) (2024). - -**Needs checking:** lab attribution and membership dates. - -**Decision:** Pending - -Candidate ID: `openalex-w7122726170` - -## 7. Human Decision-making is Susceptible to AI-driven Manipulation - -**Year:** 2025 · **Venue:** arXiv (Cornell University) - -**Matched lab coauthors:** Advait Bhat, Tim Althoff - -**Full author list:** Sahand Sabour; June M. Liu; Siyang Liu; Yao, Chris Z.; Shiyao Cui; Xuanming Zhang; Wen Zhang; Yaru Cao; Advait Bhat; Jinan Guan; Wei Wu; Rada Mihalcea; Wang, Hongning; Tim Althoff; Tatia M. C. Lee; Minlie Huang - -[Paper](http://arxiv.org/abs/2502.07663) · [Source record](https://openalex.org/W4407425790) - -**Needs checking:** lab attribution and membership dates. - -**Decision:** Pending - -Candidate ID: `openalex-w4407425790` - -## 8. LSM-2: Learning from Incomplete Wearable Sensor Data - -**Year:** 2025 · **Venue:** arXiv (Cornell University) - -**Matched lab coauthors:** Ken Gu, Tim Althoff - -**Full author list:** Maxwell A. Xu; Girish Narayanswamy; Kumar Ayush; Dimitris Spathis; Shun Liao; Shyam A. Tailor; Ahmed Hosny Saleh Metwally; A. Ali Heydari; Yuwei Zhang; Jake Garrison; Samy Abdel-Ghaffar; Xuhai Xu; Ken Gu; Jacob E. Sunshine; Ming‐Zher Poh; Yun Liu; Tim Althoff; Shrikanth Narayanan; Pushmeet Kohli; Mark Malhotra; Shwetak Patel; Yuzhe Yang; James M. Rehg; Xin Liu; Daniel McDuff - -[Paper](http://arxiv.org/abs/2506.05321) · [Source record](https://openalex.org/W4416076746) - -**Check author identity:** Ken Gu. - -**Needs checking:** lab attribution and membership dates. - -**Decision:** Pending - -Candidate ID: `openalex-w4416076746` - -## 9. Perceptions of Moderators as a Large-Scale Measure of Online Community Governance - -**Year:** 2025 · **Venue:** Proceedings of the ACM on Human-Computer Interaction - -**Matched lab coauthors:** Galen Weld, Tim Althoff - -**Full author list:** Galen Weld; Leon Leibmann; Amy X. Zhang; Tim Althoff - -[Paper](https://doi.org/10.1145/3757644) · [Source record](https://openalex.org/W4415250179) - -Other version: [Perceptions of Moderators as a Large-Scale Measure of Online Community Governance](http://arxiv.org/abs/2401.16610) (2024). - -**Check author identity:** Galen Weld. - -**Needs checking:** lab attribution and membership dates. - -**Decision:** Pending - -Candidate ID: `openalex-w4415250179` - -## 10. RADAR: Benchmarking Language Models on Imperfect Tabular Data - -**Year:** 2025 · **Venue:** NeurIPS - -**Matched lab coauthors:** Ken Gu, Tim Althoff - -**Full author list:** Ken Gu; Zhihan Zhang; Kate Lin; Yuwei Zhang; Akshay Paruchuri; Hong Yu; Mehran Kazemi; Kumar Ayush; A. Ali Heydari; Maxwell A. Xu; Girish Narayanswamy; Yun Liu; Ming-Zher Poh; Yuzhe Yang; Mark Malhotra; Shwetak Patel; Hamid Palangi; Xuhai Xu; Daniel McDuff; Tim Althoff; Xin Liu - -[Paper](https://doi.org/10.52202/085713-3699) · [Source record](https://openalex.org/W4417255562) - -**Reviewer-corrected metadata:** authorNames, authors, type, venue. - -**Year discrepancy:** Scholar displays 2026; the conference record is NeurIPS 2025. Keep the conference year unless review establishes otherwise. - -**Needs checking:** lab attribution and membership dates. - -**Decision:** Pending - -Candidate ID: `openalex-w4417255562` - -## 11. Reddit Rules and Rulers: Quantifying the Link Between Rules and Perceptions of Governance Across Thousands of Communities - -**Year:** 2025 · **Venue:** Proceedings of the International AAAI Conference on Web and Social Media - -**Matched lab coauthors:** Galen Weld, Tim Althoff - -**Full author list:** Leon Leibmann; Galen Weld; Amy X. Zhang; Tim Althoff - -[Paper](https://doi.org/10.1609/icwsm.v19i1.35863) · [Source record](https://openalex.org/W4411121032) - -Other version: [Reddit Rules and Rulers: Quantifying the Link Between Rules and Perceptions of Governance across Thousands of Communities](http://arxiv.org/abs/2501.14163) (2025). - -**Check author identity:** Galen Weld. - -**Needs checking:** lab attribution and membership dates. - -**Decision:** Pending - -Candidate ID: `openalex-w4411121032` - -## 12. Self-Improving VLM Judges Without Human Annotations - -**Year:** 2025 · **Venue:** arXiv (Cornell University) - -**Matched lab coauthors:** Inna Lin, Tim Althoff - -**Full author list:** Inna Wanyin Lin; Yushi Hu; Shuyue Stella Li; Scott Geng; Pang Wei Koh; Luke Zettlemoyer; Tim Althoff; Marjan Ghazvininejad - -[Paper](http://arxiv.org/abs/2512.05145) · [Source record](https://openalex.org/W4417141639) - -**Needs checking:** lab attribution and membership dates. - -**Decision:** Pending - -Candidate ID: `openalex-w4417141639` diff --git a/maintenance/review.json b/maintenance/review.json index 63decd6..9275d46 100644 --- a/maintenance/review.json +++ b/maintenance/review.json @@ -20950,7 +20950,7 @@ }, { "id": "openalex-w4407425790", - "status": "pending", + "status": "accepted", "firstSeen": "2026-09-05", "fingerprint": "6dccd3e3239b800d2116c7c98a6db1072aaeb25babebec6b9216016926050973", "sourceUrl": "https://openalex.org/W4407425790", @@ -20988,7 +20988,7 @@ "advaitbhat", "timalthoff" ], - "targetId": null, + "targetId": "openalex-w4407425790", "base": {}, "changes": { "title": "Human Decision-making is Susceptible to AI-driven Manipulation", @@ -21037,6 +21037,21 @@ ], "status": "needs-membership-review", "reason": "Enough lab identities match, but publication-time membership needs confirmation." + }, + "reviewedOn": "2026-09-09", + "policyReview": { + "policyVersion": 2, + "labRelevanceStatus": "needs-membership-review", + "personIds": [ + "advaitbhat", + "timalthoff" + ], + "evidenceUrls": [ + "https://arxiv.org/abs/2502.07663", + "https://behavioral-data.github.io/team/" + ], + "overrideReason": "Primary arXiv record verifies Tim Althoff and Advait Bhat; the maintainer approved preprint inclusion and the public lab roster verifies Advait's lab membership, whose dates are not encoded.", + "reviewedOn": "2026-09-09" } }, { @@ -21547,7 +21562,7 @@ }, { "id": "openalex-w4411121032", - "status": "pending", + "status": "accepted", "firstSeen": "2026-09-05", "fingerprint": "1139ecdfcb594be91dad8844e762f15a7d40814c4548581637f5b00e3634b01e", "sourceUrl": "https://openalex.org/W4411121032", @@ -21573,7 +21588,7 @@ "galenweld", "timalthoff" ], - "targetId": null, + "targetId": "openalex-w4411121032", "base": {}, "changes": { "title": "Reddit Rules and Rulers: Quantifying the Link Between Rules and Perceptions of Governance Across Thousands of Communities", @@ -21613,6 +21628,21 @@ "status": "needs-membership-review", "reason": "Enough lab identities match, but publication-time membership needs confirmation." }, + "reviewedOn": "2026-09-09", + "policyReview": { + "policyVersion": 2, + "labRelevanceStatus": "needs-membership-review", + "personIds": [ + "galenweld", + "timalthoff" + ], + "evidenceUrls": [ + "https://ojs.aaai.org/index.php/ICWSM/article/view/35863", + "https://behavioral-data.github.io/moderator_perceptions_public/" + ], + "overrideReason": "Official ICWSM record verifies Tim Althoff and Galen Weld; the maintainer approved inclusion and the lab project site establishes the work as a Behavioral Data Science publication.", + "reviewedOn": "2026-09-09" + }, "alternateRecordIds": [ "openalex-w4406840671" ] @@ -22666,7 +22696,7 @@ }, { "id": "openalex-w4415250179", - "status": "pending", + "status": "accepted", "firstSeen": "2026-09-05", "fingerprint": "6dc03dae5ec2c69204b8e9fbb7ee7b9295bd36566cd6a04c20960b1e25ccbd30", "sourceUrl": "https://openalex.org/W4415250179", @@ -22692,7 +22722,7 @@ "galenweld", "timalthoff" ], - "targetId": null, + "targetId": "openalex-w4415250179", "base": {}, "changes": { "title": "Perceptions of Moderators as a Large-Scale Measure of Online Community Governance", @@ -22732,6 +22762,21 @@ "status": "needs-membership-review", "reason": "Enough lab identities match, but publication-time membership needs confirmation." }, + "reviewedOn": "2026-09-09", + "policyReview": { + "policyVersion": 2, + "labRelevanceStatus": "needs-membership-review", + "personIds": [ + "galenweld", + "timalthoff" + ], + "evidenceUrls": [ + "https://doi.org/10.1145/3757644", + "https://behavioral-data.github.io/moderator_discourse_public/" + ], + "overrideReason": "Publisher-linked lab project record verifies Tim Althoff and Galen Weld; the maintainer approved inclusion and the lab project site establishes the work as a Behavioral Data Science publication.", + "reviewedOn": "2026-09-09" + }, "alternateRecordIds": [ "openalex-w4391420916" ] @@ -23547,7 +23592,7 @@ }, { "id": "openalex-w4416076746", - "status": "pending", + "status": "accepted", "firstSeen": "2026-09-05", "fingerprint": "fcd987d92d609d4ac3e64e687a2378db76b8afa6214a85173e1ec66ead968cc3", "sourceUrl": "https://openalex.org/W4416076746", @@ -23594,7 +23639,7 @@ "kengu", "timalthoff" ], - "targetId": null, + "targetId": "openalex-w4416076746", "base": {}, "changes": { "title": "LSM-2: Learning from Incomplete Wearable Sensor Data", @@ -23654,6 +23699,21 @@ ], "status": "needs-membership-review", "reason": "Enough lab identities match, but publication-time membership needs confirmation." + }, + "reviewedOn": "2026-09-09", + "policyReview": { + "policyVersion": 2, + "labRelevanceStatus": "needs-membership-review", + "personIds": [ + "kengu", + "timalthoff" + ], + "evidenceUrls": [ + "https://arxiv.org/abs/2506.05321", + "https://www.kenqgu.com/" + ], + "overrideReason": "Primary arXiv record verifies Tim Althoff and Ken Gu; the maintainer approved preprint inclusion and Ken's public profile verifies his UW PhD advising relationship with Tim, while exact lab membership dates are not encoded.", + "reviewedOn": "2026-09-09" } }, { @@ -24123,7 +24183,7 @@ }, { "id": "openalex-w4416863183", - "status": "pending", + "status": "accepted", "firstSeen": "2026-09-05", "fingerprint": "f32d27edb3c129fb6f7497faa99f3c71ce6c7e998e3c23dd50e2c803d26cffcc", "sourceUrl": "https://openalex.org/W4416863183", @@ -24151,7 +24211,7 @@ "galenweld", "timalthoff" ], - "targetId": null, + "targetId": "openalex-w4416863183", "base": {}, "changes": { "title": "How Conversational Structure and Style Shape Online Community Experiences", @@ -24192,6 +24252,21 @@ ], "status": "needs-membership-review", "reason": "Enough lab identities match, but publication-time membership needs confirmation." + }, + "reviewedOn": "2026-09-09", + "policyReview": { + "policyVersion": 2, + "labRelevanceStatus": "needs-membership-review", + "personIds": [ + "galenweld", + "timalthoff" + ], + "evidenceUrls": [ + "https://ojs.aaai.org/index.php/ICWSM/article/view/42762", + "https://behavioral-data.github.io/team/" + ], + "overrideReason": "Proceedings record verifies Tim Althoff and Galen Weld; the maintainer approved inclusion and the public lab roster verifies Galen as a past lab member, while exact membership dates are not encoded.", + "reviewedOn": "2026-09-09" } }, { @@ -24342,7 +24417,7 @@ }, { "id": "openalex-w4417141639", - "status": "pending", + "status": "accepted", "firstSeen": "2026-09-05", "fingerprint": "246b5561ae5899e09464dd40bc5d7451ffa24579fe8755d42a58b47740f4ab19", "sourceUrl": "https://openalex.org/W4417141639", @@ -24372,7 +24447,7 @@ "innalin", "timalthoff" ], - "targetId": null, + "targetId": "openalex-w4417141639", "base": {}, "changes": { "title": "Self-Improving VLM Judges Without Human Annotations", @@ -24413,6 +24488,21 @@ ], "status": "needs-membership-review", "reason": "Enough lab identities match, but publication-time membership needs confirmation." + }, + "reviewedOn": "2026-09-09", + "policyReview": { + "policyVersion": 2, + "labRelevanceStatus": "needs-membership-review", + "personIds": [ + "innalin", + "timalthoff" + ], + "evidenceUrls": [ + "https://arxiv.org/abs/2512.05145", + "https://behavioral-data.github.io/team/" + ], + "overrideReason": "Primary arXiv record verifies Tim Althoff and Inna Wanyin Lin; the maintainer approved preprint inclusion and the public lab roster verifies Inna's lab membership, whose dates are not encoded.", + "reviewedOn": "2026-09-09" } }, { @@ -24491,7 +24581,7 @@ }, { "id": "openalex-w4417255562", - "status": "pending", + "status": "accepted", "firstSeen": "2026-09-05", "fingerprint": "1efc440ebe188e0ed32d6bfcf0c0488700ab129016c4c7cab43705026cd0f4cd", "sourceUrl": "https://openalex.org/W4417255562", @@ -24533,7 +24623,7 @@ "kengu", "timalthoff" ], - "targetId": null, + "targetId": "openalex-w4417255562", "base": {}, "changes": { "title": "RADAR: Benchmarking Language Models on Imperfect Tabular Data", @@ -24587,6 +24677,21 @@ ], "status": "needs-membership-review", "reason": "Enough lab identities match, but publication-time membership needs confirmation." + }, + "reviewedOn": "2026-09-09", + "policyReview": { + "policyVersion": 2, + "labRelevanceStatus": "needs-membership-review", + "personIds": [ + "kengu", + "timalthoff" + ], + "evidenceUrls": [ + "https://arxiv.org/abs/2506.08249", + "https://www.kenqgu.com/" + ], + "overrideReason": "Primary arXiv record verifies Tim Althoff and Ken Gu and identifies the NeurIPS publication; the maintainer approved inclusion and Ken's profile verifies the advising relationship with Tim.", + "reviewedOn": "2026-09-09" } }, { @@ -25327,7 +25432,7 @@ }, { "id": "openalex-w7122726170", - "status": "pending", + "status": "accepted", "firstSeen": "2026-09-05", "fingerprint": "83c22c33abc6247e6b1d0a8f30329da901bd958f1b8a7d9592fde6772d311373", "sourceUrl": "https://openalex.org/W7122726170", @@ -25369,7 +25474,7 @@ "mikemerrill", "timalthoff" ], - "targetId": null, + "targetId": "openalex-w7122726170", "base": {}, "changes": { "title": "Transforming wearable data into personal health insights using large language model agents", @@ -25423,6 +25528,21 @@ "status": "needs-membership-review", "reason": "Enough lab identities match, but publication-time membership needs confirmation." }, + "reviewedOn": "2026-09-09", + "policyReview": { + "policyVersion": 2, + "labRelevanceStatus": "needs-membership-review", + "personIds": [ + "mikemerrill", + "timalthoff" + ], + "evidenceUrls": [ + "https://www.nature.com/articles/s41467-025-67922-y", + "https://arxiv.org/abs/2406.06464" + ], + "overrideReason": "Publisher and preprint records verify Tim Althoff and Mike Merrill; the maintainer approved inclusion and the 2024 preprint establishes this lab collaboration before final journal publication.", + "reviewedOn": "2026-09-09" + }, "alternateRecordIds": [ "openalex-w4399596965" ] @@ -27121,7 +27241,7 @@ }, { "id": "openalex-w7162039229", - "status": "pending", + "status": "accepted", "firstSeen": "2026-09-05", "fingerprint": "cb9cfb0d2aa61dfe6050193624e9b8a1d74878096b36a648424a1ae350de1260", "sourceUrl": "https://openalex.org/W7162039229", @@ -27150,7 +27270,7 @@ "innalin", "timalthoff" ], - "targetId": null, + "targetId": "openalex-w7162039229", "base": {}, "changes": { "title": "CandorMD: An AI-Assisted Audio Simulation and Feedback System for Training Clinicians for Medical Error Disclosure", @@ -27191,6 +27311,21 @@ "status": "needs-membership-review", "reason": "Enough lab identities match, but publication-time membership needs confirmation." }, + "reviewedOn": "2026-09-09", + "policyReview": { + "policyVersion": 2, + "labRelevanceStatus": "needs-membership-review", + "personIds": [ + "innalin", + "timalthoff" + ], + "evidenceUrls": [ + "https://arxiv.org/abs/2605.20701", + "https://behavioral-data.github.io/team/" + ], + "overrideReason": "Primary arXiv record verifies Tim Althoff and Inna Wanyin Lin; the maintainer approved inclusion and the public lab roster verifies Inna's lab membership, whose dates are not encoded.", + "reviewedOn": "2026-09-09" + }, "alternateRecordIds": [ "openalex-w7162150239" ] @@ -27841,7 +27976,7 @@ }, { "id": "openalex-w7166852917", - "status": "pending", + "status": "accepted", "firstSeen": "2026-09-05", "fingerprint": "fb149350cb1035cc8ab81b0b39dab94543f68a8165db3f2c0ca4fd466eb11ddc", "sourceUrl": "https://openalex.org/W7166852917", @@ -27869,7 +28004,7 @@ "mikemerrill", "timalthoff" ], - "targetId": null, + "targetId": "openalex-w7166852917", "base": {}, "changes": { "title": "Inferring Events from Time Series using Language Models", @@ -27908,11 +28043,26 @@ ], "status": "needs-membership-review", "reason": "Enough lab identities match, but publication-time membership needs confirmation." + }, + "reviewedOn": "2026-09-09", + "policyReview": { + "policyVersion": 2, + "labRelevanceStatus": "needs-membership-review", + "personIds": [ + "mikemerrill", + "timalthoff" + ], + "evidenceUrls": [ + "https://aclanthology.org/2026.acl-long.157/", + "https://mikemerrill.io/" + ], + "overrideReason": "ACL Anthology verifies Tim Althoff and Mike Merrill; the maintainer approved inclusion and Mike's public biography verifies that Tim advised his UW PhD, while exact lab membership dates are not encoded.", + "reviewedOn": "2026-09-09" } }, { "id": "openalex-w7170643724", - "status": "pending", + "status": "accepted", "firstSeen": "2026-09-05", "fingerprint": "c61bd7d747f8d5872cd05ae9bfdb2d6dccef0e247ff3ab8ff20af921aaabb5c4", "sourceUrl": "https://openalex.org/W7170643724", @@ -27954,7 +28104,7 @@ "kengu", "timalthoff" ], - "targetId": null, + "targetId": "openalex-w7170643724", "base": {}, "changes": { "title": "Capable language models can outgrow the benefits of collaboration", @@ -28010,6 +28160,21 @@ "status": "needs-membership-review", "reason": "Enough lab identities match, but publication-time membership needs confirmation." }, + "reviewedOn": "2026-09-09", + "policyReview": { + "policyVersion": 2, + "labRelevanceStatus": "needs-membership-review", + "personIds": [ + "kengu", + "timalthoff" + ], + "evidenceUrls": [ + "https://www.nature.com/articles/s42256-026-01268-y", + "https://www.kenqgu.com/" + ], + "overrideReason": "Publisher record verifies Tim Althoff and Ken Gu; the maintainer approved inclusion and Ken's public profile verifies his UW PhD advising relationship with Tim, while exact lab membership dates are not encoded.", + "reviewedOn": "2026-09-09" + }, "alternateRecordIds": [ "openalex-w7125481708" ] diff --git a/maintenance/scholar-supplement.json b/maintenance/scholar-supplement.json index c1fad3c..43025c1 100644 --- a/maintenance/scholar-supplement.json +++ b/maintenance/scholar-supplement.json @@ -3,14 +3,27 @@ "candidates": [ { "id": "manual-synthworlds-2026", - "status": "pending", + "status": "accepted", "title": "SynthWorlds: Controlled Parallel Worlds for Disentangling Reasoning and Knowledge in Language Models", "year": 2026, "venue": "ICLR 2026", "authorNames": ["Ken Gu", "Advait Bhat", "Mike Merrill", "Robert West", "Xin Liu", "Daniel McDuff", "Tim Althoff"], "matchedPersonIds": ["kengu", "advaitbhat", "mikemerrill", "timalthoff"], "sourceUrl": "https://proceedings.iclr.cc/paper_files/paper/2026/hash/b213d870740582dd6af77bbdaed900c9-Abstract-Conference.html", - "note": "Found on Tim and Advait's Google Scholar profiles; title and full authors verified in ICLR proceedings. OpenAlex record not found. Review in chat; add through the manual content workflow after approval." + "note": "Found on Tim and Advait's Google Scholar profiles; title and full authors verified in ICLR proceedings. OpenAlex record not found. Added through the manual content workflow after maintainer approval.", + "targetId": "manual-synthworlds-2026", + "reviewedOn": "2026-09-09", + "policyReview": { + "policyVersion": 2, + "labRelevanceStatus": "needs-membership-review", + "personIds": ["kengu", "advaitbhat", "mikemerrill", "timalthoff"], + "evidenceUrls": [ + "https://proceedings.iclr.cc/paper_files/paper/2026/hash/b213d870740582dd6af77bbdaed900c9-Abstract-Conference.html", + "https://behavioral-data.github.io/team/" + ], + "overrideReason": "ICLR proceedings verify Tim Althoff with Ken Gu, Advait Bhat, and Mike Merrill; the maintainer approved inclusion while exact membership dates remain unencoded.", + "reviewedOn": "2026-09-09" + } } ] } From c0e0f5b01bd26b21aa13e7afaa1788462e2cc107 Mon Sep 17 00:00:00 2001 From: Advait Bhat Date: Wed, 9 Sep 2026 13:28:17 -0700 Subject: [PATCH 2/4] Review weekly publication candidates --- content/publications.json | 48 ++++++++++++++--------------- maintenance/review.json | 64 ++++++++++++++++++++------------------- 2 files changed, 57 insertions(+), 55 deletions(-) diff --git a/content/publications.json b/content/publications.json index afe3b3c..52a1779 100644 --- a/content/publications.json +++ b/content/publications.json @@ -1697,7 +1697,7 @@ { "id": "openalex-w7170643724", "title": "Capable language models can outgrow the benefits of collaboration", - "authors": "Yubin Kim and Ken Gu and Chanwoo Park and Chunjong Park and Samuel Schmidgall and A. Ali Heydari and Yao Yan and Zhihan Zhang and Yuchen Zhuang and Liu Y and Mark Malhotra and Paul Pu Liang and Hae Won Park and Yuzhe Yang and Xuhai Xu and Yilun Du and Shwetak Patel and Tim Althoff and Daniel McDuff and Xin Liu", + "authors": "Yubin Kim and Ken Gu and Chanwoo Park and Chunjong Park and Samuel Schmidgall and A. Ali Heydari and Yao Yan and Zhihan Zhang and Yuchen Zhuang and Yun Liu and Mark Malhotra and Paul Pu Liang and Hae Won Park and Yuzhe Yang and Xuhai Xu and Yilun Du and Shwetak Patel and Tim Althoff and Daniel McDuff and Xin Liu", "authorNames": [ "Yubin Kim", "Ken Gu", @@ -1708,7 +1708,7 @@ "Yao Yan", "Zhihan Zhang", "Yuchen Zhuang", - "Liu Y", + "Yun Liu", "Mark Malhotra", "Paul Pu Liang", "Hae Won Park", @@ -1757,7 +1757,7 @@ "doi": "10.1609/icwsm.v20i1.42762", "url": "https://doi.org/10.1609/icwsm.v20i1.42762", "status": "published", - "type": "other", + "type": "conference", "openalexId": "W4416863183", "description": "", "highlight": false, @@ -1789,7 +1789,7 @@ "doi": "10.18653/v1/2026.acl-long.157", "url": "https://aclanthology.org/2026.acl-long.157/", "status": "published", - "type": "article", + "type": "conference", "openalexId": "W7166852917", "description": "", "highlight": false, @@ -1840,7 +1840,7 @@ { "id": "openalex-w7122726170", "title": "Transforming wearable data into personal health insights using large language model agents", - "authors": "Mike A. Merrill and Akshay Paruchuri and Naghmeh Rezaei and Geza Kovacs and Javier Perez and Yun Liu and Erik Schenck and Nova Hammerquist and Jake Sunshine and Shyam Tailor and Kumar Ayush and Hao-Wei Su and Qian He and Cory Y. McLean and Mark Malhotra and Shwetak Patel and Jiening Zhan and Tim Althoff and Daniel McDuff and Yi Liu", + "authors": "Mike A. Merrill and Akshay Paruchuri and Naghmeh Rezaei and Geza Kovacs and Javier Perez and Yun Liu and Erik Schenck and Nova Hammerquist and Jake Sunshine and Shyam Tailor and Kumar Ayush and Hao-Wei Su and Qian He and Cory Y. McLean and Mark Malhotra and Shwetak Patel and Jiening Zhan and Tim Althoff and Daniel McDuff and Xin Liu", "authorNames": [ "Mike A. Merrill", "Akshay Paruchuri", @@ -1861,7 +1861,7 @@ "Jiening Zhan", "Tim Althoff", "Daniel McDuff", - "Yi Liu" + "Xin Liu" ], "year": 2026, "venue": "Nature Communications", @@ -1886,29 +1886,29 @@ { "id": "openalex-w4407425790", "title": "Human Decision-making is Susceptible to AI-driven Manipulation", - "authors": "Sahand Sabour and June M. Liu and Siyang Liu and Yao, Chris Z. and Shiyao Cui and Xuanming Zhang and Wen Zhang and Yaru Cao and Advait Bhat and Jinan Guan and Wei Wu and Rada Mihalcea and Wang, Hongning and Tim Althoff and Tatia M. C. Lee and Minlie Huang", + "authors": "Sahand Sabour and June M. Liu and Siyang Liu and Chris Z. Yao and Shiyao Cui and Xuanming Zhang and Wen Zhang and Yaru Cao and Advait Bhat and Jian Guan and Wei Wu and Rada Mihalcea and Hongning Wang and Tim Althoff and Tatia M.C. Lee and Minlie Huang", "authorNames": [ "Sahand Sabour", "June M. Liu", "Siyang Liu", - "Yao, Chris Z.", + "Chris Z. Yao", "Shiyao Cui", "Xuanming Zhang", "Wen Zhang", "Yaru Cao", "Advait Bhat", - "Jinan Guan", + "Jian Guan", "Wei Wu", "Rada Mihalcea", - "Wang, Hongning", + "Hongning Wang", "Tim Althoff", - "Tatia M. C. Lee", + "Tatia M.C. Lee", "Minlie Huang" ], "year": 2025, "venue": "arXiv (Cornell University)", "doi": "10.48550/arxiv.2502.07663", - "url": "http://arxiv.org/abs/2502.07663", + "url": "https://arxiv.org/abs/2502.07663", "status": "preprint", "type": "preprint", "openalexId": "W4407425790", @@ -2009,7 +2009,7 @@ { "id": "openalex-w4417255562", "title": "RADAR: Benchmarking Language Models on Imperfect Tabular Data", - "authors": "Ken Gu and Zhihan Zhang and Kate Lin and Yuwei Zhang and Akshay Paruchuri and Hong Yu and Mehran Kazemi and Kumar Ayush and A. Ali Heydari and Maxwell A. Xu and Girish Narayanswamy and Yun Liu and Ming-Zher Poh and Yuzhe Yang and Mark Malhotra and Shwetak Patel and Hamid Palangi and Xuhai Xu and Daniel McDuff and Tim Althoff and Xin Liu", + "authors": "Ken Gu and Zhihan Zhang and Kate Lin and Yuwei Zhang and Akshay Paruchuri and Hong Yu and Mehran Kazemi and Kumar Ayush and A. Ali Heydari and Max Xu and Yun Liu and Ming-Zher Poh and Yuzhe Yang and Mark Malhotra and Shwetak Patel and Hamid Palangi and Xuhai \"Orson\" Xu and Daniel McDuff and Tim Althoff and Xin Liu", "authorNames": [ "Ken Gu", "Zhihan Zhang", @@ -2020,15 +2020,14 @@ "Mehran Kazemi", "Kumar Ayush", "A. Ali Heydari", - "Maxwell A. Xu", - "Girish Narayanswamy", + "Max Xu", "Yun Liu", "Ming-Zher Poh", "Yuzhe Yang", "Mark Malhotra", "Shwetak Patel", "Hamid Palangi", - "Xuhai Xu", + "Xuhai \"Orson\" Xu", "Daniel McDuff", "Tim Althoff", "Xin Liu" @@ -2038,7 +2037,7 @@ "doi": "10.52202/085713-3699", "url": "https://doi.org/10.52202/085713-3699", "status": "published", - "type": "article", + "type": "conference", "openalexId": "W4417255562", "description": "", "highlight": false, @@ -2068,7 +2067,7 @@ "doi": "10.1609/icwsm.v19i1.35863", "url": "https://doi.org/10.1609/icwsm.v19i1.35863", "status": "published", - "type": "other", + "type": "conference", "openalexId": "W4411121032", "description": "", "highlight": false, @@ -2097,13 +2096,14 @@ "Tim Althoff", "Marjan Ghazvininejad" ], - "year": 2025, - "venue": "arXiv (Cornell University)", - "doi": "10.48550/arxiv.2512.05145", - "url": "http://arxiv.org/abs/2512.05145", - "status": "preprint", - "type": "preprint", + "year": 2026, + "venue": "ICLR 2026 Workshop on Recursive Self-Improvement", + "doi": "", + "url": "https://openreview.net/forum?id=8hYSvUpJBA", + "status": "accepted", + "type": "conference", "openalexId": "W4417141639", + "arxivId": "2512.05145", "description": "", "highlight": false, "award": "", diff --git a/maintenance/review.json b/maintenance/review.json index 9275d46..7e12322 100644 --- a/maintenance/review.json +++ b/maintenance/review.json @@ -20992,29 +20992,29 @@ "base": {}, "changes": { "title": "Human Decision-making is Susceptible to AI-driven Manipulation", - "authors": "Sahand Sabour and June M. Liu and Siyang Liu and Yao, Chris Z. and Shiyao Cui and Xuanming Zhang and Wen Zhang and Yaru Cao and Advait Bhat and Jinan Guan and Wei Wu and Rada Mihalcea and Wang, Hongning and Tim Althoff and Tatia M. C. Lee and Minlie Huang", + "authors": "Sahand Sabour and June M. Liu and Siyang Liu and Chris Z. Yao and Shiyao Cui and Xuanming Zhang and Wen Zhang and Yaru Cao and Advait Bhat and Jian Guan and Wei Wu and Rada Mihalcea and Hongning Wang and Tim Althoff and Tatia M.C. Lee and Minlie Huang", "authorNames": [ "Sahand Sabour", "June M. Liu", "Siyang Liu", - "Yao, Chris Z.", + "Chris Z. Yao", "Shiyao Cui", "Xuanming Zhang", "Wen Zhang", "Yaru Cao", "Advait Bhat", - "Jinan Guan", + "Jian Guan", "Wei Wu", "Rada Mihalcea", - "Wang, Hongning", + "Hongning Wang", "Tim Althoff", - "Tatia M. C. Lee", + "Tatia M.C. Lee", "Minlie Huang" ], "year": 2025, "venue": "arXiv (Cornell University)", "doi": "10.48550/arxiv.2502.07663", - "url": "http://arxiv.org/abs/2502.07663", + "url": "https://arxiv.org/abs/2502.07663", "status": "preprint", "openalexId": "W4407425790", "type": "preprint" @@ -21050,7 +21050,7 @@ "https://arxiv.org/abs/2502.07663", "https://behavioral-data.github.io/team/" ], - "overrideReason": "Primary arXiv record verifies Tim Althoff and Advait Bhat; the maintainer approved preprint inclusion and the public lab roster verifies Advait's lab membership, whose dates are not encoded.", + "overrideReason": "Primary arXiv record verifies Tim Althoff and Advait Bhat; the maintainer approved preprint inclusion and the public lab roster verifies Advait lab membership, whose dates are not encoded.", "reviewedOn": "2026-09-09" } }, @@ -21605,7 +21605,7 @@ "url": "https://doi.org/10.1609/icwsm.v19i1.35863", "status": "published", "openalexId": "W4411121032", - "type": "other" + "type": "conference" }, "possibleDuplicates": [], "conflicts": [], @@ -21640,7 +21640,7 @@ "https://ojs.aaai.org/index.php/ICWSM/article/view/35863", "https://behavioral-data.github.io/moderator_perceptions_public/" ], - "overrideReason": "Official ICWSM record verifies Tim Althoff and Galen Weld; the maintainer approved inclusion and the lab project site establishes the work as a Behavioral Data Science publication.", + "overrideReason": "Official ICWSM proceedings verify Tim Althoff and Galen Weld; the maintainer approved inclusion and the lab project site establishes the work as a Behavioral Data Science publication.", "reviewedOn": "2026-09-09" }, "alternateRecordIds": [ @@ -24230,7 +24230,7 @@ "url": "https://doi.org/10.1609/icwsm.v20i1.42762", "status": "published", "openalexId": "W4416863183", - "type": "other" + "type": "conference" }, "possibleDuplicates": [], "conflicts": [], @@ -24265,7 +24265,7 @@ "https://ojs.aaai.org/index.php/ICWSM/article/view/42762", "https://behavioral-data.github.io/team/" ], - "overrideReason": "Proceedings record verifies Tim Althoff and Galen Weld; the maintainer approved inclusion and the public lab roster verifies Galen as a past lab member, while exact membership dates are not encoded.", + "overrideReason": "Official ICWSM proceedings verify Tim Althoff and Galen Weld; the maintainer approved inclusion and the public lab roster verifies Galen as a past lab member, while exact membership dates are not encoded.", "reviewedOn": "2026-09-09" } }, @@ -24462,13 +24462,13 @@ "Tim Althoff", "Marjan Ghazvininejad" ], - "year": 2025, - "venue": "arXiv (Cornell University)", - "doi": "10.48550/arxiv.2512.05145", - "url": "http://arxiv.org/abs/2512.05145", - "status": "preprint", + "year": 2026, + "venue": "ICLR 2026 Workshop on Recursive Self-Improvement", + "doi": "", + "url": "https://openreview.net/forum?id=8hYSvUpJBA", + "status": "accepted", "openalexId": "W4417141639", - "type": "preprint" + "type": "conference" }, "possibleDuplicates": [], "conflicts": [], @@ -24498,10 +24498,12 @@ "timalthoff" ], "evidenceUrls": [ + "https://recursive-workshop.github.io/papers.html", + "https://openreview.net/forum?id=8hYSvUpJBA", "https://arxiv.org/abs/2512.05145", "https://behavioral-data.github.io/team/" ], - "overrideReason": "Primary arXiv record verifies Tim Althoff and Inna Wanyin Lin; the maintainer approved preprint inclusion and the public lab roster verifies Inna's lab membership, whose dates are not encoded.", + "overrideReason": "The official ICLR workshop list and OpenReview identify the accepted poster, arXiv verifies Tim Althoff and Inna Wanyin Lin, and the public lab roster verifies Inna lab membership while exact dates are not encoded.", "reviewedOn": "2026-09-09" } }, @@ -24627,7 +24629,7 @@ "base": {}, "changes": { "title": "RADAR: Benchmarking Language Models on Imperfect Tabular Data", - "authors": "Ken Gu and Zhihan Zhang and Kate Lin and Yuwei Zhang and Akshay Paruchuri and Hong Yu and Mehran Kazemi and Kumar Ayush and A. Ali Heydari and Maxwell A. Xu and Girish Narayanswamy and Yun Liu and Ming-Zher Poh and Yuzhe Yang and Mark Malhotra and Shwetak Patel and Hamid Palangi and Xuhai Xu and Daniel McDuff and Tim Althoff and Xin Liu", + "authors": "Ken Gu and Zhihan Zhang and Kate Lin and Yuwei Zhang and Akshay Paruchuri and Hong Yu and Mehran Kazemi and Kumar Ayush and A. Ali Heydari and Max Xu and Yun Liu and Ming-Zher Poh and Yuzhe Yang and Mark Malhotra and Shwetak Patel and Hamid Palangi and Xuhai \"Orson\" Xu and Daniel McDuff and Tim Althoff and Xin Liu", "authorNames": [ "Ken Gu", "Zhihan Zhang", @@ -24638,15 +24640,14 @@ "Mehran Kazemi", "Kumar Ayush", "A. Ali Heydari", - "Maxwell A. Xu", - "Girish Narayanswamy", + "Max Xu", "Yun Liu", "Ming-Zher Poh", "Yuzhe Yang", "Mark Malhotra", "Shwetak Patel", "Hamid Palangi", - "Xuhai Xu", + "Xuhai \"Orson\" Xu", "Daniel McDuff", "Tim Althoff", "Xin Liu" @@ -24657,7 +24658,7 @@ "url": "https://doi.org/10.52202/085713-3699", "status": "published", "openalexId": "W4417255562", - "type": "article" + "type": "conference" }, "possibleDuplicates": [], "conflicts": [], @@ -24687,10 +24688,11 @@ "timalthoff" ], "evidenceUrls": [ + "https://proceedings.neurips.cc/paper_files/paper/2025/hash/a0434f04b6437e875d52d0b0e25c1729-Abstract-Datasets_and_Benchmarks_Track.html", "https://arxiv.org/abs/2506.08249", "https://www.kenqgu.com/" ], - "overrideReason": "Primary arXiv record verifies Tim Althoff and Ken Gu and identifies the NeurIPS publication; the maintainer approved inclusion and Ken's profile verifies the advising relationship with Tim.", + "overrideReason": "Official NeurIPS proceedings verify the final author list and Tim Althoff with Ken Gu; the maintainer approved inclusion and the Ken Gu public profile verifies the advising relationship with Tim.", "reviewedOn": "2026-09-09" } }, @@ -25478,7 +25480,7 @@ "base": {}, "changes": { "title": "Transforming wearable data into personal health insights using large language model agents", - "authors": "Mike A. Merrill and Akshay Paruchuri and Naghmeh Rezaei and Geza Kovacs and Javier Perez and Yun Liu and Erik Schenck and Nova Hammerquist and Jake Sunshine and Shyam Tailor and Kumar Ayush and Hao-Wei Su and Qian He and Cory Y. McLean and Mark Malhotra and Shwetak Patel and Jiening Zhan and Tim Althoff and Daniel McDuff and Yi Liu", + "authors": "Mike A. Merrill and Akshay Paruchuri and Naghmeh Rezaei and Geza Kovacs and Javier Perez and Yun Liu and Erik Schenck and Nova Hammerquist and Jake Sunshine and Shyam Tailor and Kumar Ayush and Hao-Wei Su and Qian He and Cory Y. McLean and Mark Malhotra and Shwetak Patel and Jiening Zhan and Tim Althoff and Daniel McDuff and Xin Liu", "authorNames": [ "Mike A. Merrill", "Akshay Paruchuri", @@ -25499,7 +25501,7 @@ "Jiening Zhan", "Tim Althoff", "Daniel McDuff", - "Yi Liu" + "Xin Liu" ], "year": 2026, "venue": "Nature Communications", @@ -28023,7 +28025,7 @@ "url": "https://aclanthology.org/2026.acl-long.157/", "status": "published", "openalexId": "W7166852917", - "type": "article" + "type": "conference" }, "possibleDuplicates": [], "conflicts": [], @@ -28056,7 +28058,7 @@ "https://aclanthology.org/2026.acl-long.157/", "https://mikemerrill.io/" ], - "overrideReason": "ACL Anthology verifies Tim Althoff and Mike Merrill; the maintainer approved inclusion and Mike's public biography verifies that Tim advised his UW PhD, while exact lab membership dates are not encoded.", + "overrideReason": "ACL Anthology verifies Tim Althoff and Mike Merrill; the maintainer approved inclusion and the Mike Merrill public biography verifies that Tim advised his UW PhD, while exact lab membership dates are not encoded.", "reviewedOn": "2026-09-09" } }, @@ -28108,7 +28110,7 @@ "base": {}, "changes": { "title": "Capable language models can outgrow the benefits of collaboration", - "authors": "Yubin Kim and Ken Gu and Chanwoo Park and Chunjong Park and Samuel Schmidgall and A. Ali Heydari and Yao Yan and Zhihan Zhang and Yuchen Zhuang and Liu Y and Mark Malhotra and Paul Pu Liang and Hae Won Park and Yuzhe Yang and Xuhai Xu and Yilun Du and Shwetak Patel and Tim Althoff and Daniel McDuff and Xin Liu", + "authors": "Yubin Kim and Ken Gu and Chanwoo Park and Chunjong Park and Samuel Schmidgall and A. Ali Heydari and Yao Yan and Zhihan Zhang and Yuchen Zhuang and Yun Liu and Mark Malhotra and Paul Pu Liang and Hae Won Park and Yuzhe Yang and Xuhai Xu and Yilun Du and Shwetak Patel and Tim Althoff and Daniel McDuff and Xin Liu", "authorNames": [ "Yubin Kim", "Ken Gu", @@ -28119,7 +28121,7 @@ "Yao Yan", "Zhihan Zhang", "Yuchen Zhuang", - "Liu Y", + "Yun Liu", "Mark Malhotra", "Paul Pu Liang", "Hae Won Park", @@ -28172,7 +28174,7 @@ "https://www.nature.com/articles/s42256-026-01268-y", "https://www.kenqgu.com/" ], - "overrideReason": "Publisher record verifies Tim Althoff and Ken Gu; the maintainer approved inclusion and Ken's public profile verifies his UW PhD advising relationship with Tim, while exact lab membership dates are not encoded.", + "overrideReason": "Publisher record verifies Tim Althoff and Ken Gu; the maintainer approved inclusion and the Ken Gu public profile verifies his UW PhD advising relationship with Tim, while exact lab membership dates are not encoded.", "reviewedOn": "2026-09-09" }, "alternateRecordIds": [ From cfff9d7ad77c6063097a41f79e598221bff7f4f5 Mon Sep 17 00:00:00 2001 From: Advait Bhat Date: Wed, 9 Sep 2026 13:29:17 -0700 Subject: [PATCH 3/4] Document corrected publication review batch --- _planning/SCHOLAR_LATEST_REVIEW.md | 2 +- docs/PUBLICATION_POLICY.md | 2 +- 2 files changed, 2 insertions(+), 2 deletions(-) diff --git a/_planning/SCHOLAR_LATEST_REVIEW.md b/_planning/SCHOLAR_LATEST_REVIEW.md index 8249fd4..f226697 100644 --- a/_planning/SCHOLAR_LATEST_REVIEW.md +++ b/_planning/SCHOLAR_LATEST_REVIEW.md @@ -22,7 +22,7 @@ The resulting 13-paper review list is in `maintenance/recent-review.md`. All rem - **SynthWorlds** — Ken Gu, Advait Bhat, Mike Merrill and Tim Althoff appear in the [ICLR 2026 proceedings](https://proceedings.iclr.cc/paper_files/paper/2026/hash/b213d870740582dd6af77bbdaed900c9-Abstract-Conference.html). No matching OpenAlex work found; retained separately without inventing an ID. - **Inferring Events from Time Series using Language Models** — Mike Merrill and Tim; [ACL 2026 proceedings](https://aclanthology.org/2026.acl-long.157/). - **CandorMD** — Inna and Tim; [arXiv](https://arxiv.org/abs/2605.20701). Two OpenAlex records share this arXiv identifier and are grouped. -- **Self-Improving VLM Judges Without Human Annotations** — Inna and Tim; [arXiv](https://arxiv.org/abs/2512.05145). +- **Self-Improving VLM Judges Without Human Annotations** — Inna and Tim; accepted as a poster at the [ICLR 2026 Workshop on Recursive Self-Improvement](https://recursive-workshop.github.io/papers.html), with the final workshop record on [OpenReview](https://openreview.net/forum?id=8hYSvUpJBA) and the earlier version retained on [arXiv](https://arxiv.org/abs/2512.05145). - **RADAR** — Ken and Tim; [Google Research publication record](https://research.google/pubs/radar-benchmarking-language-models-on-imperfect-tabular-data/), [arXiv](https://arxiv.org/abs/2506.08249). Use NeurIPS 2025 pending review, despite Scholar displaying 2026. - **Artificial Hivemind** — Margaret and Mickel; [arXiv full author list](https://arxiv.org/abs/2510.22954). Use NeurIPS 2025 pending review, despite Scholar displaying 2026. - **Transforming wearable data into personal health insights using large language model agents** — Mike Merrill and Tim; [Nature Communications, January 2026](https://www.nature.com/articles/s41467-025-67922-y). The [2024 preprint](https://arxiv.org/abs/2406.06464) reports acceptance to Nature Communications and has the same full author list and PHIA study; grouped with the journal record. diff --git a/docs/PUBLICATION_POLICY.md b/docs/PUBLICATION_POLICY.md index 7365b09..a1ceef5 100644 --- a/docs/PUBLICATION_POLICY.md +++ b/docs/PUBLICATION_POLICY.md @@ -28,4 +28,4 @@ The initial registry fetches the eight current members' primary records. It also The initial September 5 pull retrieved 383 OpenAlex work records. A follow-up check of the newest Scholar entries for all eight current members recovered missing records and author aliases. The queue now retains 390 source records, with 68 alternate versions grouped into 322 canonical review items. After the September 9 policy clarification and reassessment, 57 pending candidates need membership review and 264 do not meet the rule; Artificial Hivemind is durably rejected because Tim is not an author. The recent review sheet contains 11 OpenAlex papers dated 2025 onward plus SynthWorlds, a manually sourced ICLR 2026 candidate. These are suggestions, not approved publications. -Start with `maintenance/recent-review.md`; the full queue is in `maintenance/batch.md` and `maintenance/review.json`. Scholar findings and source links are recorded in `_planning/SCHOLAR_LATEST_REVIEW.md`. The existing 47 approved publications were unchanged by discovery. The first review set proposes 12 records, including four preprints, but those additions belong to a separate reviewed content change and are not part of this migration branch. No live-sync checkpoint was written, and discovery and scheduled workflows remain disabled during this pilot. Complete historical coverage and membership intervals remain review tasks. +Start with `maintenance/recent-review.md`; the full queue is in `maintenance/batch.md` and `maintenance/review.json`. Scholar findings and source links are recorded in `_planning/SCHOLAR_LATEST_REVIEW.md`. The existing 47 approved publications were unchanged by discovery. The first reviewed content batch adds 12 records: eight published papers, three visibly labeled preprints, and one accepted ICLR workshop poster that retains its arXiv identifier. No live-sync checkpoint was written, and discovery and scheduled workflows remain disabled during this pilot. Complete historical coverage and membership intervals remain review tasks. From fed22d71aa51932613b9b3e1abad97991dd3fb83 Mon Sep 17 00:00:00 2001 From: Advait Bhat Date: Wed, 9 Sep 2026 13:41:56 -0700 Subject: [PATCH 4/4] Refresh publication review status --- _planning/SCHOLAR_LATEST_REVIEW.md | 14 +++++++------- docs/PUBLICATION_POLICY.md | 2 +- 2 files changed, 8 insertions(+), 8 deletions(-) diff --git a/_planning/SCHOLAR_LATEST_REVIEW.md b/_planning/SCHOLAR_LATEST_REVIEW.md index f226697..28ed1bd 100644 --- a/_planning/SCHOLAR_LATEST_REVIEW.md +++ b/_planning/SCHOLAR_LATEST_REVIEW.md @@ -2,7 +2,7 @@ Checked 2026-09-05 using the built-in browser, with each confirmed profile sorted by publication date. Inspected the first page (up to 20 entries per member; fewer when the profile had fewer papers). This is a recent-paper check, not a complete career bibliography. Scholar sometimes gives a proceedings paper a later year than the conference itself and can contain duplicate versions. Truncated Scholar author lists are not evidence that no other lab author exists. -The resulting 13-paper review list is in `maintenance/recent-review.md`. All remain pending. The main OpenAlex queue retains source observations, with duplicate versions grouped rather than deleted. The source-only SynthWorlds entry is in `maintenance/scholar-supplement.json` until it is reviewed and imported through the manual content workflow. +The resulting 13-paper review list has been resolved, so `maintenance/recent-review.md` is now empty. Eleven OpenAlex candidates and the source-only SynthWorlds supplement were accepted; Artificial Hivemind was rejected because Tim is not an author. The main queue and supplement retain the source observations and decisions, with duplicate versions grouped rather than deleted. ## Profiles checked @@ -11,25 +11,25 @@ The resulting 13-paper review list is in `maintenance/recent-review.md`. All rem | Tim Althoff | [Scholar](https://scholar.google.com/citations?user=yc4nBNgAAAAJ&hl=en&sortby=pubdate) | Capable language models can outgrow the benefits of collaboration; Responsible Evaluation of AI for Mental Health; Inferring events from time series using language models (2026) | Recovered missing papers from split OpenAlex IDs. CandorMD appeared twice. SynthWorlds needed a manual source entry. Responsible Evaluation of AI for Mental Health has only Tim among the known roster after checking the full primary-source list. | | Jina Suh | [Scholar](https://scholar.google.com/citations?user=LuNehzsAAAAJ&hl=en&sortby=pubdate) | “Always Want to Use it for Everything”: Understanding Young Adults' Perceptions of AI Dependence; Psychological Influences of Conversational AI; The agony of opacity (2026) | No additional two-lab-author match verified in this check. Several full author lists remain to be checked; do not treat an abbreviated “M Li” as Margaret. | | Yige Yuan | [Scholar](https://scholar.google.com/citations?user=lf6GtCIAAAAJ&hl=en&sortby=pubdate) | Co-harness: Co-evolving harnesses and model weights for LLM agents; Rethinking evaluation of harness evolution for agents; Do We Always Need Query-Level Workflows? (2026) | No additional two-lab-author match verified. OpenAlex profile has namesakes; do not assume every record belongs to Yige. | -| Advait Bhat | [Scholar](https://scholar.google.com/citations?user=lFBWQb0AAAAJ&hl=en&sortby=pubdate) | SynthWorlds; Reactive Writers; Biased AI writing assistants shift users’ attitudes on societal issues (2026) | SynthWorlds added for review; Human Decision-making is Susceptible to AI-driven Manipulation was already present. Reactive Writers and Biased AI have only Advait among the known lab roster, so were not promoted to the lab review list. | +| Advait Bhat | [Scholar](https://scholar.google.com/citations?user=lFBWQb0AAAAJ&hl=en&sortby=pubdate) | SynthWorlds; Reactive Writers; Biased AI writing assistants shift users’ attitudes on societal issues (2026) | SynthWorlds was recovered for manual review and accepted; Human Decision-making is Susceptible to AI-driven Manipulation was already present. Reactive Writers and Biased AI have only Advait among the known lab roster, so were not promoted to the lab review list. | | Cheng Li | [Scholar](https://scholar.google.com/citations?user=083GCIwAAAAJ&hl=en&sortby=pubdate) | VEglue; Seeing is believing; CultureVLM (2025) | No 2026 entry visible. No additional two-lab-author match verified. Same-name OpenAlex records require per-paper verification. | | Deniz Nazar | [Scholar](https://scholar.google.com/citations?user=ZxQO6oMAAAAJ&hl=en&sortby=pubdate) | Beyond One Output; NLP for Social Good (2026); a 2023 paper | Beyond One Output has Deniz with Emily Reif, Cindy Yang, Jena Hwang, Noah Smith, and Jeff Heer; no second known lab author. No additional lab candidate verified. | | Inna Lin | [Scholar](https://scholar.google.com/citations?user=LRrRtfwAAAAJ&hl=en&sortby=pubdate) | Muse Spark Safety & Preparedness Report (2026); Self-Improving VLM Judges Without Human Annotations (2025) | Recovered Tim's split identity on Self-Improving VLM Judges. CandorMD was found on Tim's profile but absent from Inna's visible list. The full Muse Spark author list contains no second known lab member. | -| Margaret Li | [Scholar](https://scholar.google.com/citations?user=cUSS3fYAAAAJ&hl=en&sortby=pubdate) | Slicing and Dicing; Compute-Optimal Tokenization; Artificial Hivemind (Scholar: 2026) | Artificial Hivemind includes past member Mickel Liu, verified from arXiv. Full author lists for Compute Optimal Tokenization, FlexOlmo and Precise Information Control contain no second known lab member. | +| Margaret Li | [Scholar](https://scholar.google.com/citations?user=cUSS3fYAAAAJ&hl=en&sortby=pubdate) | Slicing and Dicing; Compute-Optimal Tokenization; Artificial Hivemind (Scholar: 2026) | Artificial Hivemind includes Margaret and past member Mickel Liu but was rejected because Tim is not an author. Full author lists for Compute Optimal Tokenization, FlexOlmo and Precise Information Control contain no second known lab member. | ## Recovered candidates and primary evidence -- **SynthWorlds** — Ken Gu, Advait Bhat, Mike Merrill and Tim Althoff appear in the [ICLR 2026 proceedings](https://proceedings.iclr.cc/paper_files/paper/2026/hash/b213d870740582dd6af77bbdaed900c9-Abstract-Conference.html). No matching OpenAlex work found; retained separately without inventing an ID. +- **SynthWorlds** — Ken Gu, Advait Bhat, Mike Merrill and Tim Althoff appear in the [ICLR 2026 proceedings](https://proceedings.iclr.cc/paper_files/paper/2026/hash/b213d870740582dd6af77bbdaed900c9-Abstract-Conference.html). No matching OpenAlex work was found; it was accepted through the manual supplement without inventing an ID. - **Inferring Events from Time Series using Language Models** — Mike Merrill and Tim; [ACL 2026 proceedings](https://aclanthology.org/2026.acl-long.157/). - **CandorMD** — Inna and Tim; [arXiv](https://arxiv.org/abs/2605.20701). Two OpenAlex records share this arXiv identifier and are grouped. - **Self-Improving VLM Judges Without Human Annotations** — Inna and Tim; accepted as a poster at the [ICLR 2026 Workshop on Recursive Self-Improvement](https://recursive-workshop.github.io/papers.html), with the final workshop record on [OpenReview](https://openreview.net/forum?id=8hYSvUpJBA) and the earlier version retained on [arXiv](https://arxiv.org/abs/2512.05145). -- **RADAR** — Ken and Tim; [Google Research publication record](https://research.google/pubs/radar-benchmarking-language-models-on-imperfect-tabular-data/), [arXiv](https://arxiv.org/abs/2506.08249). Use NeurIPS 2025 pending review, despite Scholar displaying 2026. -- **Artificial Hivemind** — Margaret and Mickel; [arXiv full author list](https://arxiv.org/abs/2510.22954). Use NeurIPS 2025 pending review, despite Scholar displaying 2026. +- **RADAR** — Ken and Tim; [final NeurIPS 2025 proceedings](https://proceedings.neurips.cc/paper_files/paper/2025/hash/a0434f04b6437e875d52d0b0e25c1729-Abstract-Datasets_and_Benchmarks_Track.html), [Google Research publication record](https://research.google/pubs/radar-benchmarking-language-models-on-imperfect-tabular-data/), and [arXiv](https://arxiv.org/abs/2506.08249). The final proceedings record was accepted despite Scholar displaying 2026. +- **Artificial Hivemind** — Margaret and Mickel; [arXiv full author list](https://arxiv.org/abs/2510.22954). Rejected because Tim is not an author. - **Transforming wearable data into personal health insights using large language model agents** — Mike Merrill and Tim; [Nature Communications, January 2026](https://www.nature.com/articles/s41467-025-67922-y). The [2024 preprint](https://arxiv.org/abs/2406.06464) reports acceptance to Nature Communications and has the same full author list and PHIA study; grouped with the journal record. ## Duplicate handling -The recent list contains one Reddit Rules and Rulers item, one CandorMD item, and one item for each verified preprint/publication pair. [Capable language models can outgrow the benefits of collaboration](https://www.nature.com/articles/s42256-026-01268-y) explicitly links [Towards a Science of Scaling Agent Systems](https://arxiv.org/abs/2512.08296) as its preprint, despite the changed title. +The recent list is empty after the first review decisions. The queue still retains one canonical Reddit Rules and Rulers item, one CandorMD item, and one canonical item for each verified preprint/publication pair. [Capable language models can outgrow the benefits of collaboration](https://www.nature.com/articles/s42256-026-01268-y) explicitly links [Towards a Science of Scaling Agent Systems](https://arxiv.org/abs/2512.08296) as its preprint, despite the changed title. Automatic grouping uses matching DOI/arXiv IDs or matching normalized titles and complete author lists. It does not use fuzzy title similarity. Dataset versions, author responses, corrections and errata need shared identifiers or explicit evidence. Manually verified groups keep their sources and reasons in `review.json`. Rejected or deferred canonical items continue to suppress alternate versions. diff --git a/docs/PUBLICATION_POLICY.md b/docs/PUBLICATION_POLICY.md index a1ceef5..9cf3b04 100644 --- a/docs/PUBLICATION_POLICY.md +++ b/docs/PUBLICATION_POLICY.md @@ -26,6 +26,6 @@ The initial registry fetches the eight current members' primary records. It also ## First local batch -The initial September 5 pull retrieved 383 OpenAlex work records. A follow-up check of the newest Scholar entries for all eight current members recovered missing records and author aliases. The queue now retains 390 source records, with 68 alternate versions grouped into 322 canonical review items. After the September 9 policy clarification and reassessment, 57 pending candidates need membership review and 264 do not meet the rule; Artificial Hivemind is durably rejected because Tim is not an author. The recent review sheet contains 11 OpenAlex papers dated 2025 onward plus SynthWorlds, a manually sourced ICLR 2026 candidate. These are suggestions, not approved publications. +The initial September 5 pull retrieved 383 OpenAlex work records. A follow-up check of the newest Scholar entries for all eight current members recovered missing records and author aliases. The queue now retains 390 source records, with 68 alternate versions grouped into 322 canonical review items. After the September 9 policy clarification and completed first review, 46 pending candidates need membership review and 264 do not meet the rule. Eleven OpenAlex papers and the manually sourced SynthWorlds record were accepted; Artificial Hivemind is durably rejected because Tim is not an author. The focused recent-review sheet is now empty. Start with `maintenance/recent-review.md`; the full queue is in `maintenance/batch.md` and `maintenance/review.json`. Scholar findings and source links are recorded in `_planning/SCHOLAR_LATEST_REVIEW.md`. The existing 47 approved publications were unchanged by discovery. The first reviewed content batch adds 12 records: eight published papers, three visibly labeled preprints, and one accepted ICLR workshop poster that retains its arXiv identifier. No live-sync checkpoint was written, and discovery and scheduled workflows remain disabled during this pilot. Complete historical coverage and membership intervals remain review tasks.