|
2026-09-23
ยง
|
| 18:36 |
<taavi@cumin1004> |
START - Cookbook sre.gerrit.read-only-toggle from gerrit2003.wikimedia.org |
[production] |
| 18:30 |
<taavi@cumin1004> |
START - Cookbook sre.gerrit.localbackup Prepare local backup on: gerrit1003.wikimedia.org |
[production] |
| 18:29 |
<lerickson@deploy1003> |
helmfile [dse-k8s-eqiad] DONE helmfile.d/services/wdqs-next: apply |
[production] |
| 18:28 |
<lerickson@deploy1003> |
helmfile [dse-k8s-eqiad] START helmfile.d/services/wdqs-next: apply |
[production] |
| 18:27 |
<wmbot~lucaswerkmeister@tools-bastion-15> |
deployed f1b263f216 (upgrade dependencies) and 32a27cf571 (fix README) |
[tools.speedpatrolling] |
| 18:14 |
<vriley@cumin1004> |
END (PASS) - Cookbook sre.hosts.reimage (exit_code=0) for host zuul1005.eqiad.wmnet with OS trixie |
[production] |
| 18:08 |
<sukhe@cumin1004> |
END (FAIL) - Cookbook sre.dns.wipe-cache (exit_code=99) idp.wikimedia.org on all recursors |
[production] |
| 18:08 |
<sukhe@cumin1004> |
START - Cookbook sre.dns.wipe-cache idp.wikimedia.org on all recursors |
[production] |
| 18:05 |
<cdanis@cumin1004> |
END (FAIL) - Cookbook sre.dns.wipe-cache (exit_code=99) _etcd-client-ssl._tcp.eqsin.wmnet on all recursors |
[production] |
| 18:05 |
<cdanis@cumin1004> |
START - Cookbook sre.dns.wipe-cache _etcd-client-ssl._tcp.eqsin.wmnet on all recursors |
[production] |
| 18:03 |
<cdanis@cumin1004> |
END (FAIL) - Cookbook sre.dns.wipe-cache (exit_code=99) _etcd-client-ssl._tcp.eqsin.wmnet on all recursors |
[production] |
| 18:03 |
<cdanis@cumin1004> |
START - Cookbook sre.dns.wipe-cache _etcd-client-ssl._tcp.eqsin.wmnet on all recursors |
[production] |
| 18:02 |
<cdanis@cumin1004> |
END (FAIL) - Cookbook sre.dns.wipe-cache (exit_code=99) _etcd-client-ssl._tcp.ulsfo.wmnet on all recursors |
[production] |
| 18:02 |
<cdanis@cumin1004> |
START - Cookbook sre.dns.wipe-cache _etcd-client-ssl._tcp.ulsfo.wmnet on all recursors |
[production] |
| 18:01 |
<cdanis@dns1005> |
END - running authdns-update |
[production] |
| 17:58 |
<cdanis@dns1005> |
START - running authdns-update |
[production] |
| 17:57 |
<vriley@cumin1004> |
END (PASS) - Cookbook sre.hosts.downtime (exit_code=0) for 2:00:00 on zuul1005.eqiad.wmnet with reason: host reimage |
[production] |
| 17:54 |
<taavi@dns1004> |
END - running authdns-update |
[production] |
| 17:53 |
<vriley@cumin1004> |
START - Cookbook sre.hosts.downtime for 2:00:00 on zuul1005.eqiad.wmnet with reason: host reimage |
[production] |
| 17:51 |
<taavi@dns1004> |
START - running authdns-update |
[production] |
| 17:46 |
<taavi@dns1004> |
END - running authdns-update |
[production] |
| 17:43 |
<taavi@dns1004> |
START - running authdns-update |
[production] |
| 17:37 |
<vriley@cumin1004> |
START - Cookbook sre.hosts.reimage for host zuul1005.eqiad.wmnet with OS trixie |
[production] |
| 17:35 |
<rzl@cumin1004> |
START - Cookbook sre.discovery.datacenter pool all active/active services in eqiad: maintenance - T439010 |
[production] |
| 17:35 |
<cdobbins@cumin1004> |
END (PASS) - Cookbook sre.hosts.reimage (exit_code=0) for host ncredir6001.drmrs.wmnet with OS trixie |
[production] |
| 17:35 |
<cdanis@cumin1004> |
END (FAIL) - Cookbook sre.dns.admin (exit_code=99) DNS admin: depool codfw [reason: no reason specified, no task ID specified] |
[production] |
| 17:35 |
<cdanis@cumin1004> |
START - Cookbook sre.dns.admin DNS admin: depool codfw [reason: no reason specified, no task ID specified] |
[production] |
| 17:24 |
<sukhe@cumin1004> |
END (PASS) - Cookbook sre.dns.admin (exit_code=0) DNS admin: depool codfw [reason: no reason specified, no task ID specified] |
[production] |
| 17:23 |
<sukhe@cumin1004> |
START - Cookbook sre.dns.admin DNS admin: depool codfw [reason: no reason specified, no task ID specified] |
[production] |
| 17:21 |
<lerickson@deploy1003> |
helmfile [dse-k8s-eqiad] DONE helmfile.d/services/wdqs: apply |
[production] |
| 17:21 |
<lerickson@deploy1003> |
helmfile [dse-k8s-eqiad] START helmfile.d/services/wdqs: apply |
[production] |
| 17:18 |
<lerickson@deploy1003> |
helmfile [dse-k8s-eqiad] DONE helmfile.d/services/wdqs: apply |
[production] |
| 17:17 |
<lerickson@deploy1003> |
helmfile [dse-k8s-eqiad] START helmfile.d/services/wdqs: apply |
[production] |
| 17:16 |
<lerickson@deploy1003> |
helmfile [dse-k8s-codfw] DONE helmfile.d/services/wdqs-next: apply |
[production] |
| 17:14 |
<sukhe@cumin1004> |
END (PASS) - Cookbook sre.hosts.downtime (exit_code=0) for 2:00:00 on cp2059.codfw.wmnet with reason: host reimage |
[production] |
| 17:11 |
<sukhe@puppetserver1001> |
conftool action : set/pooled=no; selector: name=cp2049.codfw.wmnet |
[production] |
| 17:11 |
<sukhe@puppetserver1001> |
conftool action : set/pooled=no; selector: name=cp2049 |
[production] |
| 17:10 |
<sukhe@cumin1004> |
START - Cookbook sre.hosts.downtime for 2:00:00 on cp2059.codfw.wmnet with reason: host reimage |
[production] |
| 17:07 |
<mutante> |
cloudcontrol2005-dev, cloudcontrol2006-dev, cloudcontrol2010-dev: restart zookeeper, enabled logging (/var/log/zookeeper/zookeeper.log) after gerrit:1342354 |
[production] |
| 17:02 |
<cdobbins@cumin1004> |
END (PASS) - Cookbook sre.hosts.downtime (exit_code=0) for 2:00:00 on ncredir6001.drmrs.wmnet with reason: host reimage |
[production] |
| 16:59 |
<cdobbins@cumin1004> |
START - Cookbook sre.hosts.downtime for 2:00:00 on ncredir6001.drmrs.wmnet with reason: host reimage |
[production] |
| 16:51 |
<sukhe@cumin1004> |
START - Cookbook sre.hosts.reimage for host cp2059.codfw.wmnet with OS trixie |
[production] |
| 16:51 |
<sukhe@cumin1004> |
END (FAIL) - Cookbook sre.hosts.reimage (exit_code=99) for host cp2059.codfw.wmnet with OS trixie |
[production] |
| 16:48 |
<sukhe@cumin1004> |
START - Cookbook sre.hosts.reimage for host cp2059.codfw.wmnet with OS trixie |
[production] |
| 16:39 |
<sukhe@cumin1004> |
END (ERROR) - Cookbook sre.hosts.reimage (exit_code=97) for host cp2059.codfw.wmnet with OS trixie |
[production] |
| 16:35 |
<dzahn@cumin2003> |
START - Cookbook sre.hosts.reimage for host zuul1005.eqiad.wmnet with OS trixie |
[production] |
| 16:34 |
<dzahn@cumin2003> |
END (FAIL) - Cookbook sre.hosts.reimage (exit_code=99) for host zuul1005.eqiad.wmnet with OS trixie |
[production] |
| 16:30 |
<lerickson@deploy1003> |
helmfile [dse-k8s-codfw] START helmfile.d/services/wdqs-next: apply |
[production] |
| 16:29 |
<cdobbins@cumin1004> |
START - Cookbook sre.hosts.reimage for host ncredir6001.drmrs.wmnet with OS trixie |
[production] |
| 16:10 |
<sukhe@cumin1004> |
START - Cookbook sre.hosts.reimage for host cp2059.codfw.wmnet with OS trixie |
[production] |