|
2026-09-23
ยง
|
| 19:08 |
<sukhe@puppetserver1001> |
conftool action : set/pooled=no; selector: dc=codfw,cluster=dnsbox,service=authdns-update |
[production] |
| 18:59 |
<sukhe@cumin1004> |
DONE (FAIL) - Cookbook sre.hosts.downtime (exit_code=99) for 6:00:00 on 980 hosts with reason: power is still coming back on |
[production] |
| 18:58 |
<taavi@cumin1004> |
END (PASS) - Cookbook sre.dns.wipe-cache (exit_code=0) gerrit.discovery.wmnet on all recursors |
[production] |
| 18:58 |
<taavi@cumin1004> |
START - Cookbook sre.dns.wipe-cache gerrit.discovery.wmnet on all recursors |
[production] |
| 18:50 |
<taavi@cumin1004> |
END (PASS) - Cookbook sre.gerrit.localbackup (exit_code=0) Prepare local backup on: gerrit2003.wikimedia.org |
[production] |
| 18:45 |
<sukhe@dns1004> |
END - running authdns-update |
[production] |
| 18:43 |
<sukhe@dns1004> |
START - running authdns-update |
[production] |
| 18:43 |
<taavi@cumin1004> |
START - Cookbook sre.gerrit.localbackup Prepare local backup on: gerrit2003.wikimedia.org |
[production] |
| 18:42 |
<sukhe@puppetserver1001> |
conftool action : set/pooled=yes; selector: dc=codfw,cluster=dnsbox,service=authdns-update |
[production] |
| 18:42 |
<dzahn@cumin2003> |
END (FAIL) - Cookbook sre.gerrit.localbackup (exit_code=99) Prepare local backup on: gerrit2003.wikimedia.org |
[production] |
| 18:42 |
<dzahn@cumin2003> |
START - Cookbook sre.gerrit.localbackup Prepare local backup on: gerrit2003.wikimedia.org |
[production] |
| 18:40 |
<dzahn@cumin2003> |
END (FAIL) - Cookbook sre.gerrit.localbackup (exit_code=99) Prepare local backup on: gerrit2003.wikimedia.org |
[production] |
| 18:40 |
<dzahn@cumin2003> |
START - Cookbook sre.gerrit.localbackup Prepare local backup on: gerrit2003.wikimedia.org |
[production] |
| 18:40 |
<dzahn@cumin2003> |
END (FAIL) - Cookbook sre.gerrit.localbackup (exit_code=99) Prepare local backup on: gerrit2003.wikimedia.org |
[production] |
| 18:40 |
<dzahn@cumin2003> |
START - Cookbook sre.gerrit.localbackup Prepare local backup on: gerrit2003.wikimedia.org |
[production] |
| 18:40 |
<taavi@cumin1004> |
END (PASS) - Cookbook sre.gerrit.localbackup (exit_code=0) Prepare local backup on: gerrit1003.wikimedia.org |
[production] |
| 18:38 |
<cdanis@cumin1004> |
END (PASS) - Cookbook sre.dns.wipe-cache (exit_code=0) _etcd-client-ssl._tcp.eqsin.wmnet _etcd-client-ssl._tcp.ulsfo.wmnet _etcd-client-ssl._tcp.codfw.wmnet on all recursors |
[production] |
| 18:38 |
<cdanis@cumin1004> |
START - Cookbook sre.dns.wipe-cache _etcd-client-ssl._tcp.eqsin.wmnet _etcd-client-ssl._tcp.ulsfo.wmnet _etcd-client-ssl._tcp.codfw.wmnet on all recursors |
[production] |
| 18:36 |
<taavi@cumin1004> |
END (PASS) - Cookbook sre.gerrit.read-only-toggle (exit_code=0) from gerrit1003.wikimedia.org |
[production] |
| 18:36 |
<taavi@cumin1004> |
START - Cookbook sre.gerrit.read-only-toggle from gerrit1003.wikimedia.org |
[production] |
| 18:36 |
<taavi@cumin1004> |
END (PASS) - Cookbook sre.gerrit.read-only-toggle (exit_code=0) from gerrit2003.wikimedia.org |
[production] |
| 18:36 |
<taavi@cumin1004> |
START - Cookbook sre.gerrit.read-only-toggle from gerrit2003.wikimedia.org |
[production] |
| 18:30 |
<taavi@cumin1004> |
START - Cookbook sre.gerrit.localbackup Prepare local backup on: gerrit1003.wikimedia.org |
[production] |
| 18:29 |
<lerickson@deploy1003> |
helmfile [dse-k8s-eqiad] DONE helmfile.d/services/wdqs-next: apply |
[production] |
| 18:28 |
<lerickson@deploy1003> |
helmfile [dse-k8s-eqiad] START helmfile.d/services/wdqs-next: apply |
[production] |
| 18:14 |
<vriley@cumin1004> |
END (PASS) - Cookbook sre.hosts.reimage (exit_code=0) for host zuul1005.eqiad.wmnet with OS trixie |
[production] |
| 18:08 |
<sukhe@cumin1004> |
END (FAIL) - Cookbook sre.dns.wipe-cache (exit_code=99) idp.wikimedia.org on all recursors |
[production] |
| 18:08 |
<sukhe@cumin1004> |
START - Cookbook sre.dns.wipe-cache idp.wikimedia.org on all recursors |
[production] |
| 18:05 |
<cdanis@cumin1004> |
END (FAIL) - Cookbook sre.dns.wipe-cache (exit_code=99) _etcd-client-ssl._tcp.eqsin.wmnet on all recursors |
[production] |
| 18:05 |
<cdanis@cumin1004> |
START - Cookbook sre.dns.wipe-cache _etcd-client-ssl._tcp.eqsin.wmnet on all recursors |
[production] |
| 18:03 |
<cdanis@cumin1004> |
END (FAIL) - Cookbook sre.dns.wipe-cache (exit_code=99) _etcd-client-ssl._tcp.eqsin.wmnet on all recursors |
[production] |
| 18:03 |
<cdanis@cumin1004> |
START - Cookbook sre.dns.wipe-cache _etcd-client-ssl._tcp.eqsin.wmnet on all recursors |
[production] |
| 18:02 |
<cdanis@cumin1004> |
END (FAIL) - Cookbook sre.dns.wipe-cache (exit_code=99) _etcd-client-ssl._tcp.ulsfo.wmnet on all recursors |
[production] |
| 18:02 |
<cdanis@cumin1004> |
START - Cookbook sre.dns.wipe-cache _etcd-client-ssl._tcp.ulsfo.wmnet on all recursors |
[production] |
| 18:01 |
<cdanis@dns1005> |
END - running authdns-update |
[production] |
| 17:58 |
<cdanis@dns1005> |
START - running authdns-update |
[production] |
| 17:57 |
<vriley@cumin1004> |
END (PASS) - Cookbook sre.hosts.downtime (exit_code=0) for 2:00:00 on zuul1005.eqiad.wmnet with reason: host reimage |
[production] |
| 17:54 |
<taavi@dns1004> |
END - running authdns-update |
[production] |
| 17:53 |
<vriley@cumin1004> |
START - Cookbook sre.hosts.downtime for 2:00:00 on zuul1005.eqiad.wmnet with reason: host reimage |
[production] |
| 17:51 |
<taavi@dns1004> |
START - running authdns-update |
[production] |
| 17:46 |
<taavi@dns1004> |
END - running authdns-update |
[production] |
| 17:43 |
<taavi@dns1004> |
START - running authdns-update |
[production] |
| 17:37 |
<vriley@cumin1004> |
START - Cookbook sre.hosts.reimage for host zuul1005.eqiad.wmnet with OS trixie |
[production] |
| 17:35 |
<rzl@cumin1004> |
START - Cookbook sre.discovery.datacenter pool all active/active services in eqiad: maintenance - T439010 |
[production] |
| 17:35 |
<cdobbins@cumin1004> |
END (PASS) - Cookbook sre.hosts.reimage (exit_code=0) for host ncredir6001.drmrs.wmnet with OS trixie |
[production] |
| 17:35 |
<cdanis@cumin1004> |
END (FAIL) - Cookbook sre.dns.admin (exit_code=99) DNS admin: depool codfw [reason: no reason specified, no task ID specified] |
[production] |
| 17:35 |
<cdanis@cumin1004> |
START - Cookbook sre.dns.admin DNS admin: depool codfw [reason: no reason specified, no task ID specified] |
[production] |
| 17:24 |
<sukhe@cumin1004> |
END (PASS) - Cookbook sre.dns.admin (exit_code=0) DNS admin: depool codfw [reason: no reason specified, no task ID specified] |
[production] |
| 17:23 |
<sukhe@cumin1004> |
START - Cookbook sre.dns.admin DNS admin: depool codfw [reason: no reason specified, no task ID specified] |
[production] |
| 17:21 |
<lerickson@deploy1003> |
helmfile [dse-k8s-eqiad] DONE helmfile.d/services/wdqs: apply |
[production] |