|
2026-09-23
ยง
|
| 18:08 |
<sukhe@cumin1004> |
START - Cookbook sre.dns.wipe-cache idp.wikimedia.org on all recursors |
[production] |
| 18:05 |
<cdanis@cumin1004> |
END (FAIL) - Cookbook sre.dns.wipe-cache (exit_code=99) _etcd-client-ssl._tcp.eqsin.wmnet on all recursors |
[production] |
| 18:05 |
<cdanis@cumin1004> |
START - Cookbook sre.dns.wipe-cache _etcd-client-ssl._tcp.eqsin.wmnet on all recursors |
[production] |
| 18:03 |
<cdanis@cumin1004> |
END (FAIL) - Cookbook sre.dns.wipe-cache (exit_code=99) _etcd-client-ssl._tcp.eqsin.wmnet on all recursors |
[production] |
| 18:03 |
<cdanis@cumin1004> |
START - Cookbook sre.dns.wipe-cache _etcd-client-ssl._tcp.eqsin.wmnet on all recursors |
[production] |
| 18:02 |
<cdanis@cumin1004> |
END (FAIL) - Cookbook sre.dns.wipe-cache (exit_code=99) _etcd-client-ssl._tcp.ulsfo.wmnet on all recursors |
[production] |
| 18:02 |
<cdanis@cumin1004> |
START - Cookbook sre.dns.wipe-cache _etcd-client-ssl._tcp.ulsfo.wmnet on all recursors |
[production] |
| 18:01 |
<cdanis@dns1005> |
END - running authdns-update |
[production] |
| 17:58 |
<cdanis@dns1005> |
START - running authdns-update |
[production] |
| 17:57 |
<vriley@cumin1004> |
END (PASS) - Cookbook sre.hosts.downtime (exit_code=0) for 2:00:00 on zuul1005.eqiad.wmnet with reason: host reimage |
[production] |
| 17:54 |
<taavi@dns1004> |
END - running authdns-update |
[production] |
| 17:53 |
<vriley@cumin1004> |
START - Cookbook sre.hosts.downtime for 2:00:00 on zuul1005.eqiad.wmnet with reason: host reimage |
[production] |
| 17:51 |
<taavi@dns1004> |
START - running authdns-update |
[production] |
| 17:46 |
<taavi@dns1004> |
END - running authdns-update |
[production] |
| 17:43 |
<taavi@dns1004> |
START - running authdns-update |
[production] |
| 17:37 |
<vriley@cumin1004> |
START - Cookbook sre.hosts.reimage for host zuul1005.eqiad.wmnet with OS trixie |
[production] |
| 17:35 |
<rzl@cumin1004> |
START - Cookbook sre.discovery.datacenter pool all active/active services in eqiad: maintenance - T439010 |
[production] |
| 17:35 |
<cdobbins@cumin1004> |
END (PASS) - Cookbook sre.hosts.reimage (exit_code=0) for host ncredir6001.drmrs.wmnet with OS trixie |
[production] |
| 17:35 |
<cdanis@cumin1004> |
END (FAIL) - Cookbook sre.dns.admin (exit_code=99) DNS admin: depool codfw [reason: no reason specified, no task ID specified] |
[production] |
| 17:35 |
<cdanis@cumin1004> |
START - Cookbook sre.dns.admin DNS admin: depool codfw [reason: no reason specified, no task ID specified] |
[production] |
| 17:24 |
<sukhe@cumin1004> |
END (PASS) - Cookbook sre.dns.admin (exit_code=0) DNS admin: depool codfw [reason: no reason specified, no task ID specified] |
[production] |
| 17:23 |
<sukhe@cumin1004> |
START - Cookbook sre.dns.admin DNS admin: depool codfw [reason: no reason specified, no task ID specified] |
[production] |
| 17:21 |
<lerickson@deploy1003> |
helmfile [dse-k8s-eqiad] DONE helmfile.d/services/wdqs: apply |
[production] |
| 17:21 |
<lerickson@deploy1003> |
helmfile [dse-k8s-eqiad] START helmfile.d/services/wdqs: apply |
[production] |
| 17:18 |
<lerickson@deploy1003> |
helmfile [dse-k8s-eqiad] DONE helmfile.d/services/wdqs: apply |
[production] |
| 17:17 |
<lerickson@deploy1003> |
helmfile [dse-k8s-eqiad] START helmfile.d/services/wdqs: apply |
[production] |
| 17:16 |
<lerickson@deploy1003> |
helmfile [dse-k8s-codfw] DONE helmfile.d/services/wdqs-next: apply |
[production] |
| 17:14 |
<sukhe@cumin1004> |
END (PASS) - Cookbook sre.hosts.downtime (exit_code=0) for 2:00:00 on cp2059.codfw.wmnet with reason: host reimage |
[production] |
| 17:11 |
<sukhe@puppetserver1001> |
conftool action : set/pooled=no; selector: name=cp2049.codfw.wmnet |
[production] |
| 17:11 |
<sukhe@puppetserver1001> |
conftool action : set/pooled=no; selector: name=cp2049 |
[production] |
| 17:10 |
<sukhe@cumin1004> |
START - Cookbook sre.hosts.downtime for 2:00:00 on cp2059.codfw.wmnet with reason: host reimage |
[production] |
| 17:07 |
<mutante> |
cloudcontrol2005-dev, cloudcontrol2006-dev, cloudcontrol2010-dev: restart zookeeper, enabled logging (/var/log/zookeeper/zookeeper.log) after gerrit:1342354 |
[production] |
| 17:02 |
<cdobbins@cumin1004> |
END (PASS) - Cookbook sre.hosts.downtime (exit_code=0) for 2:00:00 on ncredir6001.drmrs.wmnet with reason: host reimage |
[production] |
| 16:59 |
<cdobbins@cumin1004> |
START - Cookbook sre.hosts.downtime for 2:00:00 on ncredir6001.drmrs.wmnet with reason: host reimage |
[production] |
| 16:51 |
<sukhe@cumin1004> |
START - Cookbook sre.hosts.reimage for host cp2059.codfw.wmnet with OS trixie |
[production] |
| 16:51 |
<sukhe@cumin1004> |
END (FAIL) - Cookbook sre.hosts.reimage (exit_code=99) for host cp2059.codfw.wmnet with OS trixie |
[production] |
| 16:48 |
<sukhe@cumin1004> |
START - Cookbook sre.hosts.reimage for host cp2059.codfw.wmnet with OS trixie |
[production] |
| 16:39 |
<sukhe@cumin1004> |
END (ERROR) - Cookbook sre.hosts.reimage (exit_code=97) for host cp2059.codfw.wmnet with OS trixie |
[production] |
| 16:35 |
<dzahn@cumin2003> |
START - Cookbook sre.hosts.reimage for host zuul1005.eqiad.wmnet with OS trixie |
[production] |
| 16:34 |
<dzahn@cumin2003> |
END (FAIL) - Cookbook sre.hosts.reimage (exit_code=99) for host zuul1005.eqiad.wmnet with OS trixie |
[production] |
| 16:30 |
<lerickson@deploy1003> |
helmfile [dse-k8s-codfw] START helmfile.d/services/wdqs-next: apply |
[production] |
| 16:29 |
<cdobbins@cumin1004> |
START - Cookbook sre.hosts.reimage for host ncredir6001.drmrs.wmnet with OS trixie |
[production] |
| 16:10 |
<sukhe@cumin1004> |
START - Cookbook sre.hosts.reimage for host cp2059.codfw.wmnet with OS trixie |
[production] |
| 16:10 |
<sukhe@cumin1004> |
END (ERROR) - Cookbook sre.hosts.reimage (exit_code=97) for host cp2059.codfw.wmnet with OS trixie |
[production] |
| 15:55 |
<vgutierrez@cumin1004> |
END (PASS) - Cookbook sre.cdn.roll-upgrade-haproxy (exit_code=0) rolling upgrade of HAProxy on A:cp-upload_ulsfo and A:cp - 3.2.23 upgrade (T438828) |
[production] |
| 15:54 |
<moritzm> |
installing cjose security updates |
[production] |
| 15:54 |
<cdobbins@cumin1004> |
conftool action : set/pooled=yes; selector: name=ncredir7004.* |
[production] |
| 15:53 |
<dancy@deploy1003> |
Finished scap sync-world: testing (duration: 07m 06s) |
[production] |
| 15:52 |
<sukhe@cumin1004> |
START - Cookbook sre.hosts.reimage for host cp2059.codfw.wmnet with OS trixie |
[production] |
| 15:52 |
<sukhe@cumin1004> |
END (ERROR) - Cookbook sre.hosts.reimage (exit_code=97) for host cp2059.codfw.wmnet with OS trixie |
[production] |