forked from rook/rook
We should not use .items[0].metadata.name if the array length is 0. This is the case when nothing has been initialized yet. Instead, we should use .items[*].metadata.name, the wildcard ensures to always return 0 even if nothing is present yet. Fixes: https://github.com/rook/rook/issues/8676 Signed-off-by: Sébastien Han <seb@redhat.com>
178 lines
5.1 KiB
Bash
Executable File
178 lines
5.1 KiB
Bash
Executable File
#!/usr/bin/env bash
|
|
|
|
# Copyright 2021 The Rook Authors. All rights reserved.
|
|
#
|
|
# Licensed under the Apache License, Version 2.0 (the "License");
|
|
# you may not use this file except in compliance with the License.
|
|
# You may obtain a copy of the License at
|
|
#
|
|
# http://www.apache.org/licenses/LICENSE-2.0
|
|
#
|
|
# Unless required by applicable law or agreed to in writing, software
|
|
# distributed under the License is distributed on an "AS IS" BASIS,
|
|
# WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
|
|
# See the License for the specific language governing permissions and
|
|
# limitations under the License.
|
|
|
|
set -xEe
|
|
|
|
: "${DAEMON_TO_VALIDATE:=${1}}"
|
|
if [ -z "$DAEMON_TO_VALIDATE" ]; then
|
|
DAEMON_TO_VALIDATE=all
|
|
fi
|
|
OSD_COUNT=$2
|
|
|
|
#############
|
|
# FUNCTIONS #
|
|
#############
|
|
EXEC_COMMAND="kubectl -n rook-ceph exec $(kubectl get pod -l app=rook-ceph-tools -n rook-ceph -o jsonpath='{.items[*].metadata.name}') -- ceph --connect-timeout 3"
|
|
|
|
trap display_status SIGINT ERR
|
|
|
|
function wait_for_daemon () {
|
|
timeout=90
|
|
daemon_to_test=$1
|
|
while [ $timeout -ne 0 ]; do
|
|
if eval $daemon_to_test; then
|
|
return 0
|
|
fi
|
|
sleep 1
|
|
let timeout=timeout-1
|
|
done
|
|
|
|
return 1
|
|
}
|
|
|
|
function test_demo_mon {
|
|
# shellcheck disable=SC2046
|
|
return $(wait_for_daemon "$EXEC_COMMAND -s | grep -sq quorum")
|
|
}
|
|
|
|
function test_demo_mgr {
|
|
# shellcheck disable=SC2046
|
|
return $(wait_for_daemon "$EXEC_COMMAND -s | grep -sq 'mgr:'")
|
|
}
|
|
|
|
function test_demo_osd {
|
|
# shellcheck disable=SC2046
|
|
return $(wait_for_daemon "$EXEC_COMMAND -s | grep -sq \"$OSD_COUNT osds: $OSD_COUNT up.*, $OSD_COUNT in.*\"")
|
|
}
|
|
|
|
function test_demo_rgw {
|
|
# shellcheck disable=SC2046
|
|
return $(wait_for_daemon "$EXEC_COMMAND -s | grep -sq 'rgw:'")
|
|
}
|
|
|
|
function test_demo_mds {
|
|
echo "Waiting for the MDS to be ready"
|
|
# NOTE: metadata server always takes up to 5 sec to run
|
|
# so we first check if the pools exit, from that we assume that
|
|
# the process will start. We stop waiting after 10 seconds.
|
|
# shellcheck disable=SC2046
|
|
return $(wait_for_daemon "$EXEC_COMMAND osd dump | grep -sq cephfs && $EXEC_COMMAND -s | grep -sq up")
|
|
}
|
|
|
|
function test_demo_rbd_mirror {
|
|
# shellcheck disable=SC2046
|
|
return $(wait_for_daemon "$EXEC_COMMAND -s | grep -sq 'rbd-mirror:'")
|
|
}
|
|
|
|
function test_demo_fs_mirror {
|
|
# shellcheck disable=SC2046
|
|
return $(wait_for_daemon "$EXEC_COMMAND -s | grep -sq 'cephfs-mirror:'")
|
|
}
|
|
|
|
function test_demo_pool {
|
|
# shellcheck disable=SC2046
|
|
return $(wait_for_daemon "$EXEC_COMMAND -s | grep -sq '11 pools'")
|
|
}
|
|
|
|
function test_csi {
|
|
# shellcheck disable=SC2046
|
|
timeout 90 sh -c 'until [ $(kubectl -n rook-ceph get pods --field-selector=status.phase=Running|grep -c ^csi-) -eq 4 ]; do sleep 1; done'
|
|
if [ $? -eq 0 ]; then
|
|
return 0
|
|
fi
|
|
return 1
|
|
}
|
|
|
|
function display_status {
|
|
$EXEC_COMMAND -s > test/ceph-status.txt
|
|
$EXEC_COMMAND osd dump > test/ceph-osd-dump.txt
|
|
$EXEC_COMMAND report > test/ceph-report.txt
|
|
|
|
kubectl -n rook-ceph logs deploy/rook-ceph-operator > test/operator-logs.txt
|
|
kubectl -n rook-ceph get pods -o wide > test/pods-list.txt
|
|
kubectl -n rook-ceph describe job/"$(kubectl -n rook-ceph get job -l app=rook-ceph-osd-prepare -o jsonpath='{.items[*].metadata.name}')" > test/osd-prepare-describe.txt
|
|
kubectl -n rook-ceph log job/"$(kubectl -n rook-ceph get job -l app=rook-ceph-osd-prepare -o jsonpath='{.items[*].metadata.name}')" > test/osd-prepare-logs.txt
|
|
kubectl -n rook-ceph describe deploy/rook-ceph-osd-0 > test/rook-ceph-osd-0-describe.txt
|
|
kubectl -n rook-ceph describe deploy/rook-ceph-osd-1 > test/rook-ceph-osd-1-describe.txt
|
|
kubectl -n rook-ceph logs deploy/rook-ceph-osd-0 --all-containers > test/rook-ceph-osd-0-logs.txt
|
|
kubectl -n rook-ceph logs deploy/rook-ceph-osd-1 --all-containers > test/rook-ceph-osd-1-logs.txt
|
|
kubectl get all -n rook-ceph -o wide > test/cluster-wide.txt
|
|
kubectl get all -n rook-ceph -o yaml > test/cluster-yaml.txt
|
|
kubectl -n rook-ceph get cephcluster -o yaml > test/cephcluster.txt
|
|
sudo lsblk | sudo tee -a test/lsblk.txt
|
|
}
|
|
|
|
########
|
|
# MAIN #
|
|
########
|
|
test_csi
|
|
test_demo_mon
|
|
test_demo_mgr
|
|
|
|
if [[ "$DAEMON_TO_VALIDATE" == "all" ]]; then
|
|
daemons_list="osd mds rgw rbd_mirror fs_mirror"
|
|
else
|
|
# change commas to space
|
|
comma_to_space=${DAEMON_TO_VALIDATE//,/ }
|
|
|
|
# transform to an array
|
|
IFS=" " read -r -a array <<< "$comma_to_space"
|
|
|
|
# sort and remove potential duplicate
|
|
daemons_list=$(echo "${array[@]}" | tr ' ' '\n' | sort -u | tr '\n' ' ')
|
|
fi
|
|
|
|
for daemon in $daemons_list; do
|
|
case "$daemon" in
|
|
mon)
|
|
continue
|
|
;;
|
|
mgr)
|
|
continue
|
|
;;
|
|
osd)
|
|
test_demo_osd
|
|
;;
|
|
mds)
|
|
test_demo_mds
|
|
;;
|
|
rgw)
|
|
test_demo_rgw
|
|
;;
|
|
rbd_mirror)
|
|
test_demo_rbd_mirror
|
|
;;
|
|
fs_mirror)
|
|
test_demo_fs_mirror
|
|
;;
|
|
*)
|
|
log "ERROR: unknown daemon to validate!"
|
|
log "Available daemon are: mon mgr osd mds rgw rbd_mirror fs_mirror"
|
|
exit 1
|
|
;;
|
|
esac
|
|
done
|
|
|
|
echo "Ceph is up and running, have a look!"
|
|
$EXEC_COMMAND -s
|
|
kubectl -n rook-ceph get pods
|
|
kubectl -n rook-ceph logs "$(kubectl -n rook-ceph -l app=rook-ceph-operator get pods -o jsonpath='{.items[*].metadata.name}')"
|
|
kubectl -n rook-ceph get cephcluster -o yaml
|
|
|
|
set +eE
|
|
display_status
|
|
set -eE
|