black_hole_test.go 6.3 KB

123456789101112131415161718192021222324252627282930313233343536373839404142434445464748495051525354555657585960616263646566676869707172737475767778798081828384858687888990919293949596979899100101102103104105106107108109110111112113114115116117118119120121122123124125126127128129130131132133134135136137138139140141142143144145146147148149150151152153154155156157158159160161162163164165166167168169170171172173174175176177178179180181182183184185186187188189190191192193194195196197198199200201202203204205206207208209210211212213214
  1. // Copyright 2017 The etcd Authors
  2. //
  3. // Licensed under the Apache License, Version 2.0 (the "License");
  4. // you may not use this file except in compliance with the License.
  5. // You may obtain a copy of the License at
  6. //
  7. // http://www.apache.org/licenses/LICENSE-2.0
  8. //
  9. // Unless required by applicable law or agreed to in writing, software
  10. // distributed under the License is distributed on an "AS IS" BASIS,
  11. // WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
  12. // See the License for the specific language governing permissions and
  13. // limitations under the License.
  14. // +build !cluster_proxy
  15. package integration
  16. import (
  17. "context"
  18. "testing"
  19. "time"
  20. "go.etcd.io/etcd/clientv3"
  21. "go.etcd.io/etcd/etcdserver/api/v3rpc/rpctypes"
  22. "go.etcd.io/etcd/integration"
  23. "go.etcd.io/etcd/pkg/testutil"
  24. "google.golang.org/grpc"
  25. )
  26. // TestBalancerUnderBlackholeKeepAliveWatch tests when watch discovers it cannot talk to
  27. // blackholed endpoint, client balancer switches to healthy one.
  28. // TODO: test server-to-client keepalive ping
  29. func TestBalancerUnderBlackholeKeepAliveWatch(t *testing.T) {
  30. defer testutil.AfterTest(t)
  31. clus := integration.NewClusterV3(t, &integration.ClusterConfig{
  32. Size: 2,
  33. GRPCKeepAliveMinTime: 1 * time.Millisecond, // avoid too_many_pings
  34. })
  35. defer clus.Terminate(t)
  36. eps := []string{clus.Members[0].GRPCAddr(), clus.Members[1].GRPCAddr()}
  37. ccfg := clientv3.Config{
  38. Endpoints: []string{eps[0]},
  39. DialTimeout: 1 * time.Second,
  40. DialOptions: []grpc.DialOption{grpc.WithBlock()},
  41. DialKeepAliveTime: 1 * time.Second,
  42. DialKeepAliveTimeout: 500 * time.Millisecond,
  43. }
  44. // gRPC internal implementation related.
  45. pingInterval := ccfg.DialKeepAliveTime + ccfg.DialKeepAliveTimeout
  46. // 3s for slow machine to process watch and reset connections
  47. // TODO: only send healthy endpoint to gRPC so gRPC wont waste time to
  48. // dial for unhealthy endpoint.
  49. // then we can reduce 3s to 1s.
  50. timeout := pingInterval + integration.RequestWaitTimeout
  51. cli, err := clientv3.New(ccfg)
  52. if err != nil {
  53. t.Fatal(err)
  54. }
  55. defer cli.Close()
  56. wch := cli.Watch(context.Background(), "foo", clientv3.WithCreatedNotify())
  57. if _, ok := <-wch; !ok {
  58. t.Fatalf("watch failed on creation")
  59. }
  60. // endpoint can switch to eps[1] when it detects the failure of eps[0]
  61. cli.SetEndpoints(eps...)
  62. clus.Members[0].Blackhole()
  63. if _, err = clus.Client(1).Put(context.TODO(), "foo", "bar"); err != nil {
  64. t.Fatal(err)
  65. }
  66. select {
  67. case <-wch:
  68. case <-time.After(timeout):
  69. t.Error("took too long to receive watch events")
  70. }
  71. clus.Members[0].Unblackhole()
  72. // waiting for moving eps[0] out of unhealthy, so that it can be re-pined.
  73. time.Sleep(ccfg.DialTimeout)
  74. clus.Members[1].Blackhole()
  75. // make sure client[0] can connect to eps[0] after remove the blackhole.
  76. if _, err = clus.Client(0).Get(context.TODO(), "foo"); err != nil {
  77. t.Fatal(err)
  78. }
  79. if _, err = clus.Client(0).Put(context.TODO(), "foo", "bar1"); err != nil {
  80. t.Fatal(err)
  81. }
  82. select {
  83. case <-wch:
  84. case <-time.After(timeout):
  85. t.Error("took too long to receive watch events")
  86. }
  87. }
  88. func TestBalancerUnderBlackholeNoKeepAlivePut(t *testing.T) {
  89. testBalancerUnderBlackholeNoKeepAlive(t, func(cli *clientv3.Client, ctx context.Context) error {
  90. _, err := cli.Put(ctx, "foo", "bar")
  91. if isClientTimeout(err) || isServerCtxTimeout(err) || err == rpctypes.ErrTimeout {
  92. return errExpected
  93. }
  94. return err
  95. })
  96. }
  97. func TestBalancerUnderBlackholeNoKeepAliveDelete(t *testing.T) {
  98. testBalancerUnderBlackholeNoKeepAlive(t, func(cli *clientv3.Client, ctx context.Context) error {
  99. _, err := cli.Delete(ctx, "foo")
  100. if isClientTimeout(err) || isServerCtxTimeout(err) || err == rpctypes.ErrTimeout {
  101. return errExpected
  102. }
  103. return err
  104. })
  105. }
  106. func TestBalancerUnderBlackholeNoKeepAliveTxn(t *testing.T) {
  107. testBalancerUnderBlackholeNoKeepAlive(t, func(cli *clientv3.Client, ctx context.Context) error {
  108. _, err := cli.Txn(ctx).
  109. If(clientv3.Compare(clientv3.Version("foo"), "=", 0)).
  110. Then(clientv3.OpPut("foo", "bar")).
  111. Else(clientv3.OpPut("foo", "baz")).Commit()
  112. if isClientTimeout(err) || isServerCtxTimeout(err) || err == rpctypes.ErrTimeout {
  113. return errExpected
  114. }
  115. return err
  116. })
  117. }
  118. func TestBalancerUnderBlackholeNoKeepAliveLinearizableGet(t *testing.T) {
  119. testBalancerUnderBlackholeNoKeepAlive(t, func(cli *clientv3.Client, ctx context.Context) error {
  120. _, err := cli.Get(ctx, "a")
  121. if isClientTimeout(err) || isServerCtxTimeout(err) || err == rpctypes.ErrTimeout {
  122. return errExpected
  123. }
  124. return err
  125. })
  126. }
  127. func TestBalancerUnderBlackholeNoKeepAliveSerializableGet(t *testing.T) {
  128. testBalancerUnderBlackholeNoKeepAlive(t, func(cli *clientv3.Client, ctx context.Context) error {
  129. _, err := cli.Get(ctx, "a", clientv3.WithSerializable())
  130. if isClientTimeout(err) || isServerCtxTimeout(err) {
  131. return errExpected
  132. }
  133. return err
  134. })
  135. }
  136. // testBalancerUnderBlackholeNoKeepAlive ensures that first request to blackholed endpoint
  137. // fails due to context timeout, but succeeds on next try, with endpoint switch.
  138. func testBalancerUnderBlackholeNoKeepAlive(t *testing.T, op func(*clientv3.Client, context.Context) error) {
  139. defer testutil.AfterTest(t)
  140. clus := integration.NewClusterV3(t, &integration.ClusterConfig{
  141. Size: 2,
  142. SkipCreatingClient: true,
  143. })
  144. defer clus.Terminate(t)
  145. eps := []string{clus.Members[0].GRPCAddr(), clus.Members[1].GRPCAddr()}
  146. ccfg := clientv3.Config{
  147. Endpoints: []string{eps[0]},
  148. DialTimeout: 1 * time.Second,
  149. DialOptions: []grpc.DialOption{grpc.WithBlock()},
  150. }
  151. cli, err := clientv3.New(ccfg)
  152. if err != nil {
  153. t.Fatal(err)
  154. }
  155. defer cli.Close()
  156. // wait for eps[0] to be pinned
  157. mustWaitPinReady(t, cli)
  158. // add all eps to list, so that when the original pined one fails
  159. // the client can switch to other available eps
  160. cli.SetEndpoints(eps...)
  161. // blackhole eps[0]
  162. clus.Members[0].Blackhole()
  163. // With round robin balancer, client will make a request to a healthy endpoint
  164. // within a few requests.
  165. // TODO: first operation can succeed
  166. // when gRPC supports better retry on non-delivered request
  167. for i := 0; i < 5; i++ {
  168. ctx, cancel := context.WithTimeout(context.Background(), time.Second*5)
  169. err = op(cli, ctx)
  170. cancel()
  171. if err == nil {
  172. break
  173. } else if err == errExpected {
  174. t.Logf("#%d: current error %v", i, err)
  175. } else {
  176. t.Errorf("#%d: failed with error %v", i, err)
  177. }
  178. }
  179. if err != nil {
  180. t.Fatal(err)
  181. }
  182. }