2017-01-10 09:01:12 +00:00
|
|
|
package weed_server
|
|
|
|
|
|
|
|
import (
|
2018-07-29 04:02:56 +00:00
|
|
|
"fmt"
|
2017-01-12 21:42:53 +00:00
|
|
|
"net"
|
|
|
|
"strings"
|
2019-01-11 13:47:46 +00:00
|
|
|
"time"
|
2017-01-12 21:42:53 +00:00
|
|
|
|
2018-07-29 04:02:56 +00:00
|
|
|
"github.com/chrislusf/raft"
|
2019-09-12 13:18:21 +00:00
|
|
|
"github.com/chrislusf/seaweedfs/weed/glog"
|
|
|
|
"github.com/chrislusf/seaweedfs/weed/pb/master_pb"
|
|
|
|
"github.com/chrislusf/seaweedfs/weed/storage/needle"
|
|
|
|
"github.com/chrislusf/seaweedfs/weed/topology"
|
2017-01-12 21:42:53 +00:00
|
|
|
"google.golang.org/grpc/peer"
|
2017-01-10 09:01:12 +00:00
|
|
|
)
|
|
|
|
|
2018-05-27 18:58:00 +00:00
|
|
|
func (ms *MasterServer) SendHeartbeat(stream master_pb.Seaweed_SendHeartbeatServer) error {
|
2017-01-10 09:01:12 +00:00
|
|
|
var dn *topology.DataNode
|
|
|
|
t := ms.Topo
|
2018-06-25 06:20:27 +00:00
|
|
|
|
|
|
|
defer func() {
|
|
|
|
if dn != nil {
|
2018-07-28 06:09:55 +00:00
|
|
|
|
2018-06-25 06:20:27 +00:00
|
|
|
glog.V(0).Infof("unregister disconnected volume server %s:%d", dn.Ip, dn.Port)
|
|
|
|
t.UnRegisterDataNode(dn)
|
2018-07-28 06:09:55 +00:00
|
|
|
|
|
|
|
message := &master_pb.VolumeLocation{
|
|
|
|
Url: dn.Url(),
|
|
|
|
PublicUrl: dn.PublicUrl,
|
|
|
|
}
|
|
|
|
for _, v := range dn.GetVolumes() {
|
|
|
|
message.DeletedVids = append(message.DeletedVids, uint32(v.Id))
|
|
|
|
}
|
2019-05-26 08:05:08 +00:00
|
|
|
for _, s := range dn.GetEcShards() {
|
|
|
|
message.DeletedVids = append(message.DeletedVids, uint32(s.VolumeId))
|
|
|
|
}
|
2018-07-28 06:09:55 +00:00
|
|
|
|
|
|
|
if len(message.DeletedVids) > 0 {
|
|
|
|
ms.clientChansLock.RLock()
|
|
|
|
for _, ch := range ms.clientChans {
|
|
|
|
ch <- message
|
|
|
|
}
|
|
|
|
ms.clientChansLock.RUnlock()
|
|
|
|
}
|
|
|
|
|
2018-06-25 06:20:27 +00:00
|
|
|
}
|
|
|
|
}()
|
|
|
|
|
2017-01-10 09:01:12 +00:00
|
|
|
for {
|
|
|
|
heartbeat, err := stream.Recv()
|
2018-08-24 07:30:03 +00:00
|
|
|
if err != nil {
|
|
|
|
return err
|
|
|
|
}
|
|
|
|
|
2019-07-18 06:23:01 +00:00
|
|
|
t.Sequence.SetMax(heartbeat.MaxFileKey)
|
|
|
|
|
2018-08-24 07:30:03 +00:00
|
|
|
if dn == nil {
|
|
|
|
if heartbeat.Ip == "" {
|
|
|
|
if pr, ok := peer.FromContext(stream.Context()); ok {
|
|
|
|
if pr.Addr != net.Addr(nil) {
|
|
|
|
heartbeat.Ip = pr.Addr.String()[0:strings.LastIndex(pr.Addr.String(), ":")]
|
|
|
|
glog.V(0).Infof("remote IP address is detected as %v", heartbeat.Ip)
|
2017-01-12 21:42:53 +00:00
|
|
|
}
|
|
|
|
}
|
2017-01-10 09:01:12 +00:00
|
|
|
}
|
2018-08-24 07:30:03 +00:00
|
|
|
dcName, rackName := t.Configuration.Locate(heartbeat.Ip, heartbeat.DataCenter, heartbeat.Rack)
|
|
|
|
dc := t.GetOrCreateDataCenter(dcName)
|
|
|
|
rack := dc.GetOrCreateRack(rackName)
|
|
|
|
dn = rack.GetOrCreateDataNode(heartbeat.Ip,
|
|
|
|
int(heartbeat.Port), heartbeat.PublicUrl,
|
2019-04-05 02:27:00 +00:00
|
|
|
int64(heartbeat.MaxVolumeCount))
|
2018-08-24 07:30:03 +00:00
|
|
|
glog.V(0).Infof("added volume server %v:%d", heartbeat.GetIp(), heartbeat.GetPort())
|
|
|
|
if err := stream.Send(&master_pb.HeartbeatResponse{
|
2019-06-23 10:08:27 +00:00
|
|
|
VolumeSizeLimit: uint64(ms.option.VolumeSizeLimitMB) * 1024 * 1024,
|
2018-08-24 07:30:03 +00:00
|
|
|
}); err != nil {
|
|
|
|
return err
|
|
|
|
}
|
|
|
|
}
|
2017-01-10 09:01:12 +00:00
|
|
|
|
2019-04-26 16:32:07 +00:00
|
|
|
glog.V(4).Infof("master received heartbeat %s", heartbeat.String())
|
2018-08-24 07:30:03 +00:00
|
|
|
message := &master_pb.VolumeLocation{
|
|
|
|
Url: dn.Url(),
|
|
|
|
PublicUrl: dn.PublicUrl,
|
|
|
|
}
|
2019-04-20 18:35:20 +00:00
|
|
|
if len(heartbeat.NewVolumes) > 0 || len(heartbeat.DeletedVolumes) > 0 {
|
2018-08-24 08:26:56 +00:00
|
|
|
// process delta volume ids if exists for fast volume id updates
|
2019-04-30 03:22:19 +00:00
|
|
|
for _, volInfo := range heartbeat.NewVolumes {
|
2019-04-20 18:35:20 +00:00
|
|
|
message.NewVids = append(message.NewVids, volInfo.Id)
|
|
|
|
}
|
2019-04-30 03:22:19 +00:00
|
|
|
for _, volInfo := range heartbeat.DeletedVolumes {
|
2019-04-20 18:35:20 +00:00
|
|
|
message.DeletedVids = append(message.DeletedVids, volInfo.Id)
|
|
|
|
}
|
|
|
|
// update master internal volume layouts
|
|
|
|
t.IncrementalSyncDataNodeRegistration(heartbeat.NewVolumes, heartbeat.DeletedVolumes, dn)
|
2019-05-23 07:42:28 +00:00
|
|
|
}
|
|
|
|
|
2019-06-05 20:32:33 +00:00
|
|
|
if len(heartbeat.Volumes) > 0 || heartbeat.HasNoVolumes {
|
2018-08-24 08:26:56 +00:00
|
|
|
// process heartbeat.Volumes
|
|
|
|
newVolumes, deletedVolumes := t.SyncDataNodeRegistration(heartbeat.Volumes, dn)
|
|
|
|
|
|
|
|
for _, v := range newVolumes {
|
2019-04-30 03:22:19 +00:00
|
|
|
glog.V(0).Infof("master see new volume %d from %s", uint32(v.Id), dn.Url())
|
2018-08-24 08:26:56 +00:00
|
|
|
message.NewVids = append(message.NewVids, uint32(v.Id))
|
|
|
|
}
|
|
|
|
for _, v := range deletedVolumes {
|
2019-04-20 18:35:20 +00:00
|
|
|
glog.V(0).Infof("master see deleted volume %d from %s", uint32(v.Id), dn.Url())
|
2018-08-24 08:26:56 +00:00
|
|
|
message.DeletedVids = append(message.DeletedVids, uint32(v.Id))
|
|
|
|
}
|
2019-05-23 07:42:28 +00:00
|
|
|
}
|
|
|
|
|
2019-05-26 07:21:17 +00:00
|
|
|
if len(heartbeat.NewEcShards) > 0 || len(heartbeat.DeletedEcShards) > 0 {
|
|
|
|
|
|
|
|
// update master internal volume layouts
|
|
|
|
t.IncrementalSyncDataNodeEcShards(heartbeat.NewEcShards, heartbeat.DeletedEcShards, dn)
|
2019-05-26 07:49:15 +00:00
|
|
|
|
|
|
|
for _, s := range heartbeat.NewEcShards {
|
|
|
|
message.NewVids = append(message.NewVids, s.Id)
|
|
|
|
}
|
|
|
|
for _, s := range heartbeat.DeletedEcShards {
|
|
|
|
if dn.HasVolumesById(needle.VolumeId(s.Id)) {
|
|
|
|
continue
|
|
|
|
}
|
|
|
|
message.DeletedVids = append(message.DeletedVids, s.Id)
|
|
|
|
}
|
|
|
|
|
2019-05-26 07:21:17 +00:00
|
|
|
}
|
|
|
|
|
2019-06-05 20:32:33 +00:00
|
|
|
if len(heartbeat.EcShards) > 0 || heartbeat.HasNoEcShards {
|
2019-06-05 08:30:24 +00:00
|
|
|
glog.V(1).Infof("master recieved ec shards from %s: %+v", dn.Url(), heartbeat.EcShards)
|
2019-05-24 06:47:49 +00:00
|
|
|
newShards, deletedShards := t.SyncDataNodeEcShards(heartbeat.EcShards, dn)
|
|
|
|
|
2019-05-26 07:49:15 +00:00
|
|
|
// broadcast the ec vid changes to master clients
|
|
|
|
for _, s := range newShards {
|
|
|
|
message.NewVids = append(message.NewVids, uint32(s.VolumeId))
|
2019-05-24 06:47:49 +00:00
|
|
|
}
|
2019-05-26 07:49:15 +00:00
|
|
|
for _, s := range deletedShards {
|
|
|
|
if dn.HasVolumesById(s.VolumeId) {
|
|
|
|
continue
|
|
|
|
}
|
|
|
|
message.DeletedVids = append(message.DeletedVids, uint32(s.VolumeId))
|
2019-05-24 06:47:49 +00:00
|
|
|
}
|
|
|
|
|
2018-08-24 07:30:03 +00:00
|
|
|
}
|
2018-07-28 06:09:55 +00:00
|
|
|
|
2018-08-24 07:30:03 +00:00
|
|
|
if len(message.NewVids) > 0 || len(message.DeletedVids) > 0 {
|
|
|
|
ms.clientChansLock.RLock()
|
2019-04-20 18:35:20 +00:00
|
|
|
for host, ch := range ms.clientChans {
|
|
|
|
glog.V(0).Infof("master send to %s: %s", host, message.String())
|
2018-08-24 07:30:03 +00:00
|
|
|
ch <- message
|
2018-07-28 06:09:55 +00:00
|
|
|
}
|
2018-08-24 07:30:03 +00:00
|
|
|
ms.clientChansLock.RUnlock()
|
2017-01-10 09:01:12 +00:00
|
|
|
}
|
2017-01-18 17:34:27 +00:00
|
|
|
|
2017-01-21 21:58:56 +00:00
|
|
|
// tell the volume servers about the leader
|
|
|
|
newLeader, err := t.Leader()
|
2019-03-04 04:43:43 +00:00
|
|
|
if err != nil {
|
|
|
|
return err
|
|
|
|
}
|
|
|
|
if err := stream.Send(&master_pb.HeartbeatResponse{
|
2019-06-17 21:51:47 +00:00
|
|
|
Leader: newLeader,
|
2019-06-23 10:08:27 +00:00
|
|
|
MetricsAddress: ms.option.MetricsAddress,
|
|
|
|
MetricsIntervalSeconds: uint32(ms.option.MetricsIntervalSec),
|
2019-03-04 04:43:43 +00:00
|
|
|
}); err != nil {
|
|
|
|
return err
|
2017-01-18 17:34:27 +00:00
|
|
|
}
|
2017-01-10 09:01:12 +00:00
|
|
|
}
|
|
|
|
}
|
2018-06-01 07:39:39 +00:00
|
|
|
|
2018-07-28 08:30:03 +00:00
|
|
|
// KeepConnected keep a stream gRPC call to the master. Used by clients to know the master is up.
|
|
|
|
// And clients gets the up-to-date list of volume locations
|
2018-06-01 07:39:39 +00:00
|
|
|
func (ms *MasterServer) KeepConnected(stream master_pb.Seaweed_KeepConnectedServer) error {
|
2018-07-28 06:09:55 +00:00
|
|
|
|
|
|
|
req, err := stream.Recv()
|
|
|
|
if err != nil {
|
|
|
|
return err
|
|
|
|
}
|
|
|
|
|
2018-07-28 08:30:03 +00:00
|
|
|
if !ms.Topo.IsLeader() {
|
2019-07-31 08:54:42 +00:00
|
|
|
return ms.informNewLeader(stream)
|
2018-07-28 08:30:03 +00:00
|
|
|
}
|
|
|
|
|
2018-07-28 06:09:55 +00:00
|
|
|
// remember client address
|
|
|
|
ctx := stream.Context()
|
|
|
|
// fmt.Printf("FromContext %+v\n", ctx)
|
|
|
|
pr, ok := peer.FromContext(ctx)
|
|
|
|
if !ok {
|
|
|
|
glog.Error("failed to get peer from ctx")
|
|
|
|
return fmt.Errorf("failed to get peer from ctx")
|
|
|
|
}
|
|
|
|
if pr.Addr == net.Addr(nil) {
|
|
|
|
glog.Error("failed to get peer address")
|
|
|
|
return fmt.Errorf("failed to get peer address")
|
|
|
|
}
|
|
|
|
|
|
|
|
clientName := req.Name + pr.Addr.String()
|
|
|
|
glog.V(0).Infof("+ client %v", clientName)
|
|
|
|
|
|
|
|
messageChan := make(chan *master_pb.VolumeLocation)
|
|
|
|
stopChan := make(chan bool)
|
|
|
|
|
|
|
|
ms.clientChansLock.Lock()
|
|
|
|
ms.clientChans[clientName] = messageChan
|
|
|
|
ms.clientChansLock.Unlock()
|
|
|
|
|
|
|
|
defer func() {
|
|
|
|
glog.V(0).Infof("- client %v", clientName)
|
|
|
|
ms.clientChansLock.Lock()
|
|
|
|
delete(ms.clientChans, clientName)
|
|
|
|
ms.clientChansLock.Unlock()
|
|
|
|
}()
|
|
|
|
|
2018-07-28 08:17:35 +00:00
|
|
|
for _, message := range ms.Topo.ToVolumeLocations() {
|
|
|
|
if err := stream.Send(message); err != nil {
|
|
|
|
return err
|
|
|
|
}
|
|
|
|
}
|
|
|
|
|
2018-07-28 06:09:55 +00:00
|
|
|
go func() {
|
|
|
|
for {
|
|
|
|
_, err := stream.Recv()
|
|
|
|
if err != nil {
|
|
|
|
glog.V(2).Infof("- client %v: %v", clientName, err)
|
|
|
|
stopChan <- true
|
|
|
|
break
|
|
|
|
}
|
2018-06-01 07:39:39 +00:00
|
|
|
}
|
2018-07-28 06:09:55 +00:00
|
|
|
}()
|
|
|
|
|
2019-01-11 13:47:46 +00:00
|
|
|
ticker := time.NewTicker(5 * time.Second)
|
2018-07-28 06:09:55 +00:00
|
|
|
for {
|
|
|
|
select {
|
|
|
|
case message := <-messageChan:
|
|
|
|
if err := stream.Send(message); err != nil {
|
2018-08-24 08:26:56 +00:00
|
|
|
glog.V(0).Infof("=> client %v: %+v", clientName, message)
|
2018-07-28 06:09:55 +00:00
|
|
|
return err
|
|
|
|
}
|
2019-01-11 13:47:46 +00:00
|
|
|
case <-ticker.C:
|
|
|
|
if !ms.Topo.IsLeader() {
|
2019-07-31 08:54:42 +00:00
|
|
|
return ms.informNewLeader(stream)
|
2019-01-11 13:47:46 +00:00
|
|
|
}
|
2018-07-28 06:09:55 +00:00
|
|
|
case <-stopChan:
|
|
|
|
return nil
|
2018-06-01 07:39:39 +00:00
|
|
|
}
|
|
|
|
}
|
2018-07-28 06:09:55 +00:00
|
|
|
|
|
|
|
return nil
|
2018-06-01 07:39:39 +00:00
|
|
|
}
|
2019-07-31 08:54:42 +00:00
|
|
|
|
|
|
|
func (ms *MasterServer) informNewLeader(stream master_pb.Seaweed_KeepConnectedServer) error {
|
|
|
|
leader, err := ms.Topo.Leader()
|
|
|
|
if err != nil {
|
2019-07-31 09:09:04 +00:00
|
|
|
glog.Errorf("topo leader: %v", err)
|
2019-07-31 08:54:42 +00:00
|
|
|
return raft.NotLeaderError
|
|
|
|
}
|
|
|
|
if err := stream.Send(&master_pb.VolumeLocation{
|
|
|
|
Leader: leader,
|
|
|
|
}); err != nil {
|
|
|
|
return err
|
|
|
|
}
|
|
|
|
return nil
|
|
|
|
}
|