2016-07-03 15:10:27 +08:00
|
|
|
package storage
|
|
|
|
|
|
|
|
import (
|
|
|
|
"bytes"
|
|
|
|
"errors"
|
|
|
|
"fmt"
|
|
|
|
"io"
|
|
|
|
"os"
|
|
|
|
"time"
|
|
|
|
|
2019-09-12 21:18:21 +08:00
|
|
|
"github.com/chrislusf/seaweedfs/weed/glog"
|
2019-10-29 15:35:16 +08:00
|
|
|
"github.com/chrislusf/seaweedfs/weed/storage/backend"
|
2019-09-12 21:18:21 +08:00
|
|
|
"github.com/chrislusf/seaweedfs/weed/storage/needle"
|
2019-12-24 04:48:20 +08:00
|
|
|
"github.com/chrislusf/seaweedfs/weed/storage/super_block"
|
2019-09-12 21:18:21 +08:00
|
|
|
. "github.com/chrislusf/seaweedfs/weed/storage/types"
|
2016-07-03 15:10:27 +08:00
|
|
|
)
|
|
|
|
|
2019-01-06 11:52:17 +08:00
|
|
|
var ErrorNotFound = errors.New("not found")
|
|
|
|
|
2016-07-03 15:10:27 +08:00
|
|
|
// isFileUnchanged checks whether this needle to write is same as last one.
|
|
|
|
// It requires serialized access in the same volume.
|
2019-04-19 12:43:36 +08:00
|
|
|
func (v *Volume) isFileUnchanged(n *needle.Needle) bool {
|
2016-07-03 15:10:27 +08:00
|
|
|
if v.Ttl.String() != "" {
|
|
|
|
return false
|
|
|
|
}
|
2019-08-14 16:08:01 +08:00
|
|
|
|
2016-07-03 15:10:27 +08:00
|
|
|
nv, ok := v.nm.Get(n.Id)
|
2019-04-16 12:58:43 +08:00
|
|
|
if ok && !nv.Offset.IsZero() && nv.Size != TombstoneFileSize {
|
2019-04-19 12:43:36 +08:00
|
|
|
oldNeedle := new(needle.Needle)
|
2019-10-29 15:35:16 +08:00
|
|
|
err := oldNeedle.ReadData(v.DataBackend, nv.Offset.ToAcutalOffset(), nv.Size, v.Version())
|
2016-07-03 15:10:27 +08:00
|
|
|
if err != nil {
|
2019-04-16 12:58:43 +08:00
|
|
|
glog.V(0).Infof("Failed to check updated file at offset %d size %d: %v", nv.Offset.ToAcutalOffset(), nv.Size, err)
|
2016-07-03 15:10:27 +08:00
|
|
|
return false
|
|
|
|
}
|
2019-04-22 04:31:45 +08:00
|
|
|
if oldNeedle.Cookie == n.Cookie && oldNeedle.Checksum == n.Checksum && bytes.Equal(oldNeedle.Data, n.Data) {
|
2016-07-03 15:10:27 +08:00
|
|
|
n.DataSize = oldNeedle.DataSize
|
|
|
|
return true
|
|
|
|
}
|
|
|
|
}
|
|
|
|
return false
|
|
|
|
}
|
|
|
|
|
|
|
|
// Destroy removes everything related to this volume
|
|
|
|
func (v *Volume) Destroy() (err error) {
|
2019-08-12 15:53:50 +08:00
|
|
|
if v.isCompacting {
|
|
|
|
err = fmt.Errorf("volume %d is compacting", v.Id)
|
|
|
|
return
|
|
|
|
}
|
2020-05-03 19:22:45 +08:00
|
|
|
close(v.asyncRequestsChan)
|
2019-12-26 08:17:58 +08:00
|
|
|
storageName, storageKey := v.RemoteStorageNameKey()
|
|
|
|
if v.HasRemoteFile() && storageName != "" && storageKey != "" {
|
|
|
|
if backendStorage, found := backend.BackendStorages[storageName]; found {
|
|
|
|
backendStorage.DeleteFile(storageKey)
|
|
|
|
}
|
|
|
|
}
|
2018-11-06 00:53:38 +08:00
|
|
|
v.Close()
|
|
|
|
os.Remove(v.FileName() + ".dat")
|
2018-12-18 12:33:32 +08:00
|
|
|
os.Remove(v.FileName() + ".idx")
|
2019-12-29 03:17:39 +08:00
|
|
|
os.Remove(v.FileName() + ".vif")
|
2019-12-25 02:19:12 +08:00
|
|
|
os.Remove(v.FileName() + ".sdx")
|
2018-08-24 14:33:16 +08:00
|
|
|
os.Remove(v.FileName() + ".cpd")
|
|
|
|
os.Remove(v.FileName() + ".cpx")
|
2019-11-19 09:35:06 +08:00
|
|
|
os.RemoveAll(v.FileName() + ".ldb")
|
2016-07-03 15:10:27 +08:00
|
|
|
return
|
|
|
|
}
|
|
|
|
|
2020-05-03 19:22:45 +08:00
|
|
|
func (v *Volume) asyncRequestAppend(request *needle.AsyncRequest) {
|
|
|
|
v.asyncRequestsChan <- request
|
|
|
|
}
|
|
|
|
|
2020-05-06 21:10:39 +08:00
|
|
|
func (v *Volume) syncWrite(n *needle.Needle) (offset uint64, size uint32, isUnchanged bool, err error) {
|
2020-01-25 12:06:58 +08:00
|
|
|
// glog.V(4).Infof("writing needle %s", needle.NewFileIdFromNeedle(v.Id, n).String())
|
2020-05-06 21:10:39 +08:00
|
|
|
actualSize := needle.GetActualSize(uint32(len(n.Data)), v.Version())
|
|
|
|
|
2016-07-03 15:10:27 +08:00
|
|
|
v.dataFileAccessLock.Lock()
|
|
|
|
defer v.dataFileAccessLock.Unlock()
|
2020-05-06 21:10:39 +08:00
|
|
|
|
|
|
|
if MaxPossibleVolumeSize < v.nm.ContentSize()+uint64(actualSize) {
|
|
|
|
err = fmt.Errorf("volume size limit %d exceeded! current size is %d", MaxPossibleVolumeSize, v.ContentSize())
|
|
|
|
return
|
|
|
|
}
|
2016-07-03 15:10:27 +08:00
|
|
|
if v.isFileUnchanged(n) {
|
|
|
|
size = n.DataSize
|
2019-04-22 04:33:23 +08:00
|
|
|
isUnchanged = true
|
2016-07-03 15:10:27 +08:00
|
|
|
return
|
|
|
|
}
|
|
|
|
|
2019-07-18 14:57:34 +08:00
|
|
|
// check whether existing needle cookie matches
|
|
|
|
nv, ok := v.nm.Get(n.Id)
|
|
|
|
if ok {
|
2019-10-29 15:35:16 +08:00
|
|
|
existingNeedle, _, _, existingNeedleReadErr := needle.ReadNeedleHeader(v.DataBackend, v.Version(), nv.Offset.ToAcutalOffset())
|
2019-07-18 14:57:34 +08:00
|
|
|
if existingNeedleReadErr != nil {
|
|
|
|
err = fmt.Errorf("reading existing needle: %v", existingNeedleReadErr)
|
|
|
|
return
|
|
|
|
}
|
|
|
|
if existingNeedle.Cookie != n.Cookie {
|
|
|
|
glog.V(0).Infof("write cookie mismatch: existing %x, new %x", existingNeedle.Cookie, n.Cookie)
|
|
|
|
err = fmt.Errorf("mismatching cookie %x", n.Cookie)
|
|
|
|
return
|
|
|
|
}
|
|
|
|
}
|
|
|
|
|
|
|
|
// append to dat file
|
2018-07-24 17:44:33 +08:00
|
|
|
n.AppendAtNs = uint64(time.Now().UnixNano())
|
2019-10-29 15:35:16 +08:00
|
|
|
if offset, size, _, err = n.Append(v.DataBackend, v.Version()); err != nil {
|
2016-07-03 15:10:27 +08:00
|
|
|
return
|
|
|
|
}
|
2020-05-06 21:10:39 +08:00
|
|
|
|
2019-04-19 15:39:34 +08:00
|
|
|
v.lastAppendAtNs = n.AppendAtNs
|
2017-01-07 02:22:20 +08:00
|
|
|
|
2019-07-18 14:57:34 +08:00
|
|
|
// add to needle map
|
2019-04-09 10:40:56 +08:00
|
|
|
if !ok || uint64(nv.Offset.ToAcutalOffset()) < offset {
|
|
|
|
if err = v.nm.Put(n.Id, ToOffset(int64(offset)), n.Size); err != nil {
|
2016-07-03 15:10:27 +08:00
|
|
|
glog.V(4).Infof("failed to save in needle map %d: %v", n.Id, err)
|
|
|
|
}
|
|
|
|
}
|
2019-04-19 15:39:34 +08:00
|
|
|
if v.lastModifiedTsSeconds < n.LastModified {
|
|
|
|
v.lastModifiedTsSeconds = n.LastModified
|
2016-07-03 15:10:27 +08:00
|
|
|
}
|
|
|
|
return
|
|
|
|
}
|
|
|
|
|
2020-05-03 19:22:45 +08:00
|
|
|
func (v *Volume) writeNeedle2(n *needle.Needle, fsync bool) (offset uint64, size uint32, isUnchanged bool, err error) {
|
|
|
|
// glog.V(4).Infof("writing needle %s", needle.NewFileIdFromNeedle(v.Id, n).String())
|
|
|
|
if n.Ttl == needle.EMPTY_TTL && v.Ttl != needle.EMPTY_TTL {
|
|
|
|
n.SetHasTtl()
|
|
|
|
n.Ttl = v.Ttl
|
|
|
|
}
|
|
|
|
|
2020-05-06 21:10:39 +08:00
|
|
|
if !fsync {
|
|
|
|
return v.syncWrite(n)
|
|
|
|
} else {
|
|
|
|
asyncRequest := needle.NewAsyncRequest(n, true)
|
|
|
|
// using len(n.Data) here instead of n.Size before n.Size is populated in n.Append()
|
|
|
|
asyncRequest.ActualSize = needle.GetActualSize(uint32(len(n.Data)), v.Version())
|
2020-05-03 19:22:45 +08:00
|
|
|
|
2020-05-06 21:10:39 +08:00
|
|
|
v.asyncRequestAppend(asyncRequest)
|
|
|
|
offset, _, isUnchanged, err = asyncRequest.WaitComplete()
|
2020-05-03 19:22:45 +08:00
|
|
|
|
2020-05-06 21:10:39 +08:00
|
|
|
return
|
|
|
|
}
|
2020-05-03 19:22:45 +08:00
|
|
|
}
|
|
|
|
|
|
|
|
func (v *Volume) doWriteRequest(n *needle.Needle) (offset uint64, size uint32, isUnchanged bool, err error) {
|
|
|
|
// glog.V(4).Infof("writing needle %s", needle.NewFileIdFromNeedle(v.Id, n).String())
|
|
|
|
if v.isFileUnchanged(n) {
|
|
|
|
size = n.DataSize
|
|
|
|
isUnchanged = true
|
|
|
|
return
|
|
|
|
}
|
|
|
|
|
|
|
|
// check whether existing needle cookie matches
|
|
|
|
nv, ok := v.nm.Get(n.Id)
|
|
|
|
if ok {
|
|
|
|
existingNeedle, _, _, existingNeedleReadErr := needle.ReadNeedleHeader(v.DataBackend, v.Version(), nv.Offset.ToAcutalOffset())
|
|
|
|
if existingNeedleReadErr != nil {
|
|
|
|
err = fmt.Errorf("reading existing needle: %v", existingNeedleReadErr)
|
|
|
|
return
|
|
|
|
}
|
|
|
|
if existingNeedle.Cookie != n.Cookie {
|
|
|
|
glog.V(0).Infof("write cookie mismatch: existing %x, new %x", existingNeedle.Cookie, n.Cookie)
|
|
|
|
err = fmt.Errorf("mismatching cookie %x", n.Cookie)
|
|
|
|
return
|
|
|
|
}
|
|
|
|
}
|
|
|
|
|
|
|
|
// append to dat file
|
|
|
|
n.AppendAtNs = uint64(time.Now().UnixNano())
|
|
|
|
if offset, size, _, err = n.Append(v.DataBackend, v.Version()); err != nil {
|
|
|
|
return
|
|
|
|
}
|
|
|
|
v.lastAppendAtNs = n.AppendAtNs
|
|
|
|
|
|
|
|
// add to needle map
|
|
|
|
if !ok || uint64(nv.Offset.ToAcutalOffset()) < offset {
|
|
|
|
if err = v.nm.Put(n.Id, ToOffset(int64(offset)), n.Size); err != nil {
|
|
|
|
glog.V(4).Infof("failed to save in needle map %d: %v", n.Id, err)
|
|
|
|
}
|
|
|
|
}
|
|
|
|
if v.lastModifiedTsSeconds < n.LastModified {
|
|
|
|
v.lastModifiedTsSeconds = n.LastModified
|
|
|
|
}
|
|
|
|
return
|
|
|
|
}
|
|
|
|
|
2020-05-06 21:10:39 +08:00
|
|
|
func (v *Volume) syncDelete(n *needle.Needle) (uint32, error) {
|
2019-04-19 12:43:36 +08:00
|
|
|
glog.V(4).Infof("delete needle %s", needle.NewFileIdFromNeedle(v.Id, n).String())
|
2020-05-06 21:10:39 +08:00
|
|
|
actualSize := needle.GetActualSize(0, v.Version())
|
2016-07-03 15:10:27 +08:00
|
|
|
v.dataFileAccessLock.Lock()
|
|
|
|
defer v.dataFileAccessLock.Unlock()
|
2020-05-06 21:10:39 +08:00
|
|
|
|
|
|
|
if MaxPossibleVolumeSize < v.nm.ContentSize()+uint64(actualSize) {
|
|
|
|
err := fmt.Errorf("volume size limit %d exceeded! current size is %d", MaxPossibleVolumeSize, v.ContentSize())
|
|
|
|
return 0, err
|
|
|
|
}
|
|
|
|
|
2016-07-03 15:10:27 +08:00
|
|
|
nv, ok := v.nm.Get(n.Id)
|
|
|
|
//fmt.Println("key", n.Id, "volume offset", nv.Offset, "data_size", n.Size, "cached size", nv.Size)
|
2017-01-07 02:22:20 +08:00
|
|
|
if ok && nv.Size != TombstoneFileSize {
|
2016-07-03 15:10:27 +08:00
|
|
|
size := nv.Size
|
2019-01-01 07:08:32 +08:00
|
|
|
n.Data = nil
|
|
|
|
n.AppendAtNs = uint64(time.Now().UnixNano())
|
2019-10-29 15:35:16 +08:00
|
|
|
offset, _, _, err := n.Append(v.DataBackend, v.Version())
|
2017-01-20 23:31:11 +08:00
|
|
|
if err != nil {
|
2016-07-03 15:10:27 +08:00
|
|
|
return size, err
|
|
|
|
}
|
2019-04-19 15:39:34 +08:00
|
|
|
v.lastAppendAtNs = n.AppendAtNs
|
2019-04-09 10:40:56 +08:00
|
|
|
if err = v.nm.Delete(n.Id, ToOffset(int64(offset))); err != nil {
|
2016-07-03 15:10:27 +08:00
|
|
|
return size, err
|
|
|
|
}
|
|
|
|
return size, err
|
|
|
|
}
|
|
|
|
return 0, nil
|
|
|
|
}
|
|
|
|
|
2020-05-03 19:22:45 +08:00
|
|
|
func (v *Volume) deleteNeedle2(n *needle.Needle) (uint32, error) {
|
2020-05-06 21:10:39 +08:00
|
|
|
// todo: delete info is always appended no fsync, it may need fsync in future
|
|
|
|
fsync := false
|
2020-05-03 19:22:45 +08:00
|
|
|
|
2020-05-06 21:10:39 +08:00
|
|
|
if !fsync {
|
|
|
|
return v.syncDelete(n)
|
|
|
|
} else {
|
|
|
|
asyncRequest := needle.NewAsyncRequest(n, false)
|
|
|
|
asyncRequest.ActualSize = needle.GetActualSize(0, v.Version())
|
2020-05-03 19:22:45 +08:00
|
|
|
|
2020-05-06 21:10:39 +08:00
|
|
|
v.asyncRequestAppend(asyncRequest)
|
|
|
|
_, size, _, err := asyncRequest.WaitComplete()
|
|
|
|
|
|
|
|
return uint32(size), err
|
|
|
|
}
|
2020-05-03 19:22:45 +08:00
|
|
|
}
|
|
|
|
|
|
|
|
func (v *Volume) doDeleteRequest(n *needle.Needle) (uint32, error) {
|
|
|
|
glog.V(4).Infof("delete needle %s", needle.NewFileIdFromNeedle(v.Id, n).String())
|
|
|
|
nv, ok := v.nm.Get(n.Id)
|
|
|
|
//fmt.Println("key", n.Id, "volume offset", nv.Offset, "data_size", n.Size, "cached size", nv.Size)
|
|
|
|
if ok && nv.Size != TombstoneFileSize {
|
|
|
|
size := nv.Size
|
|
|
|
n.Data = nil
|
|
|
|
n.AppendAtNs = uint64(time.Now().UnixNano())
|
|
|
|
offset, _, _, err := n.Append(v.DataBackend, v.Version())
|
|
|
|
if err != nil {
|
|
|
|
return size, err
|
|
|
|
}
|
|
|
|
v.lastAppendAtNs = n.AppendAtNs
|
|
|
|
if err = v.nm.Delete(n.Id, ToOffset(int64(offset))); err != nil {
|
|
|
|
return size, err
|
|
|
|
}
|
|
|
|
return size, err
|
|
|
|
}
|
|
|
|
return 0, nil
|
|
|
|
}
|
|
|
|
|
2016-07-03 15:10:27 +08:00
|
|
|
// read fills in Needle content by looking up n.Id from NeedleMapper
|
2019-04-19 12:43:36 +08:00
|
|
|
func (v *Volume) readNeedle(n *needle.Needle) (int, error) {
|
2019-12-06 22:59:57 +08:00
|
|
|
v.dataFileAccessLock.RLock()
|
|
|
|
defer v.dataFileAccessLock.RUnlock()
|
2019-08-14 16:08:01 +08:00
|
|
|
|
2016-07-03 15:10:27 +08:00
|
|
|
nv, ok := v.nm.Get(n.Id)
|
2019-04-09 10:40:56 +08:00
|
|
|
if !ok || nv.Offset.IsZero() {
|
2019-07-22 04:49:59 +08:00
|
|
|
return -1, ErrorNotFound
|
2016-07-03 15:10:27 +08:00
|
|
|
}
|
2017-01-07 02:22:20 +08:00
|
|
|
if nv.Size == TombstoneFileSize {
|
2019-01-06 11:52:38 +08:00
|
|
|
return -1, errors.New("already deleted")
|
2017-01-07 02:22:20 +08:00
|
|
|
}
|
2019-01-09 01:03:28 +08:00
|
|
|
if nv.Size == 0 {
|
|
|
|
return 0, nil
|
|
|
|
}
|
2019-10-29 15:35:16 +08:00
|
|
|
err := n.ReadData(v.DataBackend, nv.Offset.ToAcutalOffset(), nv.Size, v.Version())
|
2016-07-03 15:10:27 +08:00
|
|
|
if err != nil {
|
|
|
|
return 0, err
|
|
|
|
}
|
|
|
|
bytesRead := len(n.Data)
|
|
|
|
if !n.HasTtl() {
|
|
|
|
return bytesRead, nil
|
|
|
|
}
|
|
|
|
ttlMinutes := n.Ttl.Minutes()
|
|
|
|
if ttlMinutes == 0 {
|
|
|
|
return bytesRead, nil
|
|
|
|
}
|
|
|
|
if !n.HasLastModifiedDate() {
|
|
|
|
return bytesRead, nil
|
|
|
|
}
|
|
|
|
if uint64(time.Now().Unix()) < n.LastModified+uint64(ttlMinutes*60) {
|
|
|
|
return bytesRead, nil
|
|
|
|
}
|
2019-01-06 11:52:38 +08:00
|
|
|
return -1, ErrorNotFound
|
2016-07-03 15:10:27 +08:00
|
|
|
}
|
|
|
|
|
2020-05-03 19:22:45 +08:00
|
|
|
func (v *Volume) startWorker() {
|
|
|
|
go func() {
|
|
|
|
chanClosed := false
|
|
|
|
for {
|
|
|
|
// chan closed. go thread will exit
|
|
|
|
if chanClosed {
|
|
|
|
break
|
|
|
|
}
|
|
|
|
currentRequests := make([]*needle.AsyncRequest, 0, 128)
|
|
|
|
currentBytesToWrite := int64(0)
|
|
|
|
for {
|
|
|
|
request, ok := <-v.asyncRequestsChan
|
|
|
|
//volume may be closed
|
|
|
|
if !ok {
|
|
|
|
chanClosed = true
|
|
|
|
break
|
|
|
|
}
|
|
|
|
if MaxPossibleVolumeSize < v.ContentSize()+uint64(currentBytesToWrite+request.ActualSize) {
|
|
|
|
request.Complete(0, 0, false,
|
|
|
|
fmt.Errorf("volume size limit %d exceeded! current size is %d", MaxPossibleVolumeSize, v.ContentSize()))
|
|
|
|
break
|
|
|
|
}
|
|
|
|
currentRequests = append(currentRequests, request)
|
|
|
|
currentBytesToWrite += request.ActualSize
|
|
|
|
// submit at most 4M bytes or 128 requests at one time to decrease request delay.
|
|
|
|
// it also need to break if there is no data in channel to avoid io hang.
|
|
|
|
if currentBytesToWrite >= 4*1024*1024 || len(currentRequests) >= 128 || len(v.asyncRequestsChan) == 0 {
|
|
|
|
break
|
|
|
|
}
|
|
|
|
}
|
|
|
|
if len(currentRequests) == 0 {
|
|
|
|
continue
|
|
|
|
}
|
|
|
|
v.dataFileAccessLock.Lock()
|
|
|
|
end, _, e := v.DataBackend.GetStat()
|
|
|
|
if e != nil {
|
|
|
|
for i := 0; i < len(currentRequests); i++ {
|
|
|
|
currentRequests[i].Complete(0, 0, false,
|
|
|
|
fmt.Errorf("cannot read current volume position: %v", e))
|
|
|
|
}
|
|
|
|
v.dataFileAccessLock.Unlock()
|
|
|
|
continue
|
|
|
|
}
|
|
|
|
|
|
|
|
for i := 0; i < len(currentRequests); i++ {
|
|
|
|
if currentRequests[i].IsWriteRequest {
|
|
|
|
offset, size, isUnchanged, err := v.doWriteRequest(currentRequests[i].N)
|
|
|
|
currentRequests[i].UpdateResult(offset, uint64(size), isUnchanged, err)
|
|
|
|
} else {
|
|
|
|
size, err := v.doDeleteRequest(currentRequests[i].N)
|
|
|
|
currentRequests[i].UpdateResult(0, uint64(size), false, err)
|
|
|
|
}
|
|
|
|
}
|
|
|
|
|
2020-05-06 21:10:39 +08:00
|
|
|
// if sync error, data is not reliable, we should mark the completed request as fail and rollback
|
|
|
|
if err := v.DataBackend.Sync(); err != nil {
|
|
|
|
// todo: this may generate dirty data or cause data inconsistent, may be weed need to panic?
|
|
|
|
if te := v.DataBackend.Truncate(end); te != nil {
|
|
|
|
glog.V(0).Infof("Failed to truncate %s back to %d with error: %v", v.DataBackend.Name(), end, te)
|
|
|
|
}
|
|
|
|
for i := 0; i < len(currentRequests); i++ {
|
|
|
|
if currentRequests[i].IsSucceed() {
|
|
|
|
currentRequests[i].UpdateResult(0, 0, false, err)
|
2020-05-03 19:22:45 +08:00
|
|
|
}
|
|
|
|
}
|
|
|
|
}
|
|
|
|
|
|
|
|
for i := 0; i < len(currentRequests); i++ {
|
|
|
|
currentRequests[i].Submit()
|
|
|
|
}
|
|
|
|
v.dataFileAccessLock.Unlock()
|
|
|
|
}
|
|
|
|
}()
|
|
|
|
}
|
|
|
|
|
2019-01-16 17:48:59 +08:00
|
|
|
type VolumeFileScanner interface {
|
2019-12-24 04:48:20 +08:00
|
|
|
VisitSuperBlock(super_block.SuperBlock) error
|
2019-01-16 17:48:59 +08:00
|
|
|
ReadNeedleBody() bool
|
2019-10-22 15:50:30 +08:00
|
|
|
VisitNeedle(n *needle.Needle, offset int64, needleHeader, needleBody []byte) error
|
2019-01-16 17:48:59 +08:00
|
|
|
}
|
|
|
|
|
2019-04-19 12:43:36 +08:00
|
|
|
func ScanVolumeFile(dirname string, collection string, id needle.VolumeId,
|
2016-07-03 15:10:27 +08:00
|
|
|
needleMapKind NeedleMapType,
|
2019-01-16 17:48:59 +08:00
|
|
|
volumeFileScanner VolumeFileScanner) (err error) {
|
2016-07-03 15:10:27 +08:00
|
|
|
var v *Volume
|
|
|
|
if v, err = loadVolumeWithoutIndex(dirname, collection, id, needleMapKind); err != nil {
|
2019-03-26 14:01:53 +08:00
|
|
|
return fmt.Errorf("failed to load volume %d: %v", id, err)
|
2016-07-03 15:10:27 +08:00
|
|
|
}
|
2019-12-29 04:28:58 +08:00
|
|
|
if v.volumeInfo.Version == 0 {
|
|
|
|
if err = volumeFileScanner.VisitSuperBlock(v.SuperBlock); err != nil {
|
|
|
|
return fmt.Errorf("failed to process volume %d super block: %v", id, err)
|
|
|
|
}
|
2016-07-03 15:10:27 +08:00
|
|
|
}
|
2018-11-04 15:28:24 +08:00
|
|
|
defer v.Close()
|
2016-07-03 15:10:27 +08:00
|
|
|
|
|
|
|
version := v.Version()
|
|
|
|
|
2018-06-25 02:37:08 +08:00
|
|
|
offset := int64(v.SuperBlock.BlockSize())
|
2019-03-26 14:01:53 +08:00
|
|
|
|
2019-10-29 15:35:16 +08:00
|
|
|
return ScanVolumeFileFrom(version, v.DataBackend, offset, volumeFileScanner)
|
2019-03-26 14:01:53 +08:00
|
|
|
}
|
|
|
|
|
2019-11-29 10:33:18 +08:00
|
|
|
func ScanVolumeFileFrom(version needle.Version, datBackend backend.BackendStorageFile, offset int64, volumeFileScanner VolumeFileScanner) (err error) {
|
2019-10-29 15:35:16 +08:00
|
|
|
n, nh, rest, e := needle.ReadNeedleHeader(datBackend, version, offset)
|
2016-07-03 15:10:27 +08:00
|
|
|
if e != nil {
|
2019-03-26 14:01:53 +08:00
|
|
|
if e == io.EOF {
|
|
|
|
return nil
|
|
|
|
}
|
2019-12-09 11:44:16 +08:00
|
|
|
return fmt.Errorf("cannot read %s at offset %d: %v", datBackend.Name(), offset, e)
|
2016-07-03 15:10:27 +08:00
|
|
|
}
|
|
|
|
for n != nil {
|
2019-10-22 15:50:30 +08:00
|
|
|
var needleBody []byte
|
2019-01-16 17:48:59 +08:00
|
|
|
if volumeFileScanner.ReadNeedleBody() {
|
2019-10-29 15:35:16 +08:00
|
|
|
if needleBody, err = n.ReadNeedleBody(datBackend, version, offset+NeedleHeaderSize, rest); err != nil {
|
2016-07-03 15:10:27 +08:00
|
|
|
glog.V(0).Infof("cannot read needle body: %v", err)
|
|
|
|
//err = fmt.Errorf("cannot read needle body: %v", err)
|
|
|
|
//return
|
|
|
|
}
|
|
|
|
}
|
2019-10-22 15:50:30 +08:00
|
|
|
err := volumeFileScanner.VisitNeedle(n, offset, nh, needleBody)
|
2018-07-15 11:51:17 +08:00
|
|
|
if err == io.EOF {
|
|
|
|
return nil
|
|
|
|
}
|
|
|
|
if err != nil {
|
2016-07-03 15:10:27 +08:00
|
|
|
glog.V(0).Infof("visit needle error: %v", err)
|
2019-07-18 14:22:01 +08:00
|
|
|
return fmt.Errorf("visit needle error: %v", err)
|
2016-07-03 15:10:27 +08:00
|
|
|
}
|
2019-04-19 15:39:34 +08:00
|
|
|
offset += NeedleHeaderSize + rest
|
2016-07-03 15:10:27 +08:00
|
|
|
glog.V(4).Infof("==> new entry offset %d", offset)
|
2019-10-29 15:35:16 +08:00
|
|
|
if n, nh, rest, err = needle.ReadNeedleHeader(datBackend, version, offset); err != nil {
|
2019-04-18 15:18:29 +08:00
|
|
|
if err == io.EOF {
|
|
|
|
return nil
|
|
|
|
}
|
|
|
|
return fmt.Errorf("cannot read needle header at offset %d: %v", offset, err)
|
|
|
|
}
|
|
|
|
glog.V(4).Infof("new entry needle size:%d rest:%d", n.Size, rest)
|
|
|
|
}
|
|
|
|
return nil
|
|
|
|
}
|