2016-02-09 15:12:00 -08:00
|
|
|
// Copyright (c) 2011-present, Facebook, Inc. All rights reserved.
|
2017-07-15 16:03:42 -07:00
|
|
|
// This source code is licensed under both the GPLv2 (found in the
|
|
|
|
// COPYING file in the root directory) and Apache 2.0 License
|
|
|
|
// (found in the LICENSE.Apache file in the root directory).
|
2015-02-23 17:49:23 -08:00
|
|
|
// Copyright (c) 2011 The LevelDB Authors. All rights reserved.
|
|
|
|
// Use of this source code is governed by a BSD-style license that can be
|
|
|
|
// found in the LICENSE file. See the AUTHORS file for names of contributors.
|
|
|
|
|
|
|
|
#pragma once
|
|
|
|
|
2017-04-10 15:38:34 -07:00
|
|
|
#include <cstddef>
|
|
|
|
|
2020-02-20 12:07:53 -08:00
|
|
|
#include "rocksdb/rocksdb_namespace.h"
|
|
|
|
|
|
|
|
namespace ROCKSDB_NAMESPACE {
|
2015-02-23 17:49:23 -08:00
|
|
|
|
|
|
|
class Slice;
|
2015-07-10 20:15:45 -07:00
|
|
|
class Status;
|
2015-02-23 17:49:23 -08:00
|
|
|
class ColumnFamilyHandle;
|
|
|
|
class WriteBatch;
|
|
|
|
struct SliceParts;
|
|
|
|
|
|
|
|
// Abstract base class that defines the basic interface for a write batch.
|
|
|
|
// See WriteBatch for a basic implementation and WrithBatchWithIndex for an
|
2018-03-08 10:18:34 -08:00
|
|
|
// indexed implementation.
|
2015-02-23 17:49:23 -08:00
|
|
|
class WriteBatchBase {
|
|
|
|
public:
|
|
|
|
virtual ~WriteBatchBase() {}
|
|
|
|
|
|
|
|
// Store the mapping "key->value" in the database.
|
2017-04-10 15:38:34 -07:00
|
|
|
virtual Status Put(ColumnFamilyHandle* column_family, const Slice& key,
|
|
|
|
const Slice& value) = 0;
|
|
|
|
virtual Status Put(const Slice& key, const Slice& value) = 0;
|
2015-02-23 17:49:23 -08:00
|
|
|
|
|
|
|
// Variant of Put() that gathers output like writev(2). The key and value
|
2017-05-17 23:03:54 -07:00
|
|
|
// that will be written to the database are concatenations of arrays of
|
2015-02-23 17:49:23 -08:00
|
|
|
// slices.
|
2017-04-10 15:38:34 -07:00
|
|
|
virtual Status Put(ColumnFamilyHandle* column_family, const SliceParts& key,
|
|
|
|
const SliceParts& value);
|
|
|
|
virtual Status Put(const SliceParts& key, const SliceParts& value);
|
2015-02-23 17:49:23 -08:00
|
|
|
|
|
|
|
// Merge "value" with the existing value of "key" in the database.
|
|
|
|
// "key->merge(existing, value)"
|
2017-04-10 15:38:34 -07:00
|
|
|
virtual Status Merge(ColumnFamilyHandle* column_family, const Slice& key,
|
|
|
|
const Slice& value) = 0;
|
|
|
|
virtual Status Merge(const Slice& key, const Slice& value) = 0;
|
2015-02-23 17:49:23 -08:00
|
|
|
|
2015-05-27 16:59:22 -07:00
|
|
|
// variant that takes SliceParts
|
2017-04-10 15:38:34 -07:00
|
|
|
virtual Status Merge(ColumnFamilyHandle* column_family, const SliceParts& key,
|
|
|
|
const SliceParts& value);
|
|
|
|
virtual Status Merge(const SliceParts& key, const SliceParts& value);
|
2015-05-27 16:59:22 -07:00
|
|
|
|
2015-02-23 17:49:23 -08:00
|
|
|
// If the database contains a mapping for "key", erase it. Else do nothing.
|
2017-04-10 15:38:34 -07:00
|
|
|
virtual Status Delete(ColumnFamilyHandle* column_family,
|
|
|
|
const Slice& key) = 0;
|
|
|
|
virtual Status Delete(const Slice& key) = 0;
|
2015-02-23 17:49:23 -08:00
|
|
|
|
|
|
|
// variant that takes SliceParts
|
2017-04-10 15:38:34 -07:00
|
|
|
virtual Status Delete(ColumnFamilyHandle* column_family,
|
|
|
|
const SliceParts& key);
|
|
|
|
virtual Status Delete(const SliceParts& key);
|
2015-02-23 17:49:23 -08:00
|
|
|
|
Support for SingleDelete()
Summary:
This patch fixes #7460559. It introduces SingleDelete as a new database
operation. This operation can be used to delete keys that were never
overwritten (no put following another put of the same key). If an overwritten
key is single deleted the behavior is undefined. Single deletion of a
non-existent key has no effect but multiple consecutive single deletions are
not allowed (see limitations).
In contrast to the conventional Delete() operation, the deletion entry is
removed along with the value when the two are lined up in a compaction. Note:
The semantics are similar to @igor's prototype that allowed to have this
behavior on the granularity of a column family (
https://reviews.facebook.net/D42093 ). This new patch, however, is more
aggressive when it comes to removing tombstones: It removes the SingleDelete
together with the value whenever there is no snapshot between them while the
older patch only did this when the sequence number of the deletion was older
than the earliest snapshot.
Most of the complex additions are in the Compaction Iterator, all other changes
should be relatively straightforward. The patch also includes basic support for
single deletions in db_stress and db_bench.
Limitations:
- Not compatible with cuckoo hash tables
- Single deletions cannot be used in combination with merges and normal
deletions on the same key (other keys are not affected by this)
- Consecutive single deletions are currently not allowed (and older version of
this patch supported this so it could be resurrected if needed)
Test Plan: make all check
Reviewers: yhchiang, sdong, rven, anthony, yoshinorim, igor
Reviewed By: igor
Subscribers: maykov, dhruba, leveldb
Differential Revision: https://reviews.facebook.net/D43179
2015-09-17 11:42:56 -07:00
|
|
|
// If the database contains a mapping for "key", erase it. Expects that the
|
|
|
|
// key was not overwritten. Else do nothing.
|
2017-04-10 15:38:34 -07:00
|
|
|
virtual Status SingleDelete(ColumnFamilyHandle* column_family,
|
|
|
|
const Slice& key) = 0;
|
|
|
|
virtual Status SingleDelete(const Slice& key) = 0;
|
Support for SingleDelete()
Summary:
This patch fixes #7460559. It introduces SingleDelete as a new database
operation. This operation can be used to delete keys that were never
overwritten (no put following another put of the same key). If an overwritten
key is single deleted the behavior is undefined. Single deletion of a
non-existent key has no effect but multiple consecutive single deletions are
not allowed (see limitations).
In contrast to the conventional Delete() operation, the deletion entry is
removed along with the value when the two are lined up in a compaction. Note:
The semantics are similar to @igor's prototype that allowed to have this
behavior on the granularity of a column family (
https://reviews.facebook.net/D42093 ). This new patch, however, is more
aggressive when it comes to removing tombstones: It removes the SingleDelete
together with the value whenever there is no snapshot between them while the
older patch only did this when the sequence number of the deletion was older
than the earliest snapshot.
Most of the complex additions are in the Compaction Iterator, all other changes
should be relatively straightforward. The patch also includes basic support for
single deletions in db_stress and db_bench.
Limitations:
- Not compatible with cuckoo hash tables
- Single deletions cannot be used in combination with merges and normal
deletions on the same key (other keys are not affected by this)
- Consecutive single deletions are currently not allowed (and older version of
this patch supported this so it could be resurrected if needed)
Test Plan: make all check
Reviewers: yhchiang, sdong, rven, anthony, yoshinorim, igor
Reviewed By: igor
Subscribers: maykov, dhruba, leveldb
Differential Revision: https://reviews.facebook.net/D43179
2015-09-17 11:42:56 -07:00
|
|
|
|
|
|
|
// variant that takes SliceParts
|
2017-04-10 15:38:34 -07:00
|
|
|
virtual Status SingleDelete(ColumnFamilyHandle* column_family,
|
|
|
|
const SliceParts& key);
|
|
|
|
virtual Status SingleDelete(const SliceParts& key);
|
Support for SingleDelete()
Summary:
This patch fixes #7460559. It introduces SingleDelete as a new database
operation. This operation can be used to delete keys that were never
overwritten (no put following another put of the same key). If an overwritten
key is single deleted the behavior is undefined. Single deletion of a
non-existent key has no effect but multiple consecutive single deletions are
not allowed (see limitations).
In contrast to the conventional Delete() operation, the deletion entry is
removed along with the value when the two are lined up in a compaction. Note:
The semantics are similar to @igor's prototype that allowed to have this
behavior on the granularity of a column family (
https://reviews.facebook.net/D42093 ). This new patch, however, is more
aggressive when it comes to removing tombstones: It removes the SingleDelete
together with the value whenever there is no snapshot between them while the
older patch only did this when the sequence number of the deletion was older
than the earliest snapshot.
Most of the complex additions are in the Compaction Iterator, all other changes
should be relatively straightforward. The patch also includes basic support for
single deletions in db_stress and db_bench.
Limitations:
- Not compatible with cuckoo hash tables
- Single deletions cannot be used in combination with merges and normal
deletions on the same key (other keys are not affected by this)
- Consecutive single deletions are currently not allowed (and older version of
this patch supported this so it could be resurrected if needed)
Test Plan: make all check
Reviewers: yhchiang, sdong, rven, anthony, yoshinorim, igor
Reviewed By: igor
Subscribers: maykov, dhruba, leveldb
Differential Revision: https://reviews.facebook.net/D43179
2015-09-17 11:42:56 -07:00
|
|
|
|
2019-01-31 14:27:21 -08:00
|
|
|
// If the database contains mappings in the range ["begin_key", "end_key"),
|
2016-08-16 08:16:04 -07:00
|
|
|
// erase them. Else do nothing.
|
2017-04-10 15:38:34 -07:00
|
|
|
virtual Status DeleteRange(ColumnFamilyHandle* column_family,
|
|
|
|
const Slice& begin_key, const Slice& end_key) = 0;
|
|
|
|
virtual Status DeleteRange(const Slice& begin_key, const Slice& end_key) = 0;
|
2016-08-16 08:16:04 -07:00
|
|
|
|
|
|
|
// variant that takes SliceParts
|
2017-04-10 15:38:34 -07:00
|
|
|
virtual Status DeleteRange(ColumnFamilyHandle* column_family,
|
|
|
|
const SliceParts& begin_key,
|
|
|
|
const SliceParts& end_key);
|
|
|
|
virtual Status DeleteRange(const SliceParts& begin_key,
|
|
|
|
const SliceParts& end_key);
|
2016-08-16 08:16:04 -07:00
|
|
|
|
2015-02-23 17:49:23 -08:00
|
|
|
// Append a blob of arbitrary size to the records in this batch. The blob will
|
|
|
|
// be stored in the transaction log but not in any other file. In particular,
|
|
|
|
// it will not be persisted to the SST files. When iterating over this
|
|
|
|
// WriteBatch, WriteBatch::Handler::LogData will be called with the contents
|
|
|
|
// of the blob as it is encountered. Blobs, puts, deletes, and merges will be
|
2017-05-17 23:03:54 -07:00
|
|
|
// encountered in the same order in which they were inserted. The blob will
|
2015-02-23 17:49:23 -08:00
|
|
|
// NOT consume sequence number(s) and will NOT increase the count of the batch
|
|
|
|
//
|
|
|
|
// Example application: add timestamps to the transaction log for use in
|
|
|
|
// replication.
|
2017-04-10 15:38:34 -07:00
|
|
|
virtual Status PutLogData(const Slice& blob) = 0;
|
2015-02-23 17:49:23 -08:00
|
|
|
|
|
|
|
// Clear all updates buffered in this batch.
|
|
|
|
virtual void Clear() = 0;
|
|
|
|
|
|
|
|
// Covert this batch into a WriteBatch. This is an abstracted way of
|
|
|
|
// converting any WriteBatchBase(eg WriteBatchWithIndex) into a basic
|
|
|
|
// WriteBatch.
|
|
|
|
virtual WriteBatch* GetWriteBatch() = 0;
|
2015-07-10 20:15:45 -07:00
|
|
|
|
|
|
|
// Records the state of the batch for future calls to RollbackToSavePoint().
|
|
|
|
// May be called multiple times to set multiple save points.
|
|
|
|
virtual void SetSavePoint() = 0;
|
|
|
|
|
|
|
|
// Remove all entries in this batch (Put, Merge, Delete, PutLogData) since the
|
|
|
|
// most recent call to SetSavePoint() and removes the most recent save point.
|
|
|
|
// If there is no previous call to SetSavePoint(), behaves the same as
|
|
|
|
// Clear().
|
|
|
|
virtual Status RollbackToSavePoint() = 0;
|
2017-04-10 15:38:34 -07:00
|
|
|
|
2017-05-03 10:54:07 -07:00
|
|
|
// Pop the most recent save point.
|
|
|
|
// If there is no previous call to SetSavePoint(), Status::NotFound()
|
|
|
|
// will be returned.
|
|
|
|
// Otherwise returns Status::OK().
|
|
|
|
virtual Status PopSavePoint() = 0;
|
|
|
|
|
2017-04-10 15:38:34 -07:00
|
|
|
// Sets the maximum size of the write batch in bytes. 0 means no limit.
|
|
|
|
virtual void SetMaxBytes(size_t max_bytes) = 0;
|
2015-02-23 17:49:23 -08:00
|
|
|
};
|
|
|
|
|
2020-02-20 12:07:53 -08:00
|
|
|
} // namespace ROCKSDB_NAMESPACE
|