/usr/local/lib64/python3.6/site-packages/caffe2/python/operator_test
NameSizeModeActions
__pycache__/-0755rm
activation_ops_test.py96910644editdlrm
adadelta_test.py79320644editdlrm
adagrad_test.py75860644editdlrm
adagrad_test_helper.py51810644editdlrm
adam_test.py215590644editdlrm
affine_channel_op_test.py37840644editdlrm
alias_with_name_test.py9380644editdlrm
apmeter_test.py27380644editdlrm
arg_ops_test.py19170644editdlrm
assert_test.py7970644editdlrm
async_net_barrier_test.py9460644editdlrm
atomic_ops_test.py41040644editdlrm
basic_rnn_test.py47200644editdlrm
batch_box_cox_test.py50800644editdlrm
batch_bucketize_op_test.py37300644editdlrm
batch_moments_op_test.py27950644editdlrm
batch_sparse_to_dense_op_test.py41990644editdlrm
bbox_transform_test.py122580644editdlrm
bisect_percentile_op_test.py62270644editdlrm
blobs_queue_db_test.py32400644editdlrm
boolean_mask_test.py163890644editdlrm
boolean_unmask_test.py17110644editdlrm
box_with_nms_limit_op_test.py87500644editdlrm
bucketize_op_test.py9300644editdlrm
cast_op_test.py16000644editdlrm
ceil_op_test.py8880644editdlrm
channel_backprop_stats_op_test.py21310644editdlrm
channel_shuffle_test.py17940644editdlrm
channel_stats_op_test.py26390644editdlrm
checkpoint_test.py15000644editdlrm
clip_op_test.py19840644editdlrm
clip_tensor_op_test.py20760644editdlrm
collect_and_distribute_fpn_rpn_proposals_op_test.py112690644editdlrm
concat_op_cost_test.py28580644editdlrm
concat_split_op_test.py72660644editdlrm
conditional_test.py9950644editdlrm
conftest.py14460644editdlrm
conv_test.py324730644editdlrm
conv_transpose_test.py159450644editdlrm
copy_ops_test.py73740644editdlrm
copy_rows_to_tensor_op_test.py25260644editdlrm
cosine_embedding_criterion_op_test.py19530644editdlrm
counter_ops_test.py33480644editdlrm
crf_test.py53150644editdlrm
cross_entropy_ops_test.py100850644editdlrm
ctc_beam_search_decoder_op_test.py51970644editdlrm
ctc_greedy_decoder_op_test.py47430644editdlrm
cudnn_recurrent_test.py58170644editdlrm
dataset_ops_test.py238470644editdlrm
data_couple_op_test.py8580644editdlrm
decay_adagrad_test.py26940644editdlrm
deform_conv_test.py192760644editdlrm
dense_vector_to_id_list_op_test.py20440644editdlrm
depthwise_3x3_conv_test.py18630644editdlrm
detectron_keypoints.py79730644editdlrm
distance_op_test.py43510644editdlrm
dropout_op_test.py29710644editdlrm
duplicate_operands_test.py7340644editdlrm
elementwise_linear_op_test.py13820644editdlrm
elementwise_logical_ops_test.py46170644editdlrm
elementwise_ops_test.py333540644editdlrm
elementwise_op_broadcast_test.py174660644editdlrm
emptysample_ops_test.py19770644editdlrm
enforce_finite_op_test.py12860644editdlrm
ensure_clipped_test.py15050644editdlrm
ensure_cpu_output_op_test.py12440644editdlrm
erf_op_test.py7490644editdlrm
expand_op_test.py21090644editdlrm
fc_operator_test.py37200644editdlrm
feature_maps_ops_test.py214920644editdlrm
filler_ops_test.py84760644editdlrm
find_op_test.py13160644editdlrm
flatten_op_test.py9220644editdlrm
flexible_top_k_test.py26090644editdlrm
floor_op_test.py8940644editdlrm
fused_nbit_rowwise_conversion_ops_test.py140770644editdlrm
fused_nbit_rowwise_test_helper.py26930644editdlrm
gather_ops_test.py92160644editdlrm
gather_ranges_op_test.py91250644editdlrm
given_tensor_byte_string_to_uint8_fill_op_test.py13920644editdlrm
given_tensor_fill_op_test.py15030644editdlrm
glu_op_test.py12120644editdlrm
group_conv_test.py28700644editdlrm
group_norm_op_test.py52520644editdlrm
gru_test.py129320644editdlrm
heatmap_max_keypoint_op_test.py47700644editdlrm
histogram_test.py30970644editdlrm
hsm_test.py94560644editdlrm
hyperbolic_ops_test.py14720644editdlrm
im2col_col2im_test.py43110644editdlrm
image_input_op_test.py173450644editdlrm
index_hash_ops_test.py28850644editdlrm
index_ops_test.py45970644editdlrm
instance_norm_test.py99170644editdlrm
integral_image_ops_test.py34190644editdlrm
jsd_ops_test.py10440644editdlrm
key_split_ops_test.py12890644editdlrm
lars_test.py13540644editdlrm
layer_norm_op_test.py149830644editdlrm
leaky_relu_test.py56390644editdlrm
learning_rate_adaption_op_test.py28370644editdlrm
learning_rate_op_test.py86520644editdlrm
lengths_pad_op_test.py16250644editdlrm
lengths_reducer_fused_nbit_rowwise_ops_test.py154950644editdlrm
lengths_tile_op_test.py13320644editdlrm
lengths_top_k_ops_test.py23710644editdlrm
length_split_op_test.py48680644editdlrm
listwise_l2r_operator_test.py87400644editdlrm
load_save_test.py332410644editdlrm
locally_connected_op_test.py77610644editdlrm
loss_ops_test.py9020644editdlrm
lpnorm_op_test.py27250644editdlrm
map_ops_test.py22490644editdlrm
margin_ranking_criterion_op_test.py18160644editdlrm
math_ops_test.py16030644editdlrm
matmul_op_test.py100960644editdlrm
mean_op_test.py14690644editdlrm
merge_id_lists_op_test.py29890644editdlrm
mkl_conv_op_test.py15470644editdlrm
mkl_packed_fc_op_test.py26470644editdlrm
mod_op_test.py14590644editdlrm
moments_op_test.py17220644editdlrm
momentum_sgd_test.py64800644editdlrm
mpi_test.py81540644editdlrm
mul_gradient_benchmark.py15090644editdlrm
negate_gradient_op_test.py15180644editdlrm
ngram_ops_test.py23270644editdlrm
normalize_op_test.py16790644editdlrm
numpy_tile_op_test.py19240644editdlrm
one_hot_ops_test.py74780644editdlrm
onnx_while_test.py30700644editdlrm
order_switch_test.py13060644editdlrm
pack_ops_test.py126340644editdlrm
pack_rnn_sequence_op_test.py28910644editdlrm
pad_test.py13770644editdlrm
partition_ops_test.py68380644editdlrm
percentile_op_test.py44270644editdlrm
piecewise_linear_transform_test.py61870644editdlrm
pooling_test.py165080644editdlrm
prepend_dim_test.py15050644editdlrm
python_op_test.py13120644editdlrm
quantile_test.py32760644editdlrm
rand_quantization_op_speed_test.py31280644editdlrm
rank_loss_operator_test.py57520644editdlrm
rebatching_queue_test.py90470644editdlrm
record_queue_test.py31250644editdlrm
recurrent_network_test.py140480644editdlrm
recurrent_net_executor_test.py109220644editdlrm
reduce_ops_test.py173410644editdlrm
reduction_ops_test.py46640644editdlrm
reshape_ops_test.py82110644editdlrm
resize_op_test.py94170644editdlrm
rmac_regions_op_test.py31780644editdlrm
rms_norm_op_test.py13250644editdlrm
rnn_cell_test.py597070644editdlrm
roi_align_rotated_op_test.py75670644editdlrm
rowwise_counter_test.py22050644editdlrm
scale_op_test.py21770644editdlrm
segment_ops_test.py257450644editdlrm
self_binning_histogram_test.py129150644editdlrm
selu_op_test.py32320644editdlrm
sequence_ops_test.py160000644editdlrm
shape_inference_test.py257080644editdlrm
sinusoid_position_encoding_op_test.py23080644editdlrm
softmax_ops_test.py236850644editdlrm
softplus_op_test.py5160644editdlrm
sparse_dropout_with_replacement_op_test.py28850644editdlrm
sparse_gradient_checker_test.py12940644editdlrm
sparse_itemwise_dropout_with_replacement_op_test.py29130644editdlrm
sparse_lengths_sum_benchmark.py41590644editdlrm
sparse_lp_regularizer_test.py25530644editdlrm
sparse_normalize_test.py31360644editdlrm
sparse_ops_test.py34690644editdlrm
sparse_to_dense_mask_op_test.py36930644editdlrm
spatial_bn_op_test.py201820644editdlrm
specialized_segment_ops_test.py117750644editdlrm
split_op_cost_test.py86450644editdlrm
square_root_divide_op_test.py21790644editdlrm
stats_ops_test.py17890644editdlrm
stats_put_ops_test.py65960644editdlrm
storm_test.py65070644editdlrm
string_ops_test.py41540644editdlrm
text_file_reader_test.py25170644editdlrm
thresholded_relu_op_test.py23230644editdlrm
tile_op_test.py38870644editdlrm
top_k_test.py91130644editdlrm
torch_integration_test.py399410644editdlrm
transpose_op_test.py27220644editdlrm
trigonometric_op_test.py17150644editdlrm
unique_ops_test.py22550644editdlrm
unique_uniform_fill_op_test.py13350644editdlrm
unsafe_coalesce_test.py29400644editdlrm
upsample_op_test.py73080644editdlrm
utility_ops_test.py150540644editdlrm
video_input_op_test.py105030644editdlrm
weighted_multi_sample_test.py19970644editdlrm
weighted_sample_test.py27390644editdlrm
weighted_sum_test.py30520644editdlrm
weight_scale_test.py20570644editdlrm
wngrad_test.py82790644editdlrm
__init__.py00644editdlrm
Edit: /usr/local/lib64/python3.6/site-packages/caffe2/python/operator_test/softmax_ops_test.py (23685B)
from caffe2.python import core, workspace from hypothesis import given, settings import caffe2.python.hypothesis_test_util as hu import caffe2.python.serialized_test.serialized_test_util as serial import hypothesis.strategies as st import numpy as np import unittest class TestSoftmaxOps(serial.SerializedTestCase): @serial.given(n=st.sampled_from([0, 2, 4, 71, 103]), D=st.sampled_from([0, 4, 8, 64, 79, 256, 333]), engine=st.sampled_from([None, 'CUDNN']), **hu.gcs) def test_softmax(self, n, D, engine, gc, dc): # n = number of examples, D = |labels| # Initialize X and add 1e-2 for numerical stability X = np.random.rand(n, D).astype(np.float32) X = X + 1e-2 # Reference implementation of cross entropy with soft labels def label_softmax(X): probs = np.zeros((n, D)) rowmax = np.zeros(n) if D == 0: return [probs] for i in range(n): rowmax[i] = max(X[i, ]) # We need to subtract the max to avoid numerical issues probs[i] = X[i] - rowmax[i] exps = np.exp(probs[i, ]) norm = sum(exps) probs[i, ] = exps / norm return [probs] op = core.CreateOperator( "Softmax", ["X"], ["probs"], engine=engine ) self.assertReferenceChecks( device_option=gc, op=op, inputs=[X], reference=label_softmax, ) @given(n=st.sampled_from([0, 2, 4, 71, 103, 555, 751, 1201]), D=st.sampled_from([0, 4, 8, 64, 79, 256, 333, 1000]), engine=st.sampled_from([None, 'CUDNN']), **hu.gcs) @settings(deadline=10000) def test_softmax_grad(self, n, D, engine, gc, dc): # n = number of examples, D = |labels| # Initialize X and add 1e-2 for numerical stability Y = np.random.rand(n, D).astype(np.float32) dY = np.random.rand(n, D).astype(np.float32) Y = Y + 1e-2 # Reference implementation of cross entropy with soft labels def label_softmax_grad(X, dY): dX = Y * 0.0 for i in range(n): d = np.dot(Y[i, :], dY[i, :]) dX[i, :] = Y[i, :] * (dY[i, :] - d) return [dX] op = core.CreateOperator( "SoftmaxGradient", ["Y", "dY"], ["dX"], engine=engine ) self.assertReferenceChecks( device_option=gc, op=op, inputs=[Y, dY], reference=label_softmax_grad, ) @given(axis=st.integers(min_value=1, max_value=4), engine=st.sampled_from([None, 'CUDNN']), **hu.gcs) def test_softmax_axis(self, axis, engine, gc, dc): np.random.seed(1) X = np.random.randn(1, 2, 3, 2, 1).astype(np.float32) X = X + 1e-2 def prod(xs): p = 1 for x in xs: p *= x return p N = prod(list(X.shape)[:axis]) D = prod(list(X.shape)[axis:]) # Reference implementation of cross entropy with soft labels def label_softmax(X): X_ = X.reshape(N, D) probs = np.zeros((N, D)) rowmax = np.zeros(N) for i in range(N): rowmax[i] = max(X_[i, ]) # We need to subtract the max to avoid numerical issues probs[i] = X_[i] - rowmax[i] exps = np.exp(probs[i, ]) norm = sum(exps) probs[i, ] = exps / norm return [probs.reshape(*X.shape)] op = core.CreateOperator( "Softmax", ["X"], ["probs"], axis=axis, engine=engine, ) self.assertReferenceChecks( device_option=gc, op=op, inputs=[X], reference=label_softmax, ) self.assertGradientChecks( gc, op, [X], 0, [0], stepsize=1e-4, threshold=1e-2) @given(n=st.integers(2, 10), D=st.integers(4, 16), only_loss=st.booleans(), **hu.gcs) @settings(deadline=10000) def test_softmax_with_loss(self, n, D, gc, only_loss, dc): # n = number of examples, D = |labels| # Initialize X and add 1e-2 for numerical stability np.random.seed(2603) X = np.random.rand(n, D).astype(np.float32) X = X + 1e-2 # Initialize label label = (np.random.rand(n) * D).astype(np.int32) # Reference implementation of cross entropy with soft labels def label_softmax_crossent(X, label): probs = np.zeros((n, D)) rowmax = np.zeros(n) for i in range(n): rowmax[i] = max(X[i, ]) # We need to subtract the max to avoid numerical issues probs[i] = X[i] - rowmax[i] exps = np.exp(probs[i, ]) norm = sum(exps) probs[i, ] = exps / norm label_xent = [-np.log(max(probs[i][label[i]], 1e-20)) for i in range(n)] avgloss = np.sum(label_xent) / float(n) return (probs, avgloss) op = core.CreateOperator( "SoftmaxWithLoss", ["X", "label"], ["probs", "avgloss"], only_loss=only_loss, ) self.assertReferenceChecks( device_option=gc, op=op, inputs=[X, label], reference=label_softmax_crossent, ) self.assertGradientChecks( gc, op, [X, label], 0, [1], stepsize=1e-4, threshold=1e-2) @given( n=st.integers(2, 5), D=st.integers(4, 16), only_loss=st.booleans(), label_prob=st.booleans(), **hu.gcs ) @settings(deadline=10000) def test_softmax_with_loss_axis_2( self, n, D, only_loss, label_prob, gc, dc ): np.random.seed(2603) X = np.random.rand(n, n, D).astype(np.float32) X = X + 1e-2 if label_prob: label = np.random.rand(n, n, D).astype(np.float32) label /= label.sum(axis=2, keepdims=True) else: label = (np.random.rand(n, n) * D).astype(np.int32) # Reference implementation of cross entropy with soft labels def label_softmax_crossent(X, label): probs = np.zeros((n, n, D)) rowmax = np.zeros((n, n)) for i in range(n): for j in range(n): rowmax[i, j] = max(X[i, j, ]) # We need to subtract the max to avoid numerical issues probs[i, j] = X[i, j] - rowmax[i, j] exps = np.exp(probs[i, j, ]) norm = sum(exps) probs[i, j, ] = exps / norm label_xent = 0 for i in range(n): for j in range(n): if label_prob: for k in range(D): label_xent += ( -np.log(max(probs[i, j, k], 1e-20)) * label[i, j, k] ) else: label_xent += -np.log(max(probs[i, j, label[i, j]], 1e-20)) avgloss = label_xent / float(n * n) return (probs, avgloss) op = core.CreateOperator( "SoftmaxWithLoss", ["X", "label"], ["probs", "avgloss"], only_loss=only_loss, label_prob=label_prob, axis=2, ) self.assertReferenceChecks( device_option=gc, op=op, inputs=[X, label], reference=label_softmax_crossent, ) self.assertGradientChecks( gc, op, [X, label], 0, [1], stepsize=1e-4, threshold=1e-2) @unittest.skipIf(not workspace.has_gpu_support, "No gpu support") @given(**hu.gcs_gpu_only) def test_softmax_with_loss_large(self, gc, dc): np.random.seed(2603) for n in [32]: for D in [1000, 2000, 20000]: # n = number of examples, D = |labels| # Initialize X and add 1e-2 for numerical stability X = np.random.rand(n, D).astype(np.float32) X = X + 1e-2 # Initialize label label = (np.random.rand(n) * D).astype(np.int32) # Reference implementation of cross entropy with soft labels def label_softmax_crossent(X, label): probs = np.zeros((n, D)) rowmax = np.zeros(n) for i in range(n): rowmax[i] = max(X[i, ]) # We need to subtract the max to avoid numerical issues probs[i] = X[i] - rowmax[i] exps = np.exp(probs[i, ]) norm = sum(exps) probs[i, ] = exps / norm label_xent = [-np.log(max(probs[i][label[i]], 1e-20)) for i in range(n)] avgloss = np.sum(label_xent) / float(n) return (probs, avgloss) op = core.CreateOperator( "SoftmaxWithLoss", ["X", "label"], ["probs", "avgloss"] ) self.assertReferenceChecks( device_option=gc, op=op, inputs=[X, label], reference=label_softmax_crossent, ) @given(n=st.integers(2, 10), D=st.integers(4, 16), **hu.gcs) @settings(deadline=None) def test_softmax_with_loss_label_prob(self, n, D, gc, dc): # n = number of examples, D = |labels| # Initialize X and add 1e-2 for numerical stability np.random.seed(2603) X = np.random.rand(n, D).astype(np.float32) X = X + 1e-2 # Initialize label label = np.random.rand(D, n).astype(np.float32) # normalize labels to sum to 1 label /= np.sum(label, axis=0) label = label.transpose() # Reference implementation of cross entropy with soft labels def label_softmax_crossent(X, label): probs = np.zeros((n, D)) rowmax = np.zeros(n) for i in range(n): rowmax[i] = max(X[i, ]) # We need to subtract the max to avoid numerical issues probs[i] = X[i] - rowmax[i] exps = np.exp(probs[i, ]) norm = sum(exps) probs[i, ] = exps / norm label_xent = np.zeros(X.shape) for i in range(n): for j in range(D): label_xent[i][j] = -np.log( max(probs[i, j], 1e-20)) * label[i, j] avgloss = np.sum(label_xent) / float(n) return (probs, avgloss) op = core.CreateOperator( "SoftmaxWithLoss", ["X", "label"], ["probs", "avgloss"], label_prob=1 ) self.assertReferenceChecks( device_option=gc, op=op, inputs=[X, label], reference=label_softmax_crossent, ) self.assertGradientChecks( gc, op, [X, label], 0, [1], stepsize=1e-4, threshold=1e-2) @given( n=st.integers(2, 10), D=st.integers(4, 16), only_loss=st.booleans(), **hu.gcs) @settings(deadline=None) def test_softmax_with_loss_weighted(self, n, D, only_loss, gc, dc): # n = number of examples, D = |labels| # Initialize X and add 1e-2 for numerical stability np.random.seed(2603) X = np.random.rand(n, D).astype(np.float32) X = X + 1e-2 # Initialize label label = (np.random.rand(n) * D).astype(np.int32) # Init weights (weight by sample) weights = np.random.rand(n).astype(np.float32) # Reference implementation of cross entropy with soft labels def label_softmax_crossent_weighted(X, label, weights): probs = np.zeros((n, D)) rowmax = np.zeros(n) for i in range(n): rowmax[i] = max(X[i, ]) # We need to subtract the max to avoid numerical issues probs[i] = X[i] - rowmax[i] exps = np.exp(probs[i, ]) norm = sum(exps) probs[i, ] = exps / norm label_xent = [-weights[i] * np.log(max(probs[i][label[i]], 1e-20)) for i in range(n)] avgloss = np.sum(label_xent) / sum(weights) return (probs, avgloss) op = core.CreateOperator( "SoftmaxWithLoss", ["X", "label", "weights"], ["probs", "avgloss"], only_loss=only_loss, ) self.assertReferenceChecks( device_option=gc, op=op, inputs=[X, label, weights], reference=label_softmax_crossent_weighted, ) self.assertGradientChecks( gc, op, [X, label, weights], 0, [1], stepsize=1e-4, threshold=1e-2) @given(n=st.integers(2, 10), D=st.integers(4, 16), **hu.gcs) @settings(deadline=None) def test_softmax_with_loss_label_prob_weighted(self, n, D, gc, dc): # n = number of examples, D = |labels| # Initialize X and add 1e-2 for numerical stability X = np.random.rand(n, D).astype(np.float32) X = X + 1e-2 # Initialize label label = np.random.rand(D, n).astype(np.float32) # normalize labels to sum to 1 label /= np.sum(label, axis=0) label = label.transpose() # Init weights (weight by sample) weights = np.random.rand(n).astype(np.float32) # Reference implementation of cross entropy with soft labels def label_softmax_crossent_weighted(X, label, weights): probs = np.zeros((n, D)) rowmax = np.zeros(n) for i in range(n): rowmax[i] = max(X[i, ]) # We need to subtract the max to avoid numerical issues probs[i] = X[i] - rowmax[i] exps = np.exp(probs[i, ]) norm = sum(exps) probs[i, ] = exps / norm label_xent = np.zeros(X.shape) for i in range(n): for j in range(D): label_xent[i][j] = -np.log( max(probs[i, j], 1e-20)) * label[i, j] * weights[i] avgloss = np.sum(label_xent) / sum(weights) return (probs, avgloss) op = core.CreateOperator( "SoftmaxWithLoss", ["X", "label", "weights"], ["probs", "avgloss"], label_prob=1, ) self.assertReferenceChecks( device_option=gc, op=op, inputs=[X, label, weights], reference=label_softmax_crossent_weighted, ) self.assertGradientChecks( gc, op, [X, label, weights], 0, [1], stepsize=1e-4, threshold=1e-2) @given(n=st.integers(2, 5), D=st.integers(2, 4), weighted=st.booleans(), **hu.gcs) @settings(deadline=None, max_examples=50) def test_spatial_softmax_with_loss(self, n, D, weighted, gc, dc): # n = number of examples, D = |labels| # Initialize X and add 1e-2 for numerical stability W = 18 H = 12 np.random.seed(2603) X = np.random.rand(n, D, H, W).astype(np.float32) X = X + 1e-2 weighted = True weights = None if weighted: weights = np.random.rand(n, H, W).astype(np.float32) # Initialize label. Some of the labels are (-1), i.e "DONT CARE" label = (np.random.rand(n, H, W) * (D + 1)).astype(np.int32) - 1 def label_softmax_crossent_spatial(X, label, weights=None): probs = np.zeros((n, D, H, W)) rowmax = np.zeros((n, H, W)) label_xent = np.zeros((n, H, W)) for i in range(n): for x in range(W): for y in range(H): rowmax[i, y, x] = max(X[i, :, y, x]) # We need to subtract the max to avoid numerical issues probs[i, :, y, x] = X[i, :, y, x] - rowmax[i, y, x] exps = np.exp(probs[i, :, y, x]) probs[i, :, y, x] = exps / sum(exps) label_xent[:, y, x] = \ [-np.log(max(probs[j, label[i, y, x], y, x], 1e-20)) for j in range(n)] total_xent = 0.0 total_weight = 0.0 for y in range(H): for x in range(W): for i in range(n): l = label[i, y, x] if (l != (-1)): w = 1.0 if weights is None else weights[i, y, x] total_xent += \ -np.log(max(probs[i, l, y, x], 1e-20)) * w total_weight += w print("Total weight {}".format(total_weight)) return (probs, total_xent / total_weight) op = core.CreateOperator( "SpatialSoftmaxWithLoss", ["X", "label"] + ([] if weights is None else ["weights"]), ["probs", "avgloss"], ) inputs = [X, label] + ([] if weights is None else [weights]) self.assertReferenceChecks( device_option=gc, op=op, inputs=inputs, reference=label_softmax_crossent_spatial, ) self.assertGradientChecks( gc, op, inputs, 0, [1], stepsize=1e-4, threshold=1e-2) @given(n=st.integers(4, 5), D=st.integers(3, 4), weighted=st.booleans(), **hu.gcs) def test_spatial_softmax_with_loss_allignore(self, n, D, weighted, gc, dc): # n = number of examples, D = |labels| # Initialize X and add 1e-2 for numerical stability W = 18 H = 12 np.random.seed(2603) X = np.random.rand(n, D, H, W).astype(np.float32) X = X + 1e-2 weighted = True weights = None if weighted: weights = np.random.rand(n, H, W).astype(np.float32) # Initialize label. All labels as "DONT CARE" label = np.zeros((n, H, W)).astype(np.int32) - 1 print(label) def label_softmax_crossent_spatial(X, label, weights=None): probs = np.zeros((n, D, H, W)) rowmax = np.zeros((n, H, W)) label_xent = np.zeros((n, H, W)) for i in range(n): for x in range(W): for y in range(H): rowmax[i, y, x] = max(X[i, :, y, x]) # We need to subtract the max to avoid numerical issues probs[i, :, y, x] = X[i, :, y, x] - rowmax[i, y, x] exps = np.exp(probs[i, :, y, x]) probs[i, :, y, x] = exps / sum(exps) label_xent[:, y, x] = \ [-np.log(max(probs[j, label[i, y, x], y, x], 1e-20)) for j in range(n)] return (probs, 0.0) op = core.CreateOperator( "SpatialSoftmaxWithLoss", ["X", "label"] + ([] if weights is None else ["weights"]), ["probs", "avgloss"], ) inputs = [X, label] + ([] if weights is None else [weights]) self.assertReferenceChecks( device_option=gc, op=op, inputs=inputs, reference=label_softmax_crossent_spatial, ) @given(n=st.integers(4, 5), D=st.integers(3, 4), weighted=st.booleans(), **hu.gcs) def test_softmax_with_loss_zero_weight(self, n, D, weighted, gc, dc): # n = number of examples, D = |labels| # Initialize X and add 1e-2 for numerical stability np.random.seed(2603) X = np.random.rand(n, D).astype(np.float32) X = X + 1e-2 weights = np.zeros(n).astype(np.float32) # Initialize label label = (np.random.rand(n) * D).astype(np.int32) def label_softmax_crossent(X, label, weights=None): probs = np.zeros((n, D)) rowmax = np.zeros((n)) for i in range(n): rowmax[i] = max(X[i, ]) # We need to subtract the max to avoid numerical issues probs[i] = X[i] - rowmax[i] exps = np.exp(probs[i, ]) norm = sum(exps) probs[i, ] = exps / norm return (probs, 0.0) op = core.CreateOperator( "SoftmaxWithLoss", ["X", "label", "weights"], ["probs", "avgloss"] ) inputs = [X, label] + ([] if weights is None else [weights]) self.assertReferenceChecks( device_option=gc, op=op, inputs=inputs, reference=label_softmax_crossent, ) @unittest.skipIf(not workspace.has_gpu_support, "No gpu support") def test_compare_cpugpu(self): ''' Additional test that checks CPU and GPU returns same values with larger examples. This is mainly to test the more complex GPU implementation is correct. ''' from caffe2.proto import caffe2_pb2 for _j in range(3): gpuop = core.CreateOperator( "SpatialSoftmaxWithLoss", ["X_gpu", "label_gpu"], ["probs_gpu", "avgloss_gpu"], device_option=core.DeviceOption(workspace.GpuDeviceType, 0) ) cpuop = core.CreateOperator( "SpatialSoftmaxWithLoss", ["X_cpu", "label_cpu"], ["probs_cpu", "avgloss_cpu"], device_option=core.DeviceOption(caffe2_pb2.CPU) ) n = 8 D = 4 W = 64 + int(np.random.rand(1) * 1024) H = 64 + int(np.random.rand(1) * 1024) print("W: {} H: {}".format(W, H)) X = np.random.rand(n, D, H, W).astype(np.float32) X = X + 1e-2 # Initialize label. Some of the labels are (-1), i.e "DONT CARE" label = (np.random.rand(n, H, W) * (D + 1)).astype(np.int32) - 1 gpu0 = core.DeviceOption(workspace.GpuDeviceType, 0) workspace.FeedBlob("X_cpu", X) workspace.FeedBlob("label_cpu", label) workspace.FeedBlob("X_gpu", X, device_option=gpu0) workspace.FeedBlob("label_gpu", label, device_option=gpu0) workspace.RunOperatorOnce(gpuop) workspace.RunOperatorOnce(cpuop) probs_gpu = workspace.FetchBlob("probs_gpu") probs_cpu = workspace.FetchBlob("probs_cpu") loss_gpu = workspace.FetchBlob("avgloss_gpu") loss_cpu = workspace.FetchBlob("avgloss_cpu") np.testing.assert_allclose(probs_gpu, probs_cpu, rtol=1e-4) np.testing.assert_allclose(loss_gpu, loss_cpu, rtol=1e-1) if __name__ == "__main__": import unittest import random random.seed(2603) unittest.main()